{"as_of":"2026-08-10T11:40:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:dd9412805feea0d56e5b1ae8c01c07e62b1b17dafd6be78b3e55be1acf1ff038","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":6,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":6,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":6,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":6,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T00:50:40.370322Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T03:49:29.875637Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2411.15708","last_updated":"2024-11-24T04:26:04Z","snapshot_observed_at":"2026-08-10T03:12:30.342926Z","submitted_at":"2024-11-24T04:26:04Z","title":"LLaMA-MoE v2: Exploring Sparsity of LLaMA from Perspective of Mixture-of-Experts with Post-Training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15708","snapshot_observed_at":"2026-08-09T00:50:40.370322Z","title":"Llama-moe v2: Exploring sparsity of llama from perspec- tive of mixture-of-experts with post-training,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03772","last_updated":"2025-03-20T06:38:41Z","snapshot_observed_at":"2026-08-10T03:12:03.704635Z","submitted_at":"2025-02-06T04:17:02Z","title":"A Retrospective Systematic Study on Hierarchical Sparse Query Transformer-assisted Ultrasound Screening for Early Hepatocellular Carcinoma","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-09T00:50:40.370322Z"},"links":{"cited_paper":"/paper/2411.15708","citing_paper":"/paper/2502.03772"},"observation_digest":"sha256:d3f0e673d8b9d4df5ae2587f7f7128116140df8a50d8483f99e70643bc1e66ee","observation_id":"5793c7b5-7f93-4846-9b7a-f8fbe9f3ce0e","resolution":{"observed_at":"2026-08-09T00:50:40.370322Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15708","last_updated":"2024-11-24T04:26:04Z","snapshot_observed_at":"2026-08-10T03:12:30.342926Z","submitted_at":"2024-11-24T04:26:04Z","title":"LLaMA-MoE v2: Exploring Sparsity of LLaMA from Perspective of Mixture-of-Experts with Post-Training","version":1},"cited_work":{"arxiv_id":"2411.15708","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.15708","snapshot_observed_at":"2026-07-04T03:49:29.875637Z","title":"Llama-moe v2: Exploring sparsity of llama from perspective of mixture-of-experts with post-training","venue":null,"work_id":"edde3884-dd00-485b-b770-36999dfdef84","year":2024},"citing_paper":{"arxiv_id":"2502.04416","last_updated":"2026-04-23T00:51:26Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-06T14:05:30Z","title":"Analytical FFN-to-MoE Restructuring via Activation Pattern Analysis","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-23T04:08:29.089438Z"},"links":{"cited_paper":"/paper/2411.15708","citing_paper":"/paper/2502.04416"},"observation_digest":"sha256:bdf5764b142ce5a48d0ca791ccddcf613e964a6c6c75e241ec02e27b6474fc8b","observation_id":"80ce4cae-154d-4955-abb1-e94867c992a2","resolution":{"observed_at":"2026-05-23T04:12:31.083706Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15708","last_updated":"2024-11-24T04:26:04Z","snapshot_observed_at":"2026-08-10T03:12:30.342926Z","submitted_at":"2024-11-24T04:26:04Z","title":"LLaMA-MoE v2: Exploring Sparsity of LLaMA from Perspective of Mixture-of-Experts with Post-Training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15708","snapshot_observed_at":"2026-08-02T21:47:26.933964Z","title":"Llama-moe v2: Exploring sparsity of llama from perspective of mixture-of-experts with post- training.arXiv preprint arXiv:2411.15708, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.19213","last_updated":"2026-06-05T08:10:24Z","snapshot_observed_at":"2026-08-10T04:10:51.935941Z","submitted_at":"2026-02-22T14:48:42Z","title":"SegMoTE: Token-Level Mixture of Experts for Medical Image Segmentation","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-02T21:47:26.933964Z"},"links":{"cited_paper":"/paper/2411.15708","citing_paper":"/paper/2602.19213"},"observation_digest":"sha256:be59175f90eb7f6586ff257e7937c16e4c744fc4160a920079f227122a1bdd2f","observation_id":"a3cb4d6b-c2dd-413a-b985-6b5e5fd98a76","resolution":{"observed_at":"2026-08-02T21:47:26.933964Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15708","last_updated":"2024-11-24T04:26:04Z","snapshot_observed_at":"2026-08-10T03:12:30.342926Z","submitted_at":"2024-11-24T04:26:04Z","title":"LLaMA-MoE v2: Exploring Sparsity of LLaMA from Perspective of Mixture-of-Experts with Post-Training","version":1},"cited_work":{"arxiv_id":"2411.15708","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.15708","snapshot_observed_at":"2026-07-04T03:49:29.875637Z","title":"Llama-moe v2: Exploring sparsity of llama from perspective of mixture-of-experts with post-training","venue":null,"work_id":"edde3884-dd00-485b-b770-36999dfdef84","year":2024},"citing_paper":{"arxiv_id":"2605.05225","last_updated":"2026-06-05T13:47:55Z","snapshot_observed_at":"2026-08-02T06:50:35.896043Z","submitted_at":"2026-04-19T07:25:39Z","title":"MACS: Modality-Aware Capacity Scaling for Efficient Multimodal MoE Inference","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-05-11T01:29:26.298131Z"},"links":{"cited_paper":"/paper/2411.15708","citing_paper":"/paper/2605.05225"},"observation_digest":"sha256:1c639cf00b548fcc57d91728672d22c0a7adeee7d6ef3168a340f634b8e5fbdc","observation_id":"b22427e4-2144-442c-a260-4de405b9b4ab","resolution":{"observed_at":"2026-05-11T01:45:52.062735Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15708","last_updated":"2024-11-24T04:26:04Z","snapshot_observed_at":"2026-08-10T03:12:30.342926Z","submitted_at":"2024-11-24T04:26:04Z","title":"LLaMA-MoE v2: Exploring Sparsity of LLaMA from Perspective of Mixture-of-Experts with Post-Training","version":1},"cited_work":{"arxiv_id":"2411.15708","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.15708","snapshot_observed_at":"2026-07-04T03:49:29.875637Z","title":"Llama-moe v2: Exploring sparsity of llama from perspective of mixture-of-experts with post-training","venue":null,"work_id":"edde3884-dd00-485b-b770-36999dfdef84","year":2024},"citing_paper":{"arxiv_id":"2605.24681","last_updated":"2026-05-23T17:35:43Z","snapshot_observed_at":"2026-07-06T23:34:43.210149Z","submitted_at":"2026-05-23T17:35:43Z","title":"Mix-MoE: Improving Multilingual Machine Translation of Large Language Models through Mixed MoEs","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-06-30T13:22:40.017922Z"},"links":{"cited_paper":"/paper/2411.15708","citing_paper":"/paper/2605.24681"},"observation_digest":"sha256:cf999e0c3b8af44b37bbf9666fdae871ef0002f626be1549634190785624c2dd","observation_id":"b53b3918-19bd-483b-a301-bf0f97f2f32a","resolution":{"observed_at":"2026-06-30T13:24:40.055476Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15708","last_updated":"2024-11-24T04:26:04Z","snapshot_observed_at":"2026-08-10T03:12:30.342926Z","submitted_at":"2024-11-24T04:26:04Z","title":"LLaMA-MoE v2: Exploring Sparsity of LLaMA from Perspective of Mixture-of-Experts with Post-Training","version":1},"cited_work":{"arxiv_id":"2411.15708","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.15708","snapshot_observed_at":"2026-07-04T03:49:29.875637Z","title":"Llama-moe v2: Exploring sparsity of llama from perspective of mixture-of-experts with post-training","venue":null,"work_id":"edde3884-dd00-485b-b770-36999dfdef84","year":2024},"citing_paper":{"arxiv_id":"2606.20945","last_updated":"2026-06-23T11:48:31Z","snapshot_observed_at":"2026-07-06T23:56:02.397254Z","submitted_at":"2026-06-18T21:18:42Z","title":"Grouped Query Experts: Mixture-of-Experts on GQA Self-Attention","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-26T17:40:06.036019Z"},"links":{"cited_paper":"/paper/2411.15708","citing_paper":"/paper/2606.20945"},"observation_digest":"sha256:bd419fd0ddd9c3b3d75ab422dba81f734609d365e0ec6c51bdfbe06e29fae425","observation_id":"e3a068bb-1cad-493a-9633-59777f5429fa","resolution":{"observed_at":"2026-07-04T03:49:29.877090Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2411.15708/citation-record","integrity":"/paper/2411.15708/integrity","json":"/paper/2411.15708/citation-record.json","paper":"/paper/2411.15708"},"outbound":[],"paper":{"arxiv_id":"2411.15708","last_updated":"2024-11-24T04:26:04Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-10T03:12:30.342926Z","submitted_at":"2024-11-24T04:26:04Z","title":"LLaMA-MoE v2: Exploring Sparsity of LLaMA from Perspective of Mixture-of-Experts with Post-Training"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 6 inbound Pith citation observations for arXiv:2411.15708."}