{"as_of":"2026-08-08T20:45:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:537a1eb0f6ec1c71de41aadd43625a4b3dd2cf60ef01f51edadef9036432df82","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":18,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":18,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":18,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":18,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:33:32.170655Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T13:19:50.957441Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-08-07T14:33:32.170655Z","title":"Milebench: Benchmarking mllms in long context.arXiv preprint arXiv:2404.18532, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.18605","last_updated":"2025-05-24T08:59:28Z","snapshot_observed_at":"2026-08-08T00:45:48.179953Z","submitted_at":"2025-05-24T08:59:28Z","title":"Rethinking Causal Mask Attention for Vision-Language Inference","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T14:33:32.170655Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2505.18605"},"observation_digest":"sha256:420c8e61c8221fec37b6ada312b95960adbc96a0a194b1cf9a53d93a80bc1109","observation_id":"54251bc8-7ba2-4f9c-bd6f-1ac1e72d4c61","resolution":{"observed_at":"2026-08-07T14:33:32.170655Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-08-07T11:05:19.523326Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04280","last_updated":"2025-06-04T04:21:32Z","snapshot_observed_at":"2026-08-07T21:26:08.023036Z","submitted_at":"2025-06-04T04:21:32Z","title":"Evaluating MLLMs with Multimodal Multi-image Reasoning Benchmark","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:05:19.523326Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2506.04280"},"observation_digest":"sha256:188590c3a0704d3aff1bde7af0410318c137380366073b34140b0bf6f78c3f9f","observation_id":"0d9ba4f4-81b9-478f-84f5-54dd0f3419fd","resolution":{"observed_at":"2026-08-07T11:05:19.523326Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-08-07T06:02:30.430076Z","title":"H., Yu, F., Wan, X., and Wang, B","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.06279","last_updated":"2025-06-06T17:59:06Z","snapshot_observed_at":"2026-08-08T03:56:06.075559Z","submitted_at":"2025-06-06T17:59:06Z","title":"CoMemo: LVLMs Need Image Context with Image Memory","version":1},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-08-07T06:02:30.430076Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2506.06279"},"observation_digest":"sha256:b7f29bae9fc19a4a7646de9f95e9548cad722e0d3a980fdaeb12ebd8dd44603d","observation_id":"d554ce3f-0191-4126-aca7-caf3d993cddc","resolution":{"observed_at":"2026-08-07T06:02:30.430076Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-08-07T10:20:43.291216Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.15724","last_updated":"2025-06-06T01:51:24Z","snapshot_observed_at":"2026-08-07T10:11:30.411161Z","submitted_at":"2025-06-06T01:51:24Z","title":"MadaKV: Adaptive Modality-Perception KV Cache Eviction for Efficient Multimodal Long-Context Inference","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T10:20:43.291216Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2506.15724"},"observation_digest":"sha256:79163eeffefcfc53e41f8a0106f5184e8e552f0d589b4e714b25a0a164782854","observation_id":"880fad2a-01fb-4971-8831-b07aabac3f1b","resolution":{"observed_at":"2026-08-07T10:20:43.291216Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-08-06T16:12:02.171862Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.15882","last_updated":"2025-08-04T20:48:37Z","snapshot_observed_at":"2026-08-08T04:41:23.907184Z","submitted_at":"2025-07-18T19:33:15Z","title":"Document Haystack: A Long Context Multimodal Image/Document Understanding Vision LLM Benchmark","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T16:12:02.171862Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2507.15882"},"observation_digest":"sha256:3ea1a47f1f080ce0acf3fdbe13eec48be0a19312cf2879bd44676a3a25f32a03","observation_id":"2c6eae63-2f9c-449f-b856-6080ae010be5","resolution":{"observed_at":"2026-08-06T16:12:02.171862Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-08-05T20:02:29.383622Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.15810","last_updated":"2025-08-15T08:41:33Z","snapshot_observed_at":"2026-08-05T20:02:26.951057Z","submitted_at":"2025-08-15T08:41:33Z","title":"Detecting Hope, Hate, and Emotion in Arabic Textual Speech and Multi-modal Memes Using Large Language Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-05T20:02:29.383622Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2508.15810"},"observation_digest":"sha256:dbf511c38bb55785be25b107e2e3a0dcdb3c74e367bd7c8d4e307439b01c86f7","observation_id":"201b3a3c-b695-4954-b796-3d4b35a02981","resolution":{"observed_at":"2026-08-05T20:02:29.383622Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":"2404.18532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-07-04T13:19:50.957441Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":"2bf00545-5a48-438f-9805-41788308bea9","year":2024},"citing_paper":{"arxiv_id":"2604.05887","last_updated":"2026-04-07T13:51:07Z","snapshot_observed_at":"2026-07-06T22:54:30.567988Z","submitted_at":"2026-04-07T13:51:07Z","title":"HybridKV: Hybrid KV Cache Compression for Efficient Multimodal Large Language Model Inference","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T19:58:45.374139Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2604.05887"},"observation_digest":"sha256:eb5445e35be47686757d63997bfd2f9fe01dc0830ef2c44cc60ccd418654044d","observation_id":"3af10797-c208-4084-b5cd-a7c05ee04a06","resolution":{"observed_at":"2026-05-10T22:20:48.365349Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":"2404.18532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-07-04T13:19:50.957441Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":"2bf00545-5a48-438f-9805-41788308bea9","year":2024},"citing_paper":{"arxiv_id":"2604.27389","last_updated":"2026-05-13T08:16:41Z","snapshot_observed_at":"2026-08-06T03:49:04.127433Z","submitted_at":"2026-04-30T03:59:22Z","title":"COHERENCE: Benchmarking Fine-Grained Image-Text Alignment in Interleaved Multimodal Contexts","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-07T09:47:27.891288Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2604.27389"},"observation_digest":"sha256:1113c0d58694e6d9652a59d506b53137ea438115652a39f0cd6ae1a6fa077da6","observation_id":"a0cc7951-5cd0-430c-b556-f64234d04640","resolution":{"observed_at":"2026-05-12T09:41:26.768262Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":"2404.18532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-07-04T13:19:50.957441Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":"2bf00545-5a48-438f-9805-41788308bea9","year":2024},"citing_paper":{"arxiv_id":"2604.27389","last_updated":"2026-05-13T08:16:41Z","snapshot_observed_at":"2026-08-06T03:49:04.127433Z","submitted_at":"2026-04-30T03:59:22Z","title":"COHERENCE: Benchmarking Fine-Grained Image-Text Alignment in Interleaved Multimodal Contexts","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-14T21:15:30.952247Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2604.27389"},"observation_digest":"sha256:3d9edf6178b81a422f80a5fe36d9e34ea477ca3a8809dc85f5675764d7a061ad","observation_id":"a4889efe-5619-4e05-85a1-fb03afec0c7e","resolution":{"observed_at":"2026-05-14T21:17:59.250170Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":"2404.18532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-07-04T13:19:50.957441Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":"2bf00545-5a48-438f-9805-41788308bea9","year":2024},"citing_paper":{"arxiv_id":"2605.02262","last_updated":"2026-05-04T06:17:13Z","snapshot_observed_at":"2026-08-05T23:06:05.311364Z","submitted_at":"2026-05-04T06:17:13Z","title":"WindowQuant: Mixed-Precision KV Cache Quantization based on Window-Level Similarity for VLMs Inference Optimization","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-09T16:06:26.450483Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2605.02262"},"observation_digest":"sha256:a365e2ba3b9aca75ec5719b01b67a9d610fd943b2f6f8b9703b29878643acdce","observation_id":"3bf906e0-6e71-48f3-98aa-5f0ec809dcd8","resolution":{"observed_at":"2026-05-11T16:36:07.948947Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":"2404.18532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-07-04T13:19:50.957441Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":"2bf00545-5a48-438f-9805-41788308bea9","year":2024},"citing_paper":{"arxiv_id":"2605.04075","last_updated":"2026-04-14T08:17:53Z","snapshot_observed_at":"2026-08-03T03:08:39.084770Z","submitted_at":"2026-04-14T08:17:53Z","title":"RetentiveKV: State-Space Memory for Uncertainty-Aware Multimodal KV Cache Eviction","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-05-10T15:29:17.567557Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2605.04075"},"observation_digest":"sha256:86dcb531b00f09f2241050e2b632cacb2e0f5cf13abd821d21f20a83f2c28e07","observation_id":"8bca502f-b286-4aa9-87eb-9325ef1d5719","resolution":{"observed_at":"2026-05-11T10:26:02.126197Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":"2404.18532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-07-04T13:19:50.957441Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":"2bf00545-5a48-438f-9805-41788308bea9","year":2024},"citing_paper":{"arxiv_id":"2605.11591","last_updated":"2026-05-12T06:17:01Z","snapshot_observed_at":"2026-07-06T23:23:26.077921Z","submitted_at":"2026-05-12T06:17:01Z","title":"Logit-Attention Divergence: Mitigating Position Bias in Multi-Image Retrieval via Attention-Guided Calibration","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-13T01:26:45.182597Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2605.11591"},"observation_digest":"sha256:75f88b7cb63859f026ae410a2d514740cec51e204cd535f2579e3477272b80b2","observation_id":"117c5620-2cee-4b79-8911-f7da3513b579","resolution":{"observed_at":"2026-05-13T01:27:01.748004Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":"2404.18532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-07-04T13:19:50.957441Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":"2bf00545-5a48-438f-9805-41788308bea9","year":2024},"citing_paper":{"arxiv_id":"2605.12703","last_updated":"2026-05-12T19:57:37Z","snapshot_observed_at":"2026-08-02T23:28:23.895717Z","submitted_at":"2026-05-12T19:57:37Z","title":"MMCL-Bench: Multimodal Context Learning from Visual Rules, Procedures, and Evidence","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-14T20:55:50.873271Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2605.12703"},"observation_digest":"sha256:2ad39d48a9ed7e379e76aa9f1bbb7228ba3a424d186aa8713507f512cfafcdab","observation_id":"5a435841-c8b9-4dfc-ad51-0309e22c2e8b","resolution":{"observed_at":"2026-05-14T20:59:27.482431Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":"2404.18532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-07-04T13:19:50.957441Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":"2bf00545-5a48-438f-9805-41788308bea9","year":2024},"citing_paper":{"arxiv_id":"2605.14906","last_updated":"2026-05-14T14:41:17Z","snapshot_observed_at":"2026-08-02T09:41:14.453439Z","submitted_at":"2026-05-14T14:41:17Z","title":"MemLens: Benchmarking Multimodal Long-Term Memory in Large Vision-Language Models","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-06-30T21:00:25.664841Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2605.14906"},"observation_digest":"sha256:2e58234251aab1c0d1f6dd6ab6c7d025f332fb2d2c7ce1e285ebf888f7100dc1","observation_id":"739ec28f-6829-47d1-9e38-074a0d7ff60c","resolution":{"observed_at":"2026-06-30T21:05:04.279043Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":"2404.18532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-07-04T13:19:50.957441Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":"2bf00545-5a48-438f-9805-41788308bea9","year":2024},"citing_paper":{"arxiv_id":"2605.15710","last_updated":"2026-05-15T08:00:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-15T08:00:46Z","title":"SMMBench: A Benchmark for Source-Distributed Multimodal Agent Memory","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-20T19:11:15.761831Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2605.15710"},"observation_digest":"sha256:0d567a8ce22571851f16473730c79181051fcf6549a05080a3bf0c4acc482b2f","observation_id":"6f75a050-2611-44a5-b3b8-623a2d38aa37","resolution":{"observed_at":"2026-05-20T19:13:40.544030Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":"2404.18532","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-07-04T13:19:50.957441Z","title":"Milebench: Benchmarking mllms in long context","venue":null,"work_id":"2bf00545-5a48-438f-9805-41788308bea9","year":2024},"citing_paper":{"arxiv_id":"2606.26602","last_updated":"2026-06-25T05:02:38Z","snapshot_observed_at":"2026-08-05T11:23:05.305941Z","submitted_at":"2026-06-25T05:02:38Z","title":"DiCoBench: Benchmarking Multi-Image Fine-Grained Perception via Differential and Commonality Visual Cues","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-26T05:18:01.931929Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2606.26602"},"observation_digest":"sha256:c5f2439bdd16e75c282511918321f9e85ce65f5145f129d031627c1fde0cf956","observation_id":"ad7d0583-4333-40dd-ba45-118a7ff0333c","resolution":{"observed_at":"2026-07-04T13:19:50.958652Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-08-01T02:55:23.417862Z","title":"Richard Yu, Xiang Wan, and Benyou Wang","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.25294","last_updated":"2026-07-28T05:06:43Z","snapshot_observed_at":"2026-08-06T17:57:12.013512Z","submitted_at":"2026-07-28T05:06:43Z","title":"CLBench-V: Evaluating Multimodal Context Learning from Grounding to Knowledge Acquisition","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-01T02:55:23.417862Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2607.25294"},"observation_digest":"sha256:43320f92f0e64f4dd93498bc82685139809c9483b54ec5a65dcbdb1c089e52c6","observation_id":"a8857015-3a1b-4518-9281-64cdbda5985d","resolution":{"observed_at":"2026-08-01T02:55:23.417862Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.18532","snapshot_observed_at":"2026-08-06T00:41:02.947856Z","title":"arXiv preprint arXiv:2404.18532 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00962","last_updated":"2026-08-02T03:20:54Z","snapshot_observed_at":"2026-08-08T02:24:19.882724Z","submitted_at":"2026-08-02T03:20:54Z","title":"PMMC: Prospective Multimodal Memory Compilation for Long-Term LVLM Agents","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-06T00:41:02.947856Z"},"links":{"cited_paper":"/paper/2404.18532","citing_paper":"/paper/2608.00962"},"observation_digest":"sha256:f2a71e08f1313930feac86568ad4b81ac193f98c01557cc1fc47aeaadfb7f199","observation_id":"49814009-ab6d-43c4-b880-db4335dddc74","resolution":{"observed_at":"2026-08-06T00:41:02.947856Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2404.18532/citation-record","integrity":"/paper/2404.18532/integrity","json":"/paper/2404.18532/citation-record.json","paper":"/paper/2404.18532"},"outbound":[],"paper":{"arxiv_id":"2404.18532","last_updated":"2024-05-15T05:43:30Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-04T01:28:52.686106Z","submitted_at":"2024-04-29T09:19:05Z","title":"MileBench: Benchmarking MLLMs in Long Context"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 18 inbound Pith citation observations for arXiv:2404.18532."}