{"as_of":"2026-08-10T06:47:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6798990de090ceb25e65ef994ed010fb159f1baf82cb240cba1a4c33c338f42a","coverage":[{"denominator":36,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":36,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T10:34:50.476017Z","state":"measured"},{"denominator":40,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":40,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":4,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":4,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-14T16:20:19.874172Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-16T10:17:44.669084Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"cited_work":{"arxiv_id":"2506.04997","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.04997","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Towards storage-efficient visual document retrieval: An empirical study on reducing patch-level embeddings.arXiv preprint arXiv:2506.04997, 2025a","venue":null,"work_id":"d4b6a762-b468-4ab0-b8be-002bf522ef16","year":2025},"citing_paper":{"arxiv_id":"2601.21262","last_updated":"2026-04-16T06:59:30Z","snapshot_observed_at":"2026-07-29T16:09:50.031781Z","submitted_at":"2026-01-29T04:47:27Z","title":"CausalEmbed: Auto-Regressive Multi-Vector Generation in Latent Space for Visual Document Embedding","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-16T10:14:15.589472Z"},"links":{"cited_paper":"/paper/2506.04997","citing_paper":"/paper/2601.21262"},"observation_digest":"sha256:86bf62ec7d95e52439fa1c91ebcc666c6eacc16461f17d5c4612fa44a3057b2b","observation_id":"85970e61-20e3-4abb-9f6c-5d709d506f12","resolution":{"observed_at":"2026-05-16T10:17:44.671695Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"cited_work":{"arxiv_id":"2506.04997","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.04997","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Towards storage-efficient visual document retrieval: An empirical study on reducing patch-level embeddings.arXiv preprint arXiv:2506.04997, 2025a","venue":null,"work_id":"d4b6a762-b468-4ab0-b8be-002bf522ef16","year":2025},"citing_paper":{"arxiv_id":"2604.10167","last_updated":"2026-04-11T11:31:11Z","snapshot_observed_at":"2026-07-06T22:58:51.931478Z","submitted_at":"2026-04-11T11:31:11Z","title":"Visual Late Chunking: An Empirical Study of Contextual Chunking for Efficient Visual Document Retrieval","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-10T15:26:44.498777Z"},"links":{"cited_paper":"/paper/2506.04997","citing_paper":"/paper/2604.10167"},"observation_digest":"sha256:6242824b28bcf0eac45f0e08f6947a42ff9e4af4a3a26e64b33d288ef0789853","observation_id":"5c566db8-93ae-4dd2-b054-73ab1c55c255","resolution":{"observed_at":"2026-05-11T10:31:04.261270Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.04997","snapshot_observed_at":"2026-07-11T16:32:55.757864Z","title":"arXiv preprint arXiv:2506.04997 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04605","last_updated":"2026-07-13T07:13:49Z","snapshot_observed_at":"2026-08-10T01:46:03.006177Z","submitted_at":"2026-07-06T02:19:11Z","title":"Do All Visual Tokens Matter Equally? Object-Evidence Preserving Token Merging for Vision-Language Retrieval","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-07-11T16:32:55.757864Z"},"links":{"cited_paper":"/paper/2506.04997","citing_paper":"/paper/2607.04605"},"observation_digest":"sha256:801d1a8164c7b42ce8715a80c516ff59823327e0edf4182a63884b0e10789b64","observation_id":"3717ba08-d5b5-4052-b99b-43d26e915fa0","resolution":{"observed_at":"2026-07-11T16:32:55.757864Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.04997","snapshot_observed_at":"2026-07-14T16:20:19.874172Z","title":"arXiv preprint arXiv:2506.04997 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04605","last_updated":"2026-07-13T07:13:49Z","snapshot_observed_at":"2026-08-10T01:46:03.006177Z","submitted_at":"2026-07-06T02:19:11Z","title":"Do All Visual Tokens Matter Equally? Object-Evidence Preserving Token Merging for Vision-Language Retrieval","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-07-14T16:20:19.874172Z"},"links":{"cited_paper":"/paper/2506.04997","citing_paper":"/paper/2607.04605"},"observation_digest":"sha256:bc79c400360bbadae9ece8c5bf45a787b96fd057731a0306ac22dc5bbee28220","observation_id":"72881047-42d5-48b4-ae9d-5f032e5f9f77","resolution":{"observed_at":"2026-07-14T16:20:19.874172Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2506.04997/citation-record","integrity":"/paper/2506.04997/integrity","json":"/paper/2506.04997/citation-record.json","paper":"/paper/2506.04997"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2407.07726","last_updated":"2024-10-10T17:28:23Z","snapshot_observed_at":"2026-08-08T07:16:45.596308Z","submitted_at":"2024-07-10T14:57:46Z","title":"PaliGemma: A versatile 3B VLM for transfer","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.07726","snapshot_observed_at":"2026-08-07T10:34:47.967212Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:47.967212Z"},"links":{"cited_paper":"/paper/2407.07726","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:bae2e83b94c6ce1df3e12d9c40de316bfac595308778fb1d02d1117f93b823c9","observation_id":"c4335d3a-6921-4dc5-a871-086db87061c9","resolution":{"observed_at":"2026-08-07T10:34:47.967212Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:52.639548Z","title":null,"venue":null,"work_id":"b8815268-3893-432d-afaf-032edcfc2ac9","year":2023},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.022821Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:9f336d9633d07a5101d0604d34a6d57e50607f81f1b59239c9eb03688128416b","observation_id":"514a11b6-ad70-408e-b90d-c9a9cc9db918","resolution":{"observed_at":"2026-08-07T10:34:52.679562Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:52.513183Z","title":"Rossi, Changyou Chen, and Tong Sun","venue":null,"work_id":"077f4943-598c-49b6-a75f-a5d807d0032b","year":2025},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.129493Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:7e77c8ab4d971d2832a93d512fa34d6d59e9a2e0d3c1fa88689733ef62531c11","observation_id":"973dc096-a943-40f6-ac5c-d5d9728186d4","resolution":{"observed_at":"2026-08-07T10:34:52.552822Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:52.403437Z","title":null,"venue":null,"work_id":"7b91589e-369f-4327-b239-77a293919f21","year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.214199Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:95084892cc0910f9d3a0ef2b1be671479a8c52157338809572ac6eb99b8cef7d","observation_id":"a987dba7-5de2-4296-aecc-049d4b327ee3","resolution":{"observed_at":"2026-08-07T10:34:52.469558Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.04952","last_updated":"2024-11-07T18:29:38Z","snapshot_observed_at":"2026-07-06T19:46:54.707852Z","submitted_at":"2024-11-07T18:29:38Z","title":"M3DocRAG: Multi-modal Retrieval is What You Need for Multi-page Multi-document Understanding","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.04952","snapshot_observed_at":"2026-08-07T10:34:48.298867Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.298867Z"},"links":{"cited_paper":"/paper/2411.04952","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:d13f5b0e8de8cf04b1bf3aa6cdf1ff3b38702c86ab7fe5e13af3ffff5dfff6bb","observation_id":"b5544092-a869-4a6b-9b09-20613d430e27","resolution":{"observed_at":"2026-08-07T10:34:48.298867Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.14683","last_updated":"2024-09-23T03:12:43Z","snapshot_observed_at":"2026-08-07T06:44:36.523865Z","submitted_at":"2024-09-23T03:12:43Z","title":"Reducing the Footprint of Multi-Vector Retrieval with Minimal Performance Impact via Token Pooling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.14683","snapshot_observed_at":"2026-08-07T10:34:48.341670Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.341670Z"},"links":{"cited_paper":"/paper/2409.14683","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:8cf2c3414004fb83e6760e423b71128bd6f5227dc33f4a7e43ab50654bba6968","observation_id":"3ff8675f-6b21-4411-b323-6c5dd1997472","resolution":{"observed_at":"2026-08-07T10:34:48.341670Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:52.252530Z","title":null,"venue":null,"work_id":"ed2f7d57-7af3-4bf7-916c-5aaedeb405f6","year":2025},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.398865Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:a66075b140e2fa4d96bd1dadbbba642e3e8f032b4f0792d5f55516d024fafe18","observation_id":"4a81aebb-10e6-4c09-b7c2-81ddb657f1d5","resolution":{"observed_at":"2026-08-07T10:34:52.332510Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:48.469835Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.469835Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:9e6d469cbee19a042bb202253b97f95885ee20755d8f1acb4e66510586818538","observation_id":"4d6c3dcb-c2cb-46b4-9d59-6eb677129d3f","resolution":{"observed_at":"2026-08-07T10:34:48.469835Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:48.514132Z","title":null,"venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.514132Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:6416ddd5ec297ab4314ca6df86b56124fcc87c3fb55f0eff6e1c359c0ce0a497","observation_id":"942c9fa1-6462-49b6-8772-85734335801f","resolution":{"observed_at":"2026-08-07T10:34:48.514132Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:48.591704Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.591704Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:5a0024439ca9ca91a69880a9448fd6d6dc0b235df145b640ba4c2319ffa32921","observation_id":"e6ba5de4-fad5-48ef-a921-e17af4e79d6a","resolution":{"observed_at":"2026-08-07T10:34:48.591704Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:48.660446Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.660446Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:92b90ff7902464b62a685a321dcb1c5d95f8f6cb4c529b5dd6e60cf8ec7f21ad","observation_id":"140cea8c-3e1b-40ec-9232-7fde73b63535","resolution":{"observed_at":"2026-08-07T10:34:48.660446Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:52.093963Z","title":null,"venue":null,"work_id":"894aefb7-b8d7-4f8d-b6d6-e91726966eed","year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.716111Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:001d00335a8504670ef79ce37592c81c16176d798aec09b7c1b3c9c87427a8af","observation_id":"b5b9740e-7b10-4378-9ce6-af4d37c1780b","resolution":{"observed_at":"2026-08-07T10:34:52.178094Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:48.819379Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.819379Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:e0ea231e03f867f10de072106d444de40953fdd5282b70c2a93cf55e480e9db0","observation_id":"fbf8466b-7f6b-45c8-b7c4-46c31086afc5","resolution":{"observed_at":"2026-08-07T10:34:48.819379Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.02392","last_updated":"2024-08-28T08:49:57Z","snapshot_observed_at":"2026-07-06T18:40:22.951929Z","submitted_at":"2024-07-02T16:10:55Z","title":"TokenPacker: Efficient Visual Projector for Multimodal LLM","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.02392","snapshot_observed_at":"2026-08-07T10:34:48.866029Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.866029Z"},"links":{"cited_paper":"/paper/2407.02392","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:e3a3b1c0d9f7511287a5c6b2d3c9e440e933883646b41fd28823a413dff83b11","observation_id":"f1413e60-66b2-4474-8cc5-4e90d4930ee2","resolution":{"observed_at":"2026-08-07T10:34:48.866029Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2202.07800","last_updated":"2022-04-13T23:46:08Z","snapshot_observed_at":"2026-08-03T03:15:29.208149Z","submitted_at":"2022-02-16T00:19:42Z","title":"Not All Patches are What You Need: Expediting Vision Transformers via Token Reorganizations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2202.07800","snapshot_observed_at":"2026-08-07T10:34:48.952540Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:48.952540Z"},"links":{"cited_paper":"/paper/2202.07800","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:e346fddc43f11319955907012898333f426d5c85696a9075e971166ca53213a9","observation_id":"fd4a96a3-e744-4317-88b6-1a3852736fd3","resolution":{"observed_at":"2026-08-07T10:34:48.952540Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:49.030889Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.030889Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:065403e97915fe50c721934996665d38968dc739e3d95db634a2834bf5df39dc","observation_id":"ba6c668c-259f-4a66-a4ca-b573e76dfbd1","resolution":{"observed_at":"2026-08-07T10:34:49.030889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:51.903833Z","title":null,"venue":null,"work_id":"0f797af6-5b46-401d-9004-9a482f1e350d","year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.108336Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:73d131470b55fad0858bc97ee392042e1ce0ebeef3bd75dfb8f1421d2ea68772","observation_id":"a8c716eb-8f34-4fbe-b998-d63c1d26ea2b","resolution":{"observed_at":"2026-08-07T10:34:52.005884Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2203.10244","last_updated":"2022-03-19T05:00:30Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-03-19T05:00:30Z","title":"ChartQA: A Benchmark for Question Answering about Charts with Visual and Logical Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.10244","snapshot_observed_at":"2026-08-07T10:34:49.164741Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.164741Z"},"links":{"cited_paper":"/paper/2203.10244","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:595fed0f8c2896fe2b279fd8008b31c6e61f28948a37b3dce70d5cdfbae18bc3","observation_id":"1b0633fa-1f7a-47fe-b24d-34f1789bbd26","resolution":{"observed_at":"2026-08-07T10:34:49.164741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:51.745639Z","title":null,"venue":null,"work_id":"017f2ecb-8834-4ead-8bb2-927d0aac5ca3","year":2021},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.231921Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:35f2095305e9b031fe4036405ad04634b6c808a6b566ca1ff7efe3722324660a","observation_id":"b2d25baf-5494-41e2-b3a6-6f5ee645b6f7","resolution":{"observed_at":"2026-08-07T10:34:51.837276Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:51.528444Z","title":"Manmatha, and C","venue":null,"work_id":"4d6fcbd1-c2e5-42fc-b574-e9d2f677bd61","year":2020},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.301161Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:c60d1873c3e09b01e1b4e4c7c5395adb068ea518c7e4997394c648453ab8bd50","observation_id":"58d024bb-4cf7-4fd8-9726-ccae9e18de51","resolution":{"observed_at":"2026-08-07T10:34:51.631527Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:51.302294Z","title":null,"venue":null,"work_id":"624a8c9c-7c07-4719-a616-067ed81aaa2a","year":2012},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.396840Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:211c07a73c363c374b8eac9d0b74ac1ce5d6ef30ed9659bd4ef48e29335a0dc0","observation_id":"c4c44aed-6f8f-40c7-995c-540a58f88b35","resolution":{"observed_at":"2026-08-07T10:34:51.432829Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:49.478388Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.478388Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:7264b90aef53d36a29398b06e174ec1b739c11565c9e9cb0f77e49bece577623","observation_id":"6f29794e-0b4a-4dd2-8f1f-36d6e1d43ecc","resolution":{"observed_at":"2026-08-07T10:34:49.478388Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:49.536671Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.536671Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:fbdb6dd280d97a084feb4e9eb2d354210f6c4b1e043c3d3d0f60fcf6eea04f90","observation_id":"90de8716-48c5-4eb3-8f7b-035fa594ad38","resolution":{"observed_at":"2026-08-07T10:34:49.536671Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:51.118961Z","title":null,"venue":null,"work_id":"febdd1ba-ebff-43b9-8756-2319e8245dc7","year":2023},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.610618Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:5edd958f870fea1d07c32c89501d174b1725067fc7a9a2fbbf71f810aba2c910","observation_id":"36cfd537-d3b7-4f2a-9163-e58487aba55a","resolution":{"observed_at":"2026-08-07T10:34:51.235210Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12191","last_updated":"2024-10-03T15:54:49Z","snapshot_observed_at":"2026-08-06T05:35:29.109022Z","submitted_at":"2024-09-18T17:59:32Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12191","snapshot_observed_at":"2026-08-07T10:34:49.665368Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.665368Z"},"links":{"cited_paper":"/paper/2409.12191","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:d723d5eefcbe5c214747389a2167bcb351097d22a1e0ab2ebb80294248b48ae9","observation_id":"27f5e534-644c-4f2a-875b-05c548853e35","resolution":{"observed_at":"2026-08-07T10:34:49.665368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"status/1826238","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:50.829693Z","title":null,"venue":null,"work_id":"57c99a00-1716-4566-9458-0daca8b6fa88","year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.726150Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:b3618e8a5cdbcfa4351a6d91bd2498f9cb1f3bc8ba22e27f3be40870cae6a391","observation_id":"a0814261-a18d-4dc3-9d68-ead305cccf77","resolution":{"observed_at":"2026-08-07T10:34:50.883764Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02686","last_updated":"2025-06-12T16:39:09Z","snapshot_observed_at":"2026-08-07T15:56:21.414821Z","submitted_at":"2025-05-05T14:33:49Z","title":"Sailing by the Stars: A Survey on Reward Models and Learning Strategies for Learning from Rewards","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02686","snapshot_observed_at":"2026-08-07T10:34:49.797401Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.797401Z"},"links":{"cited_paper":"/paper/2505.02686","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:001aa491e36059b462ece111fa01d19cfe92967c0b3984269ebcb6ee783a45dd","observation_id":"ba0db544-a3dc-49a4-99ab-304f93204ed0","resolution":{"observed_at":"2026-08-07T10:34:49.797401Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:49.838297Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.838297Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:6372c4ed04209b42bca6ab739be032523116a639a972bec39e05695a779a9097","observation_id":"4eb93546-8eaa-423c-a5cd-b4e8df3ee813","resolution":{"observed_at":"2026-08-07T10:34:49.838297Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.13670","last_updated":"2025-05-29T03:19:42Z","snapshot_observed_at":"2026-07-06T20:09:05.798010Z","submitted_at":"2024-12-18T09:53:12Z","title":"AntiLeakBench: Preventing Data Contamination by Automatically Constructing Benchmarks with Updated Real-World Knowledge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.13670","snapshot_observed_at":"2026-08-07T10:34:49.916745Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.916745Z"},"links":{"cited_paper":"/paper/2412.13670","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:73c2165b963068b0a3b354b0b11ed2fc85ac6e80872f97e556b242f06e89867f","observation_id":"914ecf84-8c46-4d30-b1a2-819367705d34","resolution":{"observed_at":"2026-08-07T10:34:49.916745Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.17247","last_updated":"2025-02-27T11:16:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-22T17:59:53Z","title":"PyramidDrop: Accelerating Your Large Vision-Language Models via Pyramid Visual Redundancy Reduction","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.17247","snapshot_observed_at":"2026-08-07T10:34:49.994840Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:49.994840Z"},"links":{"cited_paper":"/paper/2410.17247","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:b1c7af486eac70bf431f242858355db524b5c2bed87f8f9024976f1d130c6d34","observation_id":"629ba998-3c0a-4abc-9d38-c6eb14cba1c7","resolution":{"observed_at":"2026-08-07T10:34:49.994840Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:50.085393Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.085393Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:3e43f9f9e8c590ad88f41c64943ed3fffe468c1fcb85655508aa5f4d0fd13b3d","observation_id":"905b0ba9-1dd5-4e70-b9d2-c604e5930e05","resolution":{"observed_at":"2026-08-07T10:34:50.085393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.10594","last_updated":"2025-03-02T01:19:51Z","snapshot_observed_at":"2026-08-10T05:52:50.722724Z","submitted_at":"2024-10-14T15:04:18Z","title":"VisRAG: Vision-based Retrieval-augmented Generation on Multi-modality Documents","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.10594","snapshot_observed_at":"2026-08-07T10:34:50.174984Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.174984Z"},"links":{"cited_paper":"/paper/2410.10594","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:b5662204e4ace6600dbf47499fa31bd6d197fc5de443f91f4ffa10c6907bed8b","observation_id":"c44d13d5-7738-4f9a-86ef-9a5e82d9b65f","resolution":{"observed_at":"2026-08-07T10:34:50.174984Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.04417","last_updated":"2025-06-03T04:12:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-06T09:18:04Z","title":"SparseVLM: Visual Token Sparsification for Efficient Vision-Language Model Inference","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.04417","snapshot_observed_at":"2026-08-07T10:34:50.294599Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.294599Z"},"links":{"cited_paper":"/paper/2410.04417","citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:94fb4b47403b2fbb7aa960e1eb1203b7072262cf4c66c5fd528c7bb20c0295eb","observation_id":"89a7c29f-c639-4e22-ae57-216d82d7ed2d","resolution":{"observed_at":"2026-08-07T10:34:50.294599Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:50.374302Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.374302Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:f18bf9c22a71b6b580f1b01db74867c76bb479af6d65e12a50bceb61bad9f2c1","observation_id":"ce8dcb9c-3b8f-4db1-a1cc-8ffda0837af3","resolution":{"observed_at":"2026-08-07T10:34:50.374302Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:50.415916Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.415916Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:91dc258aadc6b2033f123127b6b0bb39795c32d59edf99b686295a327af7a570","observation_id":"9cf2edb2-42f3-42e0-bfad-e99caa105cfc","resolution":{"observed_at":"2026-08-07T10:34:50.415916Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:34:50.476017Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T10:34:50.476017Z"},"links":{"citing_paper":"/paper/2506.04997"},"observation_digest":"sha256:a4ef242469a76094d1f3739127697d9811ca619f800037757c7ed44231e49728","observation_id":"5f679184-4f66-4d4c-9347-319b15eb8377","resolution":{"observed_at":"2026-08-07T10:34:50.476017Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.04997","last_updated":"2025-06-05T13:06:01Z","latest_version":1,"primary_category":"cs.IR","snapshot_observed_at":"2026-08-10T02:31:56.018212Z","submitted_at":"2025-06-05T13:06:01Z","title":"Towards Storage-Efficient Visual Document Retrieval: An Empirical Study on Reducing Patch-Level Embeddings"},"reference_resolution":{"displayed":36,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":33,"verified_exact":1,"verified_fuzzy":2},"total_outbound_references":36},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 36 of 36 outbound references and 4 inbound Pith citation observations for arXiv:2506.04997."}