{"as_of":"2026-08-18T17:55:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:0b492d727f7895dd27994ae8ce3f68ea38f5ca46ebc25a0d8e73914f7c3d8a5e","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":14,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":14,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":14,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":14,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T10:21:54.423993Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-25T04:50:21.359486Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-08-12T19:33:00.706742Z","title":"Mini-monkey: Multi-scale adaptive cropping for multimodal large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10640","last_updated":"2024-11-16T00:14:51Z","snapshot_observed_at":"2026-08-17T21:54:44.695295Z","submitted_at":"2024-11-16T00:14:51Z","title":"BlueLM-V-3B: Algorithm and System Co-Design for Multimodal Large Language Models on Mobile Devices","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-12T19:33:00.706742Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2411.10640"},"observation_digest":"sha256:0a5a72476c77e86b98dfc35ee628494d247e7143b5a50b810154a1b9d3a1cce3","observation_id":"d4560ee1-b806-42fd-be9a-873ccef88b5d","resolution":{"observed_at":"2026-08-12T19:33:00.706742Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-08-11T21:31:44.394774Z","title":"Mini-monkey: Alleviating the semantic saw- tooth effect for lightweight mllms via complementary image pyramid","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.04449","last_updated":"2025-08-06T16:57:39Z","snapshot_observed_at":"2026-08-14T02:35:33.601057Z","submitted_at":"2024-12-05T18:58:03Z","title":"p-MoD: Building Mixture-of-Depths MLLMs via Progressive Ratio Decay","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T21:31:44.394774Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2412.04449"},"observation_digest":"sha256:c6c31af5f007e4989f71bc63ece422fc65017ef1e279ac20cd2205f235fece40","observation_id":"f4b447fe-d351-4dde-a838-546c1f5c83ef","resolution":{"observed_at":"2026-08-11T21:31:44.394774Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-08-09T21:39:21.707003Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.19036","last_updated":"2025-05-30T12:26:44Z","snapshot_observed_at":"2026-08-10T12:01:17.734125Z","submitted_at":"2025-01-31T11:09:16Z","title":"RedundancyLens: Revealing and Exploiting Visual Token Processing Redundancy for Efficient Decoder-Only MLLMs","version":3},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-09T21:39:21.707003Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2501.19036"},"observation_digest":"sha256:570e2b11ba8e08845dc4a1422f9e4fe433b9b24049f4c2f7d90f12711573cf9f","observation_id":"7828f2ba-fc4f-4316-ac5e-3fb15c75e4c4","resolution":{"observed_at":"2026-08-09T21:39:21.707003Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-08-16T10:21:54.423993Z","title":"Mini-monkey: Alleviating the semantic saw- tooth effect for lightweight mllms via complementary image pyramid","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.18397","last_updated":"2025-07-15T07:32:27Z","snapshot_observed_at":"2026-08-16T10:15:07.374373Z","submitted_at":"2025-04-25T14:48:18Z","title":"Unsupervised Visual Chain-of-Thought Reasoning via Preference Optimization","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-16T10:21:54.423993Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2504.18397"},"observation_digest":"sha256:2e7ee54bfec1953dff18143564bde767b3f25c4c355d60985506875cacf5a607","observation_id":"e5b3535b-2a3b-4b74-bbe5-529146704571","resolution":{"observed_at":"2026-08-16T10:21:54.423993Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-08-15T21:02:41.670215Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.11015","last_updated":"2025-05-27T08:00:24Z","snapshot_observed_at":"2026-08-16T22:24:17.930161Z","submitted_at":"2025-05-16T09:09:46Z","title":"WildDoc: How Far Are We from Achieving Comprehensive and Robust Document Understanding in the Wild?","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T21:02:41.670215Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2505.11015"},"observation_digest":"sha256:3d27a287a3a1d24ff4c27bae452cef03c25c05f8947bab3bc50edc583c3b2de1","observation_id":"015b35cf-7bf1-4dfc-947c-46164b3e9309","resolution":{"observed_at":"2026-08-15T21:02:41.670215Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-08-07T04:08:41.334971Z","title":"Mini-monkey: Alleviating the semantic sawtooth effect for lightweight mllms via complementary image pyramid,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.11515","last_updated":"2025-06-13T07:16:41Z","snapshot_observed_at":"2026-08-14T02:40:24.668658Z","submitted_at":"2025-06-13T07:16:41Z","title":"Manager: Aggregating Insights from Unimodal Experts in Two-Tower VLMs and MLLMs","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-07T04:08:41.334971Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2506.11515"},"observation_digest":"sha256:2fbc4d9bc603365defee49333ee83ea20de6f4547f8a1ee40c199678b3b4c38e","observation_id":"7a46b53b-59b1-4e10-a856-f00b73567ab5","resolution":{"observed_at":"2026-08-07T04:08:41.334971Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-08-06T20:53:05.239887Z","title":"Mini-monkey: Alleviating the semantic sawtooth effect for lightweight mllms via complementary image pyramid","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.01643","last_updated":"2025-07-02T12:17:23Z","snapshot_observed_at":"2026-08-17T05:41:16.319289Z","submitted_at":"2025-07-02T12:17:23Z","title":"SAILViT: Towards Robust and Generalizable Visual Backbones for MLLMs via Gradual Feature Refinement","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T20:53:05.239887Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2507.01643"},"observation_digest":"sha256:93dcb3c43a461d8ccd4f10a99486f90a52f2476386ec76732fd514a9a166a812","observation_id":"f611ea19-3a0c-48f2-88f7-795668443c75","resolution":{"observed_at":"2026-08-06T20:53:05.239887Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-08-06T16:43:49.933618Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.12883","last_updated":"2025-08-13T05:27:53Z","snapshot_observed_at":"2026-08-18T13:01:08.864365Z","submitted_at":"2025-07-17T08:09:31Z","title":"HRSeg: High-Resolution Visual Perception and Enhancement for Reasoning Segmentation","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T16:43:49.933618Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2507.12883"},"observation_digest":"sha256:ca7d4c4b0c8c7d78a612e11cda1310b1ea35dee38c9e67e92b9f1a0b83ca6776","observation_id":"2a234bd9-2c76-4500-b588-a30ff55ea878","resolution":{"observed_at":"2026-08-06T16:43:49.933618Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-08-06T15:57:01.142507Z","title":"Mini-monkey: Alleviate the saw- tooth effect by multi-scale adaptive cropping","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.14675","last_updated":"2025-07-19T16:03:34Z","snapshot_observed_at":"2026-08-17T21:12:15.287539Z","submitted_at":"2025-07-19T16:03:34Z","title":"Docopilot: Improving Multimodal Models for Document-Level Understanding","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T15:57:01.142507Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2507.14675"},"observation_digest":"sha256:007e87e900c7eaf8f67ffc9e80801d581f6855211f20a64e8a1b96c5ea8c7154","observation_id":"2b79dd09-ef27-4f58-ae5b-f6dfefd523b9","resolution":{"observed_at":"2026-08-06T15:57:01.142507Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-08-06T05:36:18.365077Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.01540","last_updated":"2025-08-03T01:49:08Z","snapshot_observed_at":"2026-08-15T08:02:02.409123Z","submitted_at":"2025-08-03T01:49:08Z","title":"MagicVL-2B: Empowering Vision-Language Models on Mobile Devices with Lightweight Visual Encoders via Curriculum Learning","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-06T05:36:18.365077Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2508.01540"},"observation_digest":"sha256:0f26d08eadb31cb5e13f702d5f1ad113ba67824f2a2eba63dbed3aee20352b1f","observation_id":"fe8d2891-1f27-40ee-8a9b-e9f96cbd28d8","resolution":{"observed_at":"2026-08-06T05:36:18.365077Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":"2408.02034","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mini- monkey: Alleviating the semantic sawtooth effect for lightweight mllms via complementary image pyramid","venue":null,"work_id":"02c1cb29-327f-4dfa-9eda-4da0177d18c0","year":2024},"citing_paper":{"arxiv_id":"2604.00161","last_updated":"2026-04-21T01:45:08Z","snapshot_observed_at":"2026-07-06T22:51:22.254518Z","submitted_at":"2026-03-31T19:09:55Z","title":"Q-Mask: Query-driven Causal Masks for Text Anchoring in OCR-Oriented Vision-Language Models","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-13T23:30:53.449935Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2604.00161"},"observation_digest":"sha256:282c6da3e948b68150feef51ab29a263e8bfe4597a65a5f83ceef28e74044d39","observation_id":"8639ac15-6145-4cde-b036-f7ba2bc56403","resolution":{"observed_at":"2026-05-13T23:33:26.686268Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":"2408.02034","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mini- monkey: Alleviating the semantic sawtooth effect for lightweight mllms via complementary image pyramid","venue":null,"work_id":"02c1cb29-327f-4dfa-9eda-4da0177d18c0","year":2024},"citing_paper":{"arxiv_id":"2604.06912","last_updated":"2026-04-08T10:12:30Z","snapshot_observed_at":"2026-08-16T09:35:47.926898Z","submitted_at":"2026-04-08T10:12:30Z","title":"Q-Zoom: Query-Aware Adaptive Perception for Efficient Multimodal Large Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T18:46:26.869644Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2604.06912"},"observation_digest":"sha256:07f2b9b845a8cda9faddaa868b96c492ccaa6f0582487642ceb848d7508ae58b","observation_id":"85000c5e-19d0-4438-ade5-10153eaed7b5","resolution":{"observed_at":"2026-05-11T00:00:51.526642Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":"2408.02034","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mini- monkey: Alleviating the semantic sawtooth effect for lightweight mllms via complementary image pyramid","venue":null,"work_id":"02c1cb29-327f-4dfa-9eda-4da0177d18c0","year":2024},"citing_paper":{"arxiv_id":"2604.08896","last_updated":"2026-04-10T02:59:38Z","snapshot_observed_at":"2026-08-18T03:11:40.585335Z","submitted_at":"2026-04-10T02:59:38Z","title":"GeoMMBench and GeoMMAgent: Toward Expert-Level Multimodal Intelligence in Geoscience and Remote Sensing","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-10T17:56:12.097628Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2604.08896"},"observation_digest":"sha256:b3e9bd11e751a262f84491a98c55090741440d57844e05933065deb7558eb2f6","observation_id":"4c8d35b5-1b68-42e7-bc82-7f5d0d3929c5","resolution":{"observed_at":"2026-05-11T05:51:00.475208Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid","version":3},"cited_work":{"arxiv_id":"2408.02034","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2408.02034","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mini- monkey: Alleviating the semantic sawtooth effect for lightweight mllms via complementary image pyramid","venue":null,"work_id":"02c1cb29-327f-4dfa-9eda-4da0177d18c0","year":2024},"citing_paper":{"arxiv_id":"2605.23655","last_updated":"2026-05-22T14:07:44Z","snapshot_observed_at":"2026-08-13T14:22:11.413921Z","submitted_at":"2026-05-22T14:07:44Z","title":"CVSearch: Empowering Multimodal LLMs with Cognitive Visual Search for High-Resolution Image Perception","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-25T04:46:52.071071Z"},"links":{"cited_paper":"/paper/2408.02034","citing_paper":"/paper/2605.23655"},"observation_digest":"sha256:dd8739f04c0931cec1135d17f38e832efeaf48d0cde7825472b7222418aaf9dd","observation_id":"7eb3020c-6987-47bc-a5ca-18fa8b34fcc0","resolution":{"observed_at":"2026-05-25T04:50:21.363120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2408.02034/citation-record","integrity":"/paper/2408.02034/integrity","json":"/paper/2408.02034/citation-record.json","paper":"/paper/2408.02034"},"outbound":[],"paper":{"arxiv_id":"2408.02034","last_updated":"2024-10-28T07:40:49Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-16T13:28:40.388520Z","submitted_at":"2024-08-04T13:55:58Z","title":"Mini-Monkey: Alleviating the Semantic Sawtooth Effect for Lightweight MLLMs via Complementary Image Pyramid"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 14 inbound Pith citation observations for arXiv:2408.02034."}