{"as_of":"2026-08-10T04:02:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2d593063012b5c1b90ff1267657a36fd1fb49060e3e2ffe3bbd66e0d29123af5","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":12,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":12,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":12,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":12,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T17:01:02.057343Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-09T17:01:02.057343Z","title":"Bring reason to vision: Understanding perception and reasoning through model merging.arXiv preprint arXiv:2505.05464,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.01015","last_updated":"2026-06-23T02:24:17Z","snapshot_observed_at":"2026-08-09T22:17:27.125284Z","submitted_at":"2025-02-03T03:18:26Z","title":"Task Vector Bases: A Unified and Scalable Framework for Compressed Task Arithmetic","version":5},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-09T17:01:02.057343Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2502.01015"},"observation_digest":"sha256:aeae571491893e85d5b5243d78e7ab779e207a0f4419bd4e6c48d832a6ceb265","observation_id":"00e8cf36-898b-4728-a796-78f4cd00c571","resolution":{"observed_at":"2026-08-09T17:01:02.057343Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":"2505.05464","doi":"10.48550/arxiv.2505.05464","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Bring reason to vision: Understanding perception and reasoning through model merging","venue":"ArXiv.org","work_id":"7e7618b6-c713-4015-85c2-92e2bb34e7bc","year":2025},"citing_paper":{"arxiv_id":"2503.09567","last_updated":"2025-07-18T15:57:54Z","snapshot_observed_at":"2026-08-08T22:33:20.124926Z","submitted_at":"2025-03-12T17:35:03Z","title":"Towards Reasoning Era: A Survey of Long Chain-of-Thought for Reasoning Large Language Models","version":5},"reference_index":99,"source":"pdf_text","source_observed_at":"2026-05-12T08:40:40.910461Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2503.09567"},"observation_digest":"sha256:cd6e25cfa2c8335a4e6baa61a8f9db363a3f0a4b90f49a8e323008f8b711cd65","observation_id":"df9585c7-f145-4952-9ca8-d51d77c80e69","resolution":{"observed_at":"2026-05-12T08:40:41.876272Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-07T15:39:50.323121Z","title":"In European Confer- ence on Computer Vision, pages 370–387","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.14481","last_updated":"2025-05-21T05:04:58Z","snapshot_observed_at":"2026-08-07T15:31:09.140183Z","submitted_at":"2025-05-20T15:14:47Z","title":"PlanGPT-VL: Enhancing Urban Planning with Domain-Specific Vision-Language Models","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-07T15:39:50.323121Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2505.14481"},"observation_digest":"sha256:44d021bfd6da054510c39c59ed1736fae84fbb4ec218facbc8add39ae4bf9a0f","observation_id":"52166a51-173b-4065-baed-1b8849038595","resolution":{"observed_at":"2026-08-07T15:39:50.323121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-07T15:13:55.819405Z","title":"Bring reason to vision: Understanding perception and reasoning through model merging","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.16151","last_updated":"2025-05-22T02:51:12Z","snapshot_observed_at":"2026-08-08T23:52:03.283523Z","submitted_at":"2025-05-22T02:51:12Z","title":"Training-Free Reasoning and Reflection in MLLMs","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T15:13:55.819405Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2505.16151"},"observation_digest":"sha256:185a7450451f0c86443368583848dc7a5c02674356455638ebf9f0813bd3685a","observation_id":"c9568ce4-e1e1-418e-97cd-37470082b171","resolution":{"observed_at":"2026-08-07T15:13:55.819405Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-05T20:44:35.625880Z","title":"Bring reason to vision: Understanding perception and reasoning through model merging","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.09945","last_updated":"2025-08-13T17:00:44Z","snapshot_observed_at":"2026-08-08T00:13:29.644971Z","submitted_at":"2025-08-13T17:00:44Z","title":"VisCodex: Unified Multimodal Code Generation via Merging Vision and Coding Models","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-05T20:44:35.625880Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2508.09945"},"observation_digest":"sha256:ad94c01c433a7c62dc6fbbcc8d2e7bde04df29583365b91c9bcdda3a2dcdd931","observation_id":"18bcdaae-1d8b-42aa-9472-b489a4c8ac21","resolution":{"observed_at":"2026-08-05T20:44:35.625880Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":"2505.05464","doi":"10.48550/arxiv.2505.05464","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Bring reason to vision: Understanding perception and reasoning through model merging","venue":"ArXiv.org","work_id":"7e7618b6-c713-4015-85c2-92e2bb34e7bc","year":2025},"citing_paper":{"arxiv_id":"2604.11399","last_updated":"2026-04-13T12:41:50Z","snapshot_observed_at":"2026-07-06T22:59:49.614780Z","submitted_at":"2026-04-13T12:41:50Z","title":"Reasoning Resides in Layers: Restoring Temporal Reasoning in Video-Language Models with Layer-Selective Merging","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T15:42:43.948462Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2604.11399"},"observation_digest":"sha256:f881bc77b62ce8c417ed2b61cae183161c0dd20f8fafa9ef44e40aae9986dd64","observation_id":"3c15fbc6-41d7-49d8-9fe3-6b4c95398690","resolution":{"observed_at":"2026-05-11T09:56:05.779148Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":"2505.05464","doi":"10.48550/arxiv.2505.05464","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Bring reason to vision: Understanding perception and reasoning through model merging","venue":"ArXiv.org","work_id":"7e7618b6-c713-4015-85c2-92e2bb34e7bc","year":2025},"citing_paper":{"arxiv_id":"2604.12890","last_updated":"2026-04-25T02:32:57Z","snapshot_observed_at":"2026-07-06T23:01:00.830207Z","submitted_at":"2026-04-14T15:40:28Z","title":"Towards Long-horizon Agentic Multimodal Search","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-05-10T15:40:32.137708Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2604.12890"},"observation_digest":"sha256:11a92c5e3413d9f10bc7ca0791a5cc7de16af17d2d9f501df66d7cc968c78355","observation_id":"5e991e75-8394-4b31-8ba6-d69525255154","resolution":{"observed_at":"2026-05-11T10:06:00.190608Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":"2505.05464","doi":"10.48550/arxiv.2505.05464","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Bring reason to vision: Understanding perception and reasoning through model merging","venue":"ArXiv.org","work_id":"7e7618b6-c713-4015-85c2-92e2bb34e7bc","year":2025},"citing_paper":{"arxiv_id":"2605.31073","last_updated":"2026-05-29T09:42:08Z","snapshot_observed_at":"2026-07-06T23:40:17.605034Z","submitted_at":"2026-05-29T09:42:08Z","title":"ConsisGuard: Aligning Safety Deliberation with Policy Enforcement in LLM Guardrails","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-06-28T23:03:01.151745Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2605.31073"},"observation_digest":"sha256:1bf107f521e0b7d96b7dcf7c96c84c0f91c899866c7718c1d2b48af069ba0222","observation_id":"dd4b5312-f3f0-48b6-aeb6-f3645440beb6","resolution":{"observed_at":"2026-06-28T23:12:46.487917Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":"2505.05464","doi":"10.48550/arxiv.2505.05464","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Bring reason to vision: Understanding perception and reasoning through model merging","venue":"ArXiv.org","work_id":"7e7618b6-c713-4015-85c2-92e2bb34e7bc","year":2025},"citing_paper":{"arxiv_id":"2606.01558","last_updated":"2026-06-01T02:02:23Z","snapshot_observed_at":"2026-08-04T16:32:29.186022Z","submitted_at":"2026-06-01T02:02:23Z","title":"Attention-guided Fine-tuning of Multimodal Large Language Models Improves Chain-of-Thought Reasoning","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-06-28T15:45:26.891621Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2606.01558"},"observation_digest":"sha256:e5d1f3e64f061c72dc3a7f1e286fd52432e43264ba58570302dc375a24526a42","observation_id":"f9a56e1a-b4b3-4e7f-852a-9ff13eb396d9","resolution":{"observed_at":"2026-07-01T22:06:16.623754Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":"2505.05464","doi":"10.48550/arxiv.2505.05464","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Bring reason to vision: Understanding perception and reasoning through model merging","venue":"ArXiv.org","work_id":"7e7618b6-c713-4015-85c2-92e2bb34e7bc","year":2025},"citing_paper":{"arxiv_id":"2606.05744","last_updated":"2026-06-04T06:17:11Z","snapshot_observed_at":"2026-08-05T22:02:20.138443Z","submitted_at":"2026-06-04T06:17:11Z","title":"PlanBench-V: A Spatial Planning Map Benchmark for Vision-Language Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-28T02:10:12.974244Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2606.05744"},"observation_digest":"sha256:89140bd57d56fa7289f72357daabd35366bff55c0ca23cc71327c2b948b2ad4a","observation_id":"613f2f3a-19a1-4ec1-a37f-867336c1a40e","resolution":{"observed_at":"2026-07-02T12:26:56.348287Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-06T00:11:24.638371Z","title":"Chollet, F","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.01534","last_updated":"2026-08-02T23:01:10Z","snapshot_observed_at":"2026-08-10T01:16:39.522280Z","submitted_at":"2026-08-02T23:01:10Z","title":"Recursive Vision Language Models for General Symbolic Reasoning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T00:11:24.638371Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2608.01534"},"observation_digest":"sha256:48adc4c90b3f2665602fa9b30cdf1b4a15d31689264c74abc3ecae81a23df62b","observation_id":"e4417778-bbfb-4451-8ae9-3e4b983a85e3","resolution":{"observed_at":"2026-08-06T00:11:24.638371Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.05464","snapshot_observed_at":"2026-08-08T11:15:06.746182Z","title":"arXiv preprint arXiv:2505.05464 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.05541","last_updated":"2026-08-06T02:39:10Z","snapshot_observed_at":"2026-08-09T23:11:02.466018Z","submitted_at":"2026-08-06T02:39:10Z","title":"Hyper-ES: Effective Evolution Strategies for LLM Reasoning via Descent Direction Merging","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-08T11:15:06.746182Z"},"links":{"cited_paper":"/paper/2505.05464","citing_paper":"/paper/2608.05541"},"observation_digest":"sha256:95515a4ceeb45ba70c62fb38161e91e968389ecb25e1994a68be895fa7dc96ce","observation_id":"97c0e90a-481d-4511-b135-8327fc93146a","resolution":{"observed_at":"2026-08-08T11:15:06.746182Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2505.05464/citation-record","integrity":"/paper/2505.05464/integrity","json":"/paper/2505.05464/citation-record.json","paper":"/paper/2505.05464"},"outbound":[],"paper":{"arxiv_id":"2505.05464","last_updated":"2025-07-15T06:09:44Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T15:48:41.620409Z","submitted_at":"2025-05-08T17:56:23Z","title":"Bring Reason to Vision: Understanding Perception and Reasoning through Model Merging"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 12 inbound Pith citation observations for arXiv:2505.05464."}