{"as_of":"2026-08-13T09:07:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3d916f0fd595bc77811533924da2bc3851ccc0fe0b6869018baf6e53b12d1529","coverage":[{"denominator":45,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":45,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:09:24.020605Z","state":"measured"},{"denominator":46,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":46,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-27T19:36:57.231932Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T21:37:25.289403Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"cited_work":{"arxiv_id":"2506.08691","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.08691","snapshot_observed_at":"2026-07-02T21:37:25.289403Z","title":null,"venue":null,"work_id":"859f9f97-29f9-4d62-ae4a-f69ebf5dcf02","year":2025},"citing_paper":{"arxiv_id":"2606.08231","last_updated":"2026-06-06T15:39:29Z","snapshot_observed_at":"2026-08-07T16:58:05.550639Z","submitted_at":"2026-06-06T15:39:29Z","title":"Test-Time Scaling in Multimodal Foundation Models: A Comprehensive Survey of Generation and Reasoning","version":1},"reference_index":109,"source":"arxiv_source","source_observed_at":"2026-06-27T19:36:57.231932Z"},"links":{"cited_paper":"/paper/2506.08691","citing_paper":"/paper/2606.08231"},"observation_digest":"sha256:314a4a79a7d1de23b79615aa59f0a005c47f3638deda920b22d4cf9d2350af91","observation_id":"cc0b75d9-5201-44ca-bc33-ab92b9cfcfb7","resolution":{"observed_at":"2026-07-02T21:37:25.290672Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.08691/citation-record","integrity":"/paper/2506.08691/integrity","json":"/paper/2506.08691/citation-record.json","paper":"/paper/2506.08691"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2405.16473","last_updated":"2024-05-26T07:56:30Z","snapshot_observed_at":"2026-08-12T23:57:52.483265Z","submitted_at":"2024-05-26T07:56:30Z","title":"M$^3$CoT: A Novel Benchmark for Multi-Domain Multi-step Multi-modal Chain-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16473","snapshot_observed_at":"2026-08-07T05:09:18.196208Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.196208Z"},"links":{"cited_paper":"/paper/2405.16473","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:63f550052b3ce37f1dd0aecd9f501954b028b84990bc7526be73a61ca49af5c2","observation_id":"23ace000-230b-4d70-a2f6-a22c80d74a04","resolution":{"observed_at":"2026-08-07T05:09:18.196208Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.00855","last_updated":"2024-10-30T14:45:00Z","snapshot_observed_at":"2026-08-13T02:35:12.376027Z","submitted_at":"2024-10-30T14:45:00Z","title":"Vision-Language Models Can Self-Improve Reasoning via Reflection","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.00855","snapshot_observed_at":"2026-08-07T05:09:18.339751Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.339751Z"},"links":{"cited_paper":"/paper/2411.00855","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:17b57f0bc56b826a8dfb63ca5b6f1341cda3a82204253c7beece4325d5e2c23e","observation_id":"7c80bb66-02a8-4420-a3c7-08d9ebd843db","resolution":{"observed_at":"2026-08-07T05:09:18.339751Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:26.778172Z","title":null,"venue":null,"work_id":"475fa8f2-4ba3-4957-b36c-5070a55b6265","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.463000Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:7774bbb46872979e69452774e61b87486330626d3c4040ae3a7dd6042ec24aba","observation_id":"a88b7e8b-8b0b-4109-b718-e29bf56f8044","resolution":{"observed_at":"2026-08-07T05:09:26.869391Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.17179","last_updated":"2024-02-09T00:13:46Z","snapshot_observed_at":"2026-08-11T02:32:39.494378Z","submitted_at":"2023-09-29T12:20:19Z","title":"Alphazero-like Tree-Search can Guide Large Language Model Decoding and Training","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.17179","snapshot_observed_at":"2026-08-07T05:09:18.616592Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.616592Z"},"links":{"cited_paper":"/paper/2309.17179","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:2f1559b321f9fc3cb3e32ab52a41cfa9c62912cd2f225bf13a1f57fab946b6af","observation_id":"4234ff1d-a7b8-47fb-b9ec-4092ced0f56f","resolution":{"observed_at":"2026-08-07T05:09:18.616592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:26.590332Z","title":null,"venue":null,"work_id":"fe5c0430-4980-4956-8043-fb32655c60aa","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.772129Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:ead4489097bca6eb2af7f21e449ac29f49e2390ad07907f0048ee782d920578f","observation_id":"1ff5ce33-eb00-477a-9bfe-dedaac308c02","resolution":{"observed_at":"2026-08-07T05:09:26.688387Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.05237","last_updated":"2025-06-04T10:07:57Z","snapshot_observed_at":"2026-08-11T20:47:50.864749Z","submitted_at":"2024-12-06T18:14:24Z","title":"MAmmoTH-VL: Eliciting Multimodal Reasoning with Instruction Tuning at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.05237","snapshot_observed_at":"2026-08-07T05:09:18.908261Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.908261Z"},"links":{"cited_paper":"/paper/2412.05237","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:cfb333c8313235d3dc7c1c7bf43b744aef3d68cd6828686b04d3e9603f4c696b","observation_id":"545bdbd5-21ae-42ef-983b-af63ca588c00","resolution":{"observed_at":"2026-08-07T05:09:18.908261Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14992","last_updated":"2023-10-23T07:24:28Z","snapshot_observed_at":"2026-07-06T15:32:25.931739Z","submitted_at":"2023-05-24T10:28:28Z","title":"Reasoning with Language Model is Planning with World Model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14992","snapshot_observed_at":"2026-08-07T05:09:18.989445Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:18.989445Z"},"links":{"cited_paper":"/paper/2305.14992","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:24d05362b7749c454ee0fe782081dc12bcd069a02c8fe61f4fbb64a809eba55e","observation_id":"6b82c278-4165-4991-bafd-c1790c45ff77","resolution":{"observed_at":"2026-08-07T05:09:18.989445Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:26.391818Z","title":null,"venue":null,"work_id":"75d4bcd4-ac69-4f20-8d8e-62667db0d859","year":2025},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:19.196937Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:010302e81750f9ce140ae07a1875fba7e5387a656f43850c85f9bfed9c47ec3c","observation_id":"d9b9faf8-5635-468d-ad9f-915166704426","resolution":{"observed_at":"2026-08-07T05:09:26.473220Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:19.387591Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:19.387591Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:5cb16883cf33a4fd73b2b0d08e3ddb4b3bc13f2d3a1a6e87fc7a7846a87baa0e","observation_id":"22e7a611-396f-4de5-9f91-6664fcbf3176","resolution":{"observed_at":"2026-08-07T05:09:19.387591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.11694","last_updated":"2024-12-31T01:38:12Z","snapshot_observed_at":"2026-08-12T18:12:00.655676Z","submitted_at":"2024-11-18T16:15:17Z","title":"Enhancing LLM Reasoning with Reward-guided Tree Search","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.11694","snapshot_observed_at":"2026-08-07T05:09:19.582648Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:19.582648Z"},"links":{"cited_paper":"/paper/2411.11694","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:6692282f97fca7a839fe009c1db8231e859e4ef2c6b699a7b8174e2daa6d8ac0","observation_id":"a4cc418c-f9fe-4f52-9f47-9608d2518715","resolution":{"observed_at":"2026-08-07T05:09:19.582648Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:19.753354Z","title":null,"venue":null,"work_id":null,"year":2006},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:19.753354Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:5f728eb8b1988d3c11e3108b24c8c0f793eafe8dd57ae2cf28e80677476c74e0","observation_id":"cf9a2bda-4e86-44c3-87b2-93b8c94e6932","resolution":{"observed_at":"2026-08-07T05:09:19.753354Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:19.869927Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:19.869927Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:2cd3099d9737fc9623d7a996fa32d354184f06eeb105fbcd69d3e91dba45c3ab","observation_id":"e48b2f53-69e0-4b15-b07c-c0d92a0647c5","resolution":{"observed_at":"2026-08-07T05:09:19.869927Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:26.204510Z","title":null,"venue":null,"work_id":"6d692479-a67d-47f1-af9f-b085dd052fc9","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.016775Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:602600d3f1912f07c14940bc4baa75529d94e91637f6f18b62664a5704b91b33","observation_id":"0838fab1-c0a1-4f3b-bca0-ad1570aa6a54","resolution":{"observed_at":"2026-08-07T05:09:26.276145Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.20050","last_updated":"2023-05-31T17:24:00Z","snapshot_observed_at":"2026-08-11T17:22:43.545531Z","submitted_at":"2023-05-31T17:24:00Z","title":"Let's Verify Step by Step","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.20050","snapshot_observed_at":"2026-08-07T05:09:20.155515Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.155515Z"},"links":{"cited_paper":"/paper/2305.20050","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:d7d81d3ba58fbc73ff5105a58de022caa94e29c8be210fa62827a225b66cf6f7","observation_id":"703522a3-f8db-4ca3-bffe-b10263c8ba16","resolution":{"observed_at":"2026-08-07T05:09:20.155515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.11236","last_updated":"2024-04-25T03:04:14Z","snapshot_observed_at":"2026-08-13T00:52:30.037289Z","submitted_at":"2024-03-17T14:49:09Z","title":"ChartThinker: A Contextual Chain-of-Thought Approach to Optimized Chart Summarization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.11236","snapshot_observed_at":"2026-08-07T05:09:20.285362Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.285362Z"},"links":{"cited_paper":"/paper/2403.11236","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:b8be4bf845a8fa59ed898e4e7098321a5f875c94be8cd3b933b5e961508c5548","observation_id":"53f62501-e248-4ba2-b045-4974a4ac66bb","resolution":{"observed_at":"2026-08-07T05:09:20.285362Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.08291","last_updated":"2023-05-15T01:18:23Z","snapshot_observed_at":"2026-07-06T15:27:02.376191Z","submitted_at":"2023-05-15T01:18:23Z","title":"Large Language Model Guided Tree-of-Thought","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.08291","snapshot_observed_at":"2026-08-07T05:09:20.396124Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.396124Z"},"links":{"cited_paper":"/paper/2305.08291","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:964862b30a66d3824009f191072f8fd3213f66432c3713eb511eff1c31041ab9","observation_id":"3053ebc7-247d-47bd-a3ba-0dca31b5db42","resolution":{"observed_at":"2026-08-07T05:09:20.396124Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.02255","last_updated":"2024-01-21T03:47:06Z","snapshot_observed_at":"2026-07-06T16:27:15.027202Z","submitted_at":"2023-10-03T17:57:24Z","title":"MathVista: Evaluating Mathematical Reasoning of Foundation Models in Visual Contexts","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.02255","snapshot_observed_at":"2026-08-07T05:09:20.523908Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.523908Z"},"links":{"cited_paper":"/paper/2310.02255","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:09b8e07c23cec38bc3d8534926bc665c577a085c521f9c6514f41c4ac9877f11","observation_id":"2343b7c5-30cf-4720-be51-62f1247f0507","resolution":{"observed_at":"2026-08-07T05:09:20.523908Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:20.636332Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.636332Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:823a9526a6efdd6820e4591280e00daa92088b7f7909d81a639a7474dfd84dff","observation_id":"cf928d51-7b97-4148-b96b-7546f86bebab","resolution":{"observed_at":"2026-08-07T05:09:20.636332Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:26.023492Z","title":null,"venue":null,"work_id":"26ff45d1-60f4-4ccf-8af5-9a3244b660c4","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.747025Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:74bd249d9febc8acd132156c6b8e531a6939f20e05b1b33ae0c24749c9415129","observation_id":"02824006-03e0-44bd-8de7-a77bf2d4f914","resolution":{"observed_at":"2026-08-07T05:09:26.094013Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:25.801322Z","title":null,"venue":null,"work_id":"d8d66345-251c-4550-a1df-20c87aa20c0e","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:20.887191Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:45671f29537f06d2a266cbc4b315fe2bca13b220b2d419c0f89bd6629c541f8b","observation_id":"64f3840c-b3e4-4392-9489-7bf1e8cf09d0","resolution":{"observed_at":"2026-08-07T05:09:25.899175Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:25.608973Z","title":null,"venue":null,"work_id":"32a2e83b-5162-432a-bc85-eafaf3b93858","year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.057769Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:ec1aa77220fd53ef857142432df485f219b614a4c672867bf165a0c366214613","observation_id":"243c15b1-9f15-4dcf-aac3-c2fadbf53751","resolution":{"observed_at":"2026-08-07T05:09:25.683150Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14804","last_updated":"2024-02-22T18:56:38Z","snapshot_observed_at":"2026-08-13T01:56:18.092262Z","submitted_at":"2024-02-22T18:56:38Z","title":"Measuring Multimodal Mathematical Reasoning with MATH-Vision Dataset","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14804","snapshot_observed_at":"2026-08-07T05:09:21.174185Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.174185Z"},"links":{"cited_paper":"/paper/2402.14804","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:c905312506e55ee40046e8d5929f0fb14a772957024b106f192ce0c0487a65cd","observation_id":"1a6999fa-faf2-47bb-a125-bdb61b8312a0","resolution":{"observed_at":"2026-08-07T05:09:21.174185Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12191","last_updated":"2024-10-03T15:54:49Z","snapshot_observed_at":"2026-08-06T05:35:29.109022Z","submitted_at":"2024-09-18T17:59:32Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12191","snapshot_observed_at":"2026-08-07T05:09:21.380191Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.380191Z"},"links":{"cited_paper":"/paper/2409.12191","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:fb73c5b2501701def24d2cb32bb3c49a3887122357b050fa6c7de25f6088604f","observation_id":"51686c0b-61a3-46c9-975f-c39cbefbb205","resolution":{"observed_at":"2026-08-07T05:09:21.380191Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2203.11171","last_updated":"2023-03-07T17:57:37Z","snapshot_observed_at":"2026-07-06T12:50:22.773056Z","submitted_at":"2022-03-21T17:48:52Z","title":"Self-Consistency Improves Chain of Thought Reasoning in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.11171","snapshot_observed_at":"2026-08-07T05:09:21.545744Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.545744Z"},"links":{"cited_paper":"/paper/2203.11171","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:16d84decef63659aed243c4e6902f5be982fb759df44589bed32204d696b0241","observation_id":"d22d9979-53ec-47a4-af77-8cb356a00a67","resolution":{"observed_at":"2026-08-07T05:09:21.545744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:21.647142Z","title":"Chi, Sharan Narang, Aakanksha Chowdhery, and Denny Zhou","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.647142Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:5b4585edc41f2c60ba433b150e5842c30b2e32e14cbfd007e8414a18a0b6c839","observation_id":"59fbe118-6147-4a95-9996-0156d9c7c597","resolution":{"observed_at":"2026-08-07T05:09:21.647142Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.18521","last_updated":"2024-06-26T17:50:11Z","snapshot_observed_at":"2026-08-12T23:34:16.774520Z","submitted_at":"2024-06-26T17:50:11Z","title":"CharXiv: Charting Gaps in Realistic Chart Understanding in Multimodal LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.18521","snapshot_observed_at":"2026-08-07T05:09:21.813886Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:21.813886Z"},"links":{"cited_paper":"/paper/2406.18521","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:a0e90506a8e34a72242e011b1fa1d0eeed1f2565efe2a2e9b44f51c0a088e415","observation_id":"bc132b93-c994-4b01-8c8a-2b858c3fdb7f","resolution":{"observed_at":"2026-08-07T05:09:21.813886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:22.026151Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.026151Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:63d75058f9cdf4841975831578b04a2cead1c75843a8704316064e58c095bda5","observation_id":"ecb2c3a1-578c-4746-853a-24f3db9e00e9","resolution":{"observed_at":"2026-08-07T05:09:22.026151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.07001","last_updated":"2024-11-06T13:56:28Z","snapshot_observed_at":"2026-08-13T00:09:46.028743Z","submitted_at":"2024-05-11T12:33:46Z","title":"ChartInsights: Evaluating Multimodal Large Language Models for Low-Level Chart Question Answering","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.07001","snapshot_observed_at":"2026-08-07T05:09:22.134598Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.134598Z"},"links":{"cited_paper":"/paper/2405.07001","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:870d01c5e5eac14d8356c5eb347cdb291856a64f97570441718fb8c4b0608328","observation_id":"a92ce3da-0fea-4661-b610-a2ca319d9984","resolution":{"observed_at":"2026-08-07T05:09:22.134598Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.10332","last_updated":"2025-03-21T12:40:26Z","snapshot_observed_at":"2026-08-12T21:35:41.232161Z","submitted_at":"2024-11-15T16:32:34Z","title":"Number it: Temporal Grounding Videos like Flipping Manga","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.10332","snapshot_observed_at":"2026-08-07T05:09:22.277066Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.277066Z"},"links":{"cited_paper":"/paper/2411.10332","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:8727788249a9aacbb204e7614cb0d86934943ae3b58a6427e8dec5596ce5fa3e","observation_id":"7e1acb8e-75d2-4094-aa64-372b6ae565b4","resolution":{"observed_at":"2026-08-07T05:09:22.277066Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:25.426778Z","title":null,"venue":null,"work_id":"07b2dbcd-14d4-4abd-b06f-1f7cf815ef29","year":2025},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.476292Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:a3bf013c49b5f5e7dfe036d633a4b2db0877f09fc1a67e870d85459c6fd907ad","observation_id":"caded182-4d53-4ea1-bd7c-7174b021f941","resolution":{"observed_at":"2026-08-07T05:09:25.508723Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15915","last_updated":"2024-06-19T03:58:32Z","snapshot_observed_at":"2026-08-13T04:54:00.506301Z","submitted_at":"2023-12-26T07:20:55Z","title":"ChartBench: A Benchmark for Complex Visual Reasoning in Charts","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15915","snapshot_observed_at":"2026-08-07T05:09:22.590714Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.590714Z"},"links":{"cited_paper":"/paper/2312.15915","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:424d7c024571c30cb51170dcec6666b56cb99845c86ec21a68845cbcd7b073ed","observation_id":"0ee7b07b-56ff-4d14-bb91-6822defa8d9e","resolution":{"observed_at":"2026-08-07T05:09:22.590714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10671","last_updated":"2024-09-10T13:25:53Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-07-15T12:35:42Z","title":"Qwen2 Technical Report","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10671","snapshot_observed_at":"2026-08-07T05:09:22.792599Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.792599Z"},"links":{"cited_paper":"/paper/2407.10671","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:e983d5172c412338445a67dad1c019011eff33ff211e70c46984b06fc19a2a76","observation_id":"9fe46db7-7db3-4442-8aa6-ac1809a3e171","resolution":{"observed_at":"2026-08-07T05:09:22.792599Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:25.251886Z","title":null,"venue":null,"work_id":"e5ec472a-be15-4093-891a-a6c26618d8a5","year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.870253Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:73f245b4a330defccb3e42eb51e3ad464d2845a72f87548fea21f8a22dcd4615","observation_id":"1aabe9c5-9116-472c-81e7-422112b144eb","resolution":{"observed_at":"2026-08-07T05:09:25.334215Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:22.962916Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:22.962916Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:cba69858b8942a870b3944e7e31f37f7be3e9859912e7985313269e89137d3b9","observation_id":"2a01d426-545e-4ba5-b37a-309808586dc8","resolution":{"observed_at":"2026-08-07T05:09:22.962916Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:23.042401Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.042401Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:bdce2d9b2e523c63d56fb522db311b559afe7478a9dec81dfbd4b63a86a878dc","observation_id":"adadffbe-0eba-41be-8c18-fa335bfa63c7","resolution":{"observed_at":"2026-08-07T05:09:23.042401Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:25.058279Z","title":null,"venue":null,"work_id":"61729eae-ad8a-4cd2-adb5-b93928547257","year":2019},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.127849Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:32265f9e4cd485d4179deba90bb95558bead22fe85fb875ba84c189363a40f7c","observation_id":"b712d1f3-ae63-4f03-beab-9d97bf352f50","resolution":{"observed_at":"2026-08-07T05:09:25.128428Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03816","last_updated":"2024-11-18T05:36:16Z","snapshot_observed_at":"2026-08-12T23:49:03.076666Z","submitted_at":"2024-06-06T07:40:00Z","title":"ReST-MCTS*: LLM Self-Training via Process Reward Guided Tree Search","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03816","snapshot_observed_at":"2026-08-07T05:09:23.240182Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.240182Z"},"links":{"cited_paper":"/paper/2406.03816","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:40db9196c2f339b66a2a7403c52fa5fed73b8fa956a2c26be16f5fe3508dfc39","observation_id":"f4a3a9e7-5c69-4fc8-bffa-56f9115de103","resolution":{"observed_at":"2026-08-07T05:09:23.240182Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02884","last_updated":"2024-11-21T07:07:59Z","snapshot_observed_at":"2026-08-12T22:31:19.889900Z","submitted_at":"2024-10-03T18:12:29Z","title":"LLaMA-Berry: Pairwise Optimization for O1-like Olympiad-Level Mathematical Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.02884","snapshot_observed_at":"2026-08-07T05:09:23.349125Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.349125Z"},"links":{"cited_paper":"/paper/2410.02884","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:07cdb36612898443ff85cb903d664aa0fb6843f97012bab706240401faf33e36","observation_id":"f87f5801-3e66-45b7-a293-8186d484b3df","resolution":{"observed_at":"2026-08-07T05:09:23.349125Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:24.850784Z","title":null,"venue":null,"work_id":"a9d87ba7-d228-4605-9680-67ec51edd9eb","year":2025},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.429504Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:27adf440239151aed45f9f4bbb86fe7cb3ca85b9e6ee474c707ef5255dbc6ffd","observation_id":"4c15f74a-18aa-48b8-a829-3e63b4303cae","resolution":{"observed_at":"2026-08-07T05:09:24.940863Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.03493","last_updated":"2022-10-07T12:28:21Z","snapshot_observed_at":"2026-07-06T14:01:50.333970Z","submitted_at":"2022-10-07T12:28:21Z","title":"Automatic Chain of Thought Prompting in Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.03493","snapshot_observed_at":"2026-08-07T05:09:23.513166Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.513166Z"},"links":{"cited_paper":"/paper/2210.03493","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:0730224ecb4db3d9a001c958acd14a24ba8bf3a3906b30ae8e7ce51e47463ad8","observation_id":"ba1fb63a-f1e6-4641-aa47-f9c3ad038dcd","resolution":{"observed_at":"2026-08-07T05:09:23.513166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.00923","last_updated":"2024-05-20T06:43:48Z","snapshot_observed_at":"2026-08-11T00:52:12.225793Z","submitted_at":"2023-02-02T07:51:19Z","title":"Multimodal Chain-of-Thought Reasoning in Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.00923","snapshot_observed_at":"2026-08-07T05:09:23.620651Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.620651Z"},"links":{"cited_paper":"/paper/2302.00923","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:ff5182a053899e991b6abe1ca87233ac8df7a5366579be2e70b82148801dcc51","observation_id":"133a0312-fd14-48e7-a0bb-71019c467bc8","resolution":{"observed_at":"2026-08-07T05:09:23.620651Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12742","last_updated":"2024-06-18T16:02:18Z","snapshot_observed_at":"2026-08-12T23:39:56.903028Z","submitted_at":"2024-06-18T16:02:18Z","title":"Benchmarking Multi-Image Understanding in Vision and Language Models: Perception, Knowledge, Reasoning, and Multi-Hop Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12742","snapshot_observed_at":"2026-08-07T05:09:23.728013Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.728013Z"},"links":{"cited_paper":"/paper/2406.12742","citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:5c80309fd4e0706eb37679b4e1af9ac118d5fce82a00eb36a6527f7f9e0e960e","observation_id":"00d1e608-90c7-4f33-8808-8973a85ceb53","resolution":{"observed_at":"2026-08-07T05:09:23.728013Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:24.665826Z","title":null,"venue":null,"work_id":"567d36f2-332f-4b77-be9b-b8a55811c128","year":2023},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.812478Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:65ae6fd0036f528625303230b6d39c8c5efe9db6b20707b7555181fc3c31c476","observation_id":"4c852993-e185-4769-a192-ca990e3e4728","resolution":{"observed_at":"2026-08-07T05:09:24.769823Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:23.911059Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:23.911059Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:fe94fddfd396087ae015cc7c01dca7e1c3d93261b0b134af006edbbba7b84913","observation_id":"b55daa42-dcb9-4598-b6f7-af149e0545d6","resolution":{"observed_at":"2026-08-07T05:09:23.911059Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:09:24.020605Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-07T05:09:24.020605Z"},"links":{"citing_paper":"/paper/2506.08691"},"observation_digest":"sha256:b33c624cede54045d550eb7aae7820a4899f7297730be54a88adbb0b2d6e0880","observation_id":"d416c5df-3027-4e69-a81a-896a28d1c2ca","resolution":{"observed_at":"2026-08-07T05:09:24.020605Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.08691","last_updated":"2025-06-10T11:02:36Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-08T00:05:08.877544Z","submitted_at":"2025-06-10T11:02:36Z","title":"VReST: Enhancing Reasoning in Large Vision-Language Models through Tree Search and Self-Reward Mechanism"},"reference_resolution":{"displayed":45,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":45,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":45},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 45 of 45 outbound references and 1 inbound Pith citation observation for arXiv:2506.08691."}