{"as_of":"2026-08-20T04:05:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b0e91cd240c24d3f0675f3fa1eba3574d5dc9bc2433ff9542584a4ff529fe3d8","coverage":[{"denominator":49,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":49,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T14:29:05.241019Z","state":"measured"},{"denominator":64,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":64,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":15,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":15,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T11:24:33.123081Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-23T04:32:32.800150Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-10T18:36:55.267690Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.11223","last_updated":"2025-06-11T13:19:22Z","snapshot_observed_at":"2026-08-17T14:20:07.900972Z","submitted_at":"2025-01-20T02:16:19Z","title":"Reasoning Language Models: A Blueprint","version":4},"reference_index":179,"source":"pdf_text","source_observed_at":"2026-08-10T18:36:55.267690Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2501.11223"},"observation_digest":"sha256:0006b288e60c2323465ce4153de20f6a12d4c5d6feba80be9f87354f86853f43","observation_id":"82193361-4bc8-42a1-bc94-2b4db7320522","resolution":{"observed_at":"2026-08-10T18:36:55.267690Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":"2412.11936","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A survey of mathematical reasoning in the era of multimodal large language model: Benchmark, method & challenges","venue":null,"work_id":"a4d9b5b2-8a9a-41b7-b8d7-7039814aac89","year":2024},"citing_paper":{"arxiv_id":"2502.02871","last_updated":"2026-04-20T02:18:01Z","snapshot_observed_at":"2026-08-17T11:21:07.508926Z","submitted_at":"2025-02-05T04:05:27Z","title":"Position: Multimodal Large Language Models Can Significantly Advance Scientific Reasoning","version":2},"reference_index":223,"source":"arxiv_source","source_observed_at":"2026-05-23T04:30:38.804702Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2502.02871"},"observation_digest":"sha256:47ff6ec2eab8028f59652f66b0fed0faf987b422b10067043a7ef8652f126a4a","observation_id":"8fe2dede-2be1-4b50-9830-bc703c188f0f","resolution":{"observed_at":"2026-05-23T04:32:32.803728Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-16T11:24:11.970200Z","title":"A survey of math- ematical reasoning in the era of multimodal large language model: Benchmark, method & challenges,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.15585","last_updated":"2025-06-09T02:36:20Z","snapshot_observed_at":"2026-08-17T21:55:02.668814Z","submitted_at":"2025-04-22T05:02:49Z","title":"A Comprehensive Survey in LLM(-Agent) Full Stack Safety: Data, Training and Deployment","version":4},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-16T11:24:11.970200Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2504.15585"},"observation_digest":"sha256:615a5336c85e8813a5053ce139ddbd7b09b7da93bb4cc2f003e7156b1239c946","observation_id":"71045df9-91b9-452a-a982-4d4a36911c30","resolution":{"observed_at":"2026-08-16T11:24:11.970200Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-15T20:51:31.553533Z","title":"A survey of mathematical reasoning in the era of multimodal large language model: Benchmark, method & challenges.arXiv preprint arXiv:2412.11936, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.11907","last_updated":"2025-05-17T08:48:40Z","snapshot_observed_at":"2026-08-18T08:56:33.882910Z","submitted_at":"2025-05-17T08:48:40Z","title":"Are Multimodal Large Language Models Ready for Omnidirectional Spatial Reasoning?","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T20:51:31.553533Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2505.11907"},"observation_digest":"sha256:810e91d76154df31e02306a227b0ee1a37e0fa842b520b53450265976ec31de9","observation_id":"6df75775-2aec-45fd-b1fb-a298f94f6a80","resolution":{"observed_at":"2026-08-15T20:51:31.553533Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-07T15:42:37.737664Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.13965","last_updated":"2025-05-20T06:05:56Z","snapshot_observed_at":"2026-08-14T19:51:03.387786Z","submitted_at":"2025-05-20T06:05:56Z","title":"CAFES: A Collaborative Multi-Agent Framework for Multi-Granular Multimodal Essay Scoring","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-08-07T15:42:37.737664Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2505.13965"},"observation_digest":"sha256:57df6602018d77c792101d4a060e3b7f13659cfa27c32289d387ee10ef300dee","observation_id":"85ca2948-0a06-4aa2-a7c7-f3df5c4d17da","resolution":{"observed_at":"2026-08-07T15:42:37.737664Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-07T15:42:42.587497Z","title":"A survey of mathematical reasoning in the era of multimodal large language model: Benchmark, method & challenges","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.14197","last_updated":"2025-05-20T10:55:26Z","snapshot_observed_at":"2026-08-13T08:14:54.005493Z","submitted_at":"2025-05-20T10:55:26Z","title":"Towards Omnidirectional Reasoning with 360-R1: A Dataset, Benchmark, and GRPO-based Method","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-07T15:42:42.587497Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2505.14197"},"observation_digest":"sha256:54dabe53a6c31ed3099e605eafa1e97795613c859538d678f604be57254cbe6f","observation_id":"a5c85ee0-ca3d-4a0d-ba8d-3ebe4f30f42f","resolution":{"observed_at":"2026-08-07T15:42:42.587497Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-07T15:39:00.698793Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.14340","last_updated":"2025-05-20T13:27:17Z","snapshot_observed_at":"2026-08-19T20:16:07.372955Z","submitted_at":"2025-05-20T13:27:17Z","title":"Plane Geometry Problem Solving with Multi-modal Reasoning: A Survey","version":1},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-08-07T15:39:00.698793Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2505.14340"},"observation_digest":"sha256:d4a5117bfdfe5c4fbeee698592390b25b8f443abac318f01a4aee80253d429aa","observation_id":"5b5a7fe5-dd8f-4268-bb38-eaeb8f7e9bd6","resolution":{"observed_at":"2026-08-07T15:39:00.698793Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-07T15:39:20.521456Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.14406","last_updated":"2025-09-09T05:59:48Z","snapshot_observed_at":"2026-08-16T10:26:23.700233Z","submitted_at":"2025-05-20T14:20:30Z","title":"Pierce the Mists, Greet the Sky: Decipher Knowledge Overshadowing via Knowledge Circuit Analysis","version":4},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-07T15:39:20.521456Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2505.14406"},"observation_digest":"sha256:1dd320afe16b78facb81f3f9d580383b4000e4eb8a3db7f8aed0b9cf2a98db50","observation_id":"8aa1ca2e-288e-453a-b26f-315745af4d4e","resolution":{"observed_at":"2026-08-07T15:39:20.521456Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-07T13:43:23.839050Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21191","last_updated":"2025-05-27T13:40:28Z","snapshot_observed_at":"2026-08-14T10:12:10.764232Z","submitted_at":"2025-05-27T13:40:28Z","title":"Unveiling Instruction-Specific Neurons & Experts: An Analytical Framework for LLM's Instruction-Following Capabilities","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-07T13:43:23.839050Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2505.21191"},"observation_digest":"sha256:02af55b835c3ba7b8ec446d8d915f5f659e8903e28f5a83c5c1ebd027268c609","observation_id":"0d5b4e13-f700-482d-b896-2f0c99511bb8","resolution":{"observed_at":"2026-08-07T13:43:23.839050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-16T11:24:33.123081Z","title":"A survey of mathematical reasoning in the era of multimodal large language model: Benchmark, method & challenges, December 2024 b","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.06282","last_updated":"2025-04-22T07:25:03Z","snapshot_observed_at":"2026-08-19T08:39:44.052134Z","submitted_at":"2025-04-22T07:25:03Z","title":"Understanding Financial Reasoning in AI: A Multimodal Benchmark and Error Learning Approach","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-16T11:24:33.123081Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2506.06282"},"observation_digest":"sha256:785c8e24bcd3455c5bd640b55bad51988860ddd122f8b9dd000c14d7a75a0f19","observation_id":"1c2ee46f-164e-4739-9ac5-5f2c1d463460","resolution":{"observed_at":"2026-08-16T11:24:33.123081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-07T05:14:47.511366Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.08446","last_updated":"2025-06-10T04:44:28Z","snapshot_observed_at":"2026-08-17T06:22:42.572723Z","submitted_at":"2025-06-10T04:44:28Z","title":"A Survey on Large Language Models for Mathematical Reasoning","version":1},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-08-07T05:14:47.511366Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2506.08446"},"observation_digest":"sha256:08f06af0c503acd7563d7fb667071247bcd750ee2e92637c07771e29ffb12cee","observation_id":"96b8d507-a2f3-417b-807c-1d4d8df6c155","resolution":{"observed_at":"2026-08-07T05:14:47.511366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-06T21:52:50.011424Z","title":"A survey of mathematical reasoning in the era of multimodal large language model: Benchmark, method & challenges,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.23128","last_updated":"2025-06-29T07:37:49Z","snapshot_observed_at":"2026-08-16T21:31:13.878902Z","submitted_at":"2025-06-29T07:37:49Z","title":"Are Large Language Models Capable of Deep Relational Reasoning? Insights from DeepSeek-R1 and Benchmark Comparisons","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T21:52:50.011424Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2506.23128"},"observation_digest":"sha256:309069af82d6f409d94e07c7a60eb0d43e1a0ccbabc84dc44a88cf92a1960db9","observation_id":"eaa1e2b9-109e-4539-9e06-d948c0c055f4","resolution":{"observed_at":"2026-08-06T21:52:50.011424Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-06T00:59:53.783089Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04088","last_updated":"2025-08-07T03:52:48Z","snapshot_observed_at":"2026-08-08T12:10:28.774282Z","submitted_at":"2025-08-06T05:10:29Z","title":"GM-PRM: A Generative Multimodal Process Reward Model for Multimodal Mathematical Reasoning","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-06T00:59:53.783089Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2508.04088"},"observation_digest":"sha256:4ee3d84d4e09a74272328c7971d9f3f2b1beaa3bdade8dc9650920e1031a093c","observation_id":"37c2ba42-2db1-4040-a1bc-a13d2cc8e47c","resolution":{"observed_at":"2026-08-06T00:59:53.783089Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":"2412.11936","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A survey of mathematical reasoning in the era of multimodal large language model: Benchmark, method & challenges","venue":null,"work_id":"a4d9b5b2-8a9a-41b7-b8d7-7039814aac89","year":2024},"citing_paper":{"arxiv_id":"2508.06226","last_updated":"2026-05-14T12:18:38Z","snapshot_observed_at":"2026-08-16T17:53:19.592430Z","submitted_at":"2025-08-08T11:11:37Z","title":"GeoLaux: A Benchmark for Evaluating MLLMs' Geometry Performance on Long-Step Problems Requiring Auxiliary Lines","version":4},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-19T00:38:43.897231Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2508.06226"},"observation_digest":"sha256:2be448c7660df74079224749750acfb60b27673fc2a7014ce8c37eee61dea2a7","observation_id":"73c582c6-f868-4444-b2fe-27fc3443fa46","resolution":{"observed_at":"2026-05-19T00:41:56.049889Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11936","snapshot_observed_at":"2026-08-01T18:29:33.508768Z","title":"arXiv:2412.11936","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.17299","last_updated":"2026-07-19T15:31:13Z","snapshot_observed_at":"2026-08-15T19:13:30.871033Z","submitted_at":"2026-07-19T15:31:13Z","title":"WAR: Workload-Aware Rollouts for Synchronous Agentic Reinforcement Learning","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-01T18:29:33.508768Z"},"links":{"cited_paper":"/paper/2412.11936","citing_paper":"/paper/2607.17299"},"observation_digest":"sha256:4784f533720cd0dc269cfb5cbf65cf307fe0e6389047c2f0a7fd3f9a4cdc3c3c","observation_id":"f34545ab-8bf3-4f76-a439-9cd855772702","resolution":{"observed_at":"2026-08-01T18:29:33.508768Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2412.11936/citation-record","integrity":"/paper/2412.11936/integrity","json":"/paper/2412.11936/citation-record.json","paper":"/paper/2412.11936"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:06.071559Z","title":null,"venue":null,"work_id":"672d0fe2-c05d-4b43-8e21-e6782c7529d3","year":2020},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.107238Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:5219cbe46bedbe27c941c5f61fb362fa322274844dbfb24ca8f4fd53b533e655","observation_id":"dba59bef-4069-472a-b74c-f344e7c58290","resolution":{"observed_at":"2026-08-11T14:29:06.076096Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:06.057135Z","title":null,"venue":null,"work_id":"dcabf1f4-f9dd-4940-a564-53985b545f3b","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.110683Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:bbb62bdc1a8caef1ae7216d7dd75bf06e347fdf08c2276ad884c1412162b69aa","observation_id":"8ecd0505-1d87-42a7-9340-5b828ef930b8","resolution":{"observed_at":"2026-08-11T14:29:06.062169Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:06.043393Z","title":null,"venue":null,"work_id":"7331ee31-3263-4109-8368-8f970dadcf9e","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.113951Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:0e625d6f8e2fca4a5d27b537620d7966fd9c1f15d573decfe067b999e7e57411","observation_id":"97be51ba-aebc-406b-bc6c-3bed9092f080","resolution":{"observed_at":"2026-08-11T14:29:06.048299Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.11074","last_updated":"2025-05-20T08:54:49Z","snapshot_observed_at":"2026-08-16T13:17:52.665057Z","submitted_at":"2024-09-17T11:03:46Z","title":"RoMath: A Mathematical Reasoning Benchmark in Romanian","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.11074","snapshot_observed_at":"2026-08-11T14:29:05.004169Z","title":"arXiv preprint arXiv:2409.11074","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.004169Z"},"links":{"cited_paper":"/paper/2409.11074","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:052fb89d53c4a005794f35ef059e410e3966dff018ab6f4752181c3f070cd541","observation_id":"39a73c8e-ddb0-44f6-bde8-197070c9fd57","resolution":{"observed_at":"2026-08-11T14:29:05.004169Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:06.009633Z","title":"Yes\" or","venue":null,"work_id":"4d3b6b56-62ce-4b4a-8362-a6ad64518262","year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.121229Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:e00c6aef58041870a6ce61b2d88ee34cceada401f53ee3c91668c99a91753953","observation_id":"99e823c4-d10a-44b7-8267-03c14e08053b","resolution":{"observed_at":"2026-08-11T14:29:06.019128Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.10690","last_updated":"2025-03-29T12:18:26Z","snapshot_observed_at":"2026-08-18T08:55:36.196427Z","submitted_at":"2024-04-16T16:10:23Z","title":"MathWriting: A Dataset For Handwritten Mathematical Expression Recognition","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.10690","snapshot_observed_at":"2026-08-11T14:29:05.013674Z","title":"arXiv preprint arXiv:2404.10690","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.013674Z"},"links":{"cited_paper":"/paper/2404.10690","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:8d83878124a8a9db6589a4d6d68e60523c67a7783dd6322fc6d80d150c5a37f5","observation_id":"8f3609ab-92a0-4350-b45f-acce94fb949a","resolution":{"observed_at":"2026-08-11T14:29:05.013674Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16720","last_updated":"2026-04-30T02:46:40Z","snapshot_observed_at":"2026-08-19T11:46:55.171293Z","submitted_at":"2024-12-21T18:04:31Z","title":"OpenAI o1 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16720","snapshot_observed_at":"2026-08-11T14:29:05.018816Z","title":"arXiv preprint arXiv:2412.16720","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.018816Z"},"links":{"cited_paper":"/paper/2412.16720","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:b1e9ced0bb64ef30dcdc9a82533e0e9d697f20899328328dac4805fe053db8d7","observation_id":"1a26adb0-27c2-4099-bdec-56177def482e","resolution":{"observed_at":"2026-08-11T14:29:05.018816Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.01036","last_updated":"2023-11-02T07:03:25Z","snapshot_observed_at":"2026-08-18T08:56:27.534967Z","submitted_at":"2023-11-02T07:03:25Z","title":"ATHENA: Mathematical Reasoning with Thought Expansion","version":1},"cited_work":{"arxiv_id":"2311.01036","doi":null,"metadata_source":"pith","pith_arxiv_id":"2311.01036","snapshot_observed_at":"2026-08-11T14:29:05.570035Z","title":"ATHENA: Mathematical Reasoning with Thought Expansion","venue":"cs.CL","work_id":"08264d71-c23a-4bc1-a3c5-3e76a63d6860","year":2023},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.029306Z"},"links":{"cited_paper":"/paper/2311.01036","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:348f482d9285042b4150c7625ef9bc01ef9c50859da20cdd1cab810c2ba8bac9","observation_id":"2ba5b79c-1a85-47ad-8470-fae1c7eee7eb","resolution":{"observed_at":"2026-08-11T14:29:05.575607Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12863","last_updated":"2025-03-10T14:24:29Z","snapshot_observed_at":"2026-08-18T11:12:13.540506Z","submitted_at":"2024-07-12T13:16:50Z","title":"Token-Supervised Value Models for Enhancing Mathematical Problem-Solving Capabilities of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12863","snapshot_observed_at":"2026-08-11T14:29:05.034781Z","title":"Advances in neural information processing systems, 35:26337–26349","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.034781Z"},"links":{"cited_paper":"/paper/2407.12863","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:f4eedc3e70b26c4c89fa72e2b93d2f063fc7aa847a8ae4d30403f88e9287598e","observation_id":"3c05723f-389a-4926-80e6-a6452dd96748","resolution":{"observed_at":"2026-08-11T14:29:05.034781Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05229","last_updated":"2025-08-27T16:24:39Z","snapshot_observed_at":"2026-08-16T08:54:56.543625Z","submitted_at":"2024-10-07T17:36:37Z","title":"GSM-Symbolic: Understanding the Limitations of Mathematical Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05229","snapshot_observed_at":"2026-08-11T14:29:05.061961Z","title":"arXiv preprint arXiv:2410.05229","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.061961Z"},"links":{"cited_paper":"/paper/2410.05229","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:e38ac1a7d9b22214dd85797a80b7d6258742c73b0642084fed9b876c314abb2a","observation_id":"5e137c8f-3e85-4ebb-8143-dd8c7d61a794","resolution":{"observed_at":"2026-08-11T14:29:05.061961Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06563","last_updated":"2025-03-02T16:38:28Z","snapshot_observed_at":"2026-08-18T08:54:22.544112Z","submitted_at":"2025-02-10T15:31:54Z","title":"Large Language Models Meet Symbolic Provers for Logical Reasoning Evaluation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06563","snapshot_observed_at":"2026-08-11T14:29:05.067239Z","title":"arXiv preprint arXiv:2502.06563","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.067239Z"},"links":{"cited_paper":"/paper/2502.06563","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:fc0abb6426855adf2184300a93b7264ca22f8aa28fbff2b230160d37154cf5ce","observation_id":"74039854-4833-48c0-b298-fe04cc271d58","resolution":{"observed_at":"2026-08-11T14:29:05.067239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.06405","last_updated":"2024-04-11T14:37:29Z","snapshot_observed_at":"2026-08-16T14:02:20.434152Z","submitted_at":"2024-04-09T15:54:00Z","title":"Wu's Method can Boost Symbolic AI to Rival Silver Medalists and AlphaGeometry to Outperform Gold Medalists at IMO Geometry","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.06405","snapshot_observed_at":"2026-08-11T14:29:05.071773Z","title":"arXiv preprint arXiv:2404.06405","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.071773Z"},"links":{"cited_paper":"/paper/2404.06405","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:b89294f8cedb0432025718f602cd5159a599a2539f2981405c5828eb160a8fe7","observation_id":"0d3b8f3c-48c8-4bfe-af81-9f0bc5cdc11f","resolution":{"observed_at":"2026-08-11T14:29:05.071773Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.07594","last_updated":"2025-01-08T02:33:37Z","snapshot_observed_at":"2026-08-16T14:44:25.026948Z","submitted_at":"2023-11-10T09:51:24Z","title":"How to Bridge the Gap between Modalities: Survey on Multimodal Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.07594","snapshot_observed_at":"2026-08-11T14:29:05.077208Z","title":"arXiv preprint arXiv:2311.07594","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.077208Z"},"links":{"cited_paper":"/paper/2311.07594","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:c3e6a0223819882df8b4d85b866e219f02c2ff04b19c092209ff55122953ead5","observation_id":"ae4675dc-c168-4270-b17a-f8975df4f775","resolution":{"observed_at":"2026-08-11T14:29:05.077208Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.14219","last_updated":"2024-10-31T03:24:06Z","snapshot_observed_at":"2026-08-18T01:08:49.307435Z","submitted_at":"2024-06-20T11:37:53Z","title":"Proving Olympiad Algebraic Inequalities without Human Demonstrations","version":2},"cited_work":{"arxiv_id":"2406.14219","doi":null,"metadata_source":"pith","pith_arxiv_id":"2406.14219","snapshot_observed_at":"2026-08-11T14:29:05.358970Z","title":"Proving Olympiad Algebraic Inequalities without Human Demonstrations","venue":"cs.AI","work_id":"b46caf82-aaa4-4302-a3c2-586b3644cd55","year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.082766Z"},"links":{"cited_paper":"/paper/2406.14219","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:d2aae82e91d1226edef6d34aff0567855f2490c2f337ef0812be2d5faa972c8a","observation_id":"5205e4a4-ebd8-4836-9aad-4842212751ab","resolution":{"observed_at":"2026-08-11T14:29:05.367816Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14333","last_updated":"2024-05-23T09:03:42Z","snapshot_observed_at":"2026-08-16T13:50:18.721796Z","submitted_at":"2024-05-23T09:03:42Z","title":"DeepSeek-Prover: Advancing Theorem Proving in LLMs through Large-Scale Synthetic Data","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.14333","snapshot_observed_at":"2026-08-11T14:29:05.092179Z","title":"In Proceedings of the 1st Workshop on Large Generative Models Meet Multimodal Applications, pages 23–33","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.092179Z"},"links":{"cited_paper":"/paper/2405.14333","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:083b719aafd866309f265ca15c608889544e4d56e28c93349c42cda57c93d6fc","observation_id":"6c1acb48-a2b3-44e6-80f5-9c98c69cdd80","resolution":{"observed_at":"2026-08-11T14:29:05.092179Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.05653","last_updated":"2023-10-03T02:48:42Z","snapshot_observed_at":"2026-08-18T06:58:51.738907Z","submitted_at":"2023-09-11T17:47:22Z","title":"MAmmoTH: Building Math Generalist Models through Hybrid Instruction Tuning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.05653","snapshot_observed_at":"2026-08-11T14:29:05.097298Z","title":"arXiv preprint arXiv:2309.05653","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.097298Z"},"links":{"cited_paper":"/paper/2309.05653","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:a973efb6956c6bb9dc25d2841b96da23545c302c1f4721729ddf512f7a8d7b21","observation_id":"e62a1cc0-3c9f-4ca8-8ad3-b97b642f8995","resolution":{"observed_at":"2026-08-11T14:29:05.097298Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.12960","last_updated":"2023-10-19T17:56:40Z","snapshot_observed_at":"2026-08-16T14:50:36.993507Z","submitted_at":"2023-10-19T17:56:40Z","title":"SEGO: Sequential Subgoal Optimization for Mathematical Problem-Solving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.12960","snapshot_observed_at":"2026-08-11T14:29:05.102595Z","title":"In Proceedings of the 28th ACM SIGKDD Conference on Knowledge Discovery and Data Mining, pages 4571–4581","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.102595Z"},"links":{"cited_paper":"/paper/2310.12960","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:176b44c7b773ca4e39546041524c53416f0b15519d1c285907b35a8982e8bc34","observation_id":"cf14c607-42d8-4e3e-a264-24b39842d7f6","resolution":{"observed_at":"2026-08-11T14:29:05.102595Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:06.028344Z","title":null,"venue":null,"work_id":"02f94d12-c6ef-45e2-b8fb-31787d09102c","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.117743Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:8addfc103ded4e914ca0bc4b34fc0686135e8c4b524dad999c9012f66736a678","observation_id":"f126ae01-c912-468f-9899-8d718e7c2bc2","resolution":{"observed_at":"2026-08-11T14:29:06.032432Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.996250Z","title":null,"venue":null,"work_id":"a0868f31-797c-4443-aa65-ecb58af557f5","year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.126110Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:178c903004b3bb85fc7aad1a4e890b2259d73ece24f6418db584c775bfa69895","observation_id":"d1ede7c3-80f4-49d0-a788-44747cc9c376","resolution":{"observed_at":"2026-08-11T14:29:06.001215Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.983279Z","title":"❷ Bottleneck in Data Diversity:","venue":null,"work_id":"0598a52e-3889-4196-b006-7cc58334fe00","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.130437Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:94fbf756081029a399b3ca88fe88e68c2e5e1a078d5324eea9dcbacbfb06e79a","observation_id":"f516fd9a-df6e-45eb-853a-b9fc42bb8151","resolution":{"observed_at":"2026-08-11T14:29:05.987025Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.963329Z","title":null,"venue":null,"work_id":"8bce21ce-72ce-4ac6-96f0-9458be0f42d9","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.135877Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:c7c2b36be40e2413269681cf76d9626edec1e9ba911aa5348443895d11aeea45","observation_id":"7c0344c6-3a91-4975-9fb3-68349f7e9834","resolution":{"observed_at":"2026-08-11T14:29:05.969714Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.949239Z","title":"❸ Bottleneck in Data Scale:","venue":null,"work_id":"f8441476-4ac7-4212-95f9-2f6a401f53ba","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.139316Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:3b7c29cce4c2433aa27a57bef9fc29b51c5ca62eac181cfa0f8cd679fda40194","observation_id":"7fbdf47a-d43f-4d72-8a29-8f07dccf4b64","resolution":{"observed_at":"2026-08-11T14:29:05.953781Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.933831Z","title":null,"venue":null,"work_id":"6d59385a-51a6-43cc-a754-2bc64ab5072a","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.143243Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:75f02cbee7cbd3ab9e4ba2932f3b22bd94166edd9e307263c92701fa8887eaa9","observation_id":"42ae1ead-5673-4829-8ea2-84d0008536a4","resolution":{"observed_at":"2026-08-11T14:29:05.938432Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.922001Z","title":"❹ Based on recent trends in the latest works, we further propose the following actionable sugges- tions to address these dataset bottlenecks:","venue":null,"work_id":"86507954-7796-467e-9696-1bd648bd22e8","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.146955Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:1821c458fd35026b737081e575e80f008ca73b0367a3dbce10ea2d8ed31acb23","observation_id":"f1fa7cb8-3a00-4e12-84cb-de91dca62478","resolution":{"observed_at":"2026-08-11T14:29:05.926069Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.908949Z","title":null,"venue":null,"work_id":"eddbf6cf-f1e8-4e65-a7b9-c81cb2659a18","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.150886Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:6dcdce0344a901179e7d611aefd3e23a5282e5ea3a21cdac8f2c14bc891477e9","observation_id":"0509f27b-61dd-4520-a39b-4d72047310c3","resolution":{"observed_at":"2026-08-11T14:29:05.913504Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.897668Z","title":null,"venue":null,"work_id":"86f7350c-ff87-44ca-a11c-10a0706651aa","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.154775Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:6aefb386c5aa27f354e6acc44b7723b91643718dc6262b80d8108c02deaec7cb","observation_id":"a6ea1e75-6233-4dea-8ead-24338f149242","resolution":{"observed_at":"2026-08-11T14:29:05.901463Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.887019Z","title":null,"venue":null,"work_id":"5780eedf-3d30-4c96-b4f1-4b4d8b444163","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.159774Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:573f83ad5559c8e4da0ffd4eb62ed464e6c22990bd624608f4d2100a2f32b983","observation_id":"34f89647-9b89-4ecc-a502-643d4901d3ba","resolution":{"observed_at":"2026-08-11T14:29:05.890428Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.874976Z","title":"spatial reason- ing","venue":null,"work_id":"d8b935ff-f49f-4193-b2b3-3797c458dcc8","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.165558Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:be92e7411bf06ab560db015d7610ac2f47b64061bab3976074c0a1abebc98d05","observation_id":"d2387a89-a605-4f65-8a6e-98ead558ba33","resolution":{"observed_at":"2026-08-11T14:29:05.879433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.860719Z","title":"Use domain-specific few-shot examples (e.g., providing figure-text associations in geometry) to guide the model in switching reasoning modes","venue":null,"work_id":"e15066c8-d9e2-4037-9730-619c66faf9bd","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.171547Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:9706a2fc12f361539882898299858025ea918ca165f73e1bf8a82ca9dcf29b53","observation_id":"87b450de-ba20-4746-b819-e6677eb04a31","resolution":{"observed_at":"2026-08-11T14:29:05.866455Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.846950Z","title":"triangle","venue":null,"work_id":"cc1e3171-62d5-45e2-8f13-e87411383a46","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.176510Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:cfa57f21dba2bcab63a246d8ea12dfdb853d0981f3ce9e201d2245a85964bc05","observation_id":"dc4caeb2-0385-4b9c-b915-3f51b631eedd","resolution":{"observed_at":"2026-08-11T14:29:05.851907Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.832776Z","title":"The limitation is that labeling error types is costly and it’s difficult to cover all long-tail errors","venue":null,"work_id":"95d78d98-1583-438d-88d7-426ef83ab00b","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.181253Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:e706154b2fd1bbde831e3cac956e2f59f343c276c9e7b900e9ab2bf3e5e3e638","observation_id":"578fe127-015b-4dfb-86e4-4ba4bc36f342","resolution":{"observed_at":"2026-08-11T14:29:05.837944Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.821120Z","title":"computation-logic-conclusion","venue":null,"work_id":"a817ab32-9d72-4ae6-906c-0f8fe3680fc4","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.200899Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:1b59709bdd4c447137caaf661591d8e29f42f1ed95c0aa859641d06a766989f6","observation_id":"72c164b1-2f2e-4abe-aa08-1a4a876ea596","resolution":{"observed_at":"2026-08-11T14:29:05.825041Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.802345Z","title":"The limitation is that tool invoca- tion delays affect real-time performance, and some errors require manually defined detec- tion rules","venue":null,"work_id":"145e20b1-6a7d-4095-889d-e28433d64b99","year":2025},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.205542Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:6f1584b4031594333ad5e05d3c80a93cf5439da5409d44532b4b416bced787ad","observation_id":"617e69fa-b12c-459d-8fc7-86a75370afc4","resolution":{"observed_at":"2026-08-11T14:29:05.807600Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.790075Z","title":null,"venue":null,"work_id":"465c24bc-c200-45c6-9e96-96dcc4351b99","year":2025},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.209565Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:2d4cebc58be4a54299c04dc52d8eb922e7f093b2f8cea8bcca597840ccb4e1ed","observation_id":"17331e94-06c2-48d6-ba1c-c777f22b8354","resolution":{"observed_at":"2026-08-11T14:29:05.794363Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.776957Z","title":null,"venue":null,"work_id":"d725daeb-96df-48f5-8f58-47463730a7af","year":2025},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.214385Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:58c134eeb009f1382063f862bf2d426a998bc92a71573773c5bdc109dfb77568","observation_id":"b3132461-8613-4f86-adea-aed2e29164c6","resolution":{"observed_at":"2026-08-11T14:29:05.780786Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.763402Z","title":"If the training data coverage is insufficient, test-time strategies may not be able to compensate (Ke et al., 2025; Chen et al., 2025c)","venue":null,"work_id":"47a9fbac-ea96-4803-bf20-6feee762abae","year":2025},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.218715Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:e6b1c0c140cc4a28e70e5ac297050d4470a180d1c9818dbe77e0c4065daf41f2","observation_id":"77173011-fb98-4b4d-95be-5e327b63df6b","resolution":{"observed_at":"2026-08-11T14:29:05.768170Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.748273Z","title":"trian- gle","venue":null,"work_id":"d17876aa-b881-4667-a987-4e0ba384f708","year":2025},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.223162Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:9e5316fef14b5c16ff68a4449d6816cfcddeb63514e76dd6d0bd1121780c82d0","observation_id":"61ec9488-f54b-4770-b841-6691ff24c0d7","resolution":{"observed_at":"2026-08-11T14:29:05.752711Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.734875Z","title":null,"venue":null,"work_id":"1bf37782-b466-466e-8de2-600ae78ae13a","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.228180Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:5df2f09503c10b66ad409d1c69bb5dc016f70d4425522f8e46e75c8a3737254b","observation_id":"3433ae2a-0a2f-4284-8b6b-2ebc63c08367","resolution":{"observed_at":"2026-08-11T14:29:05.738351Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.722900Z","title":"❸ Error Feedback Limitations:","venue":null,"work_id":"136accce-6459-4756-bf36-c809643bf445","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.232068Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:c6c4ce92484b32f32e79bf90b504d4b62f29a8010b48d4a07316263565570cc2","observation_id":"385697d0-fe96-42a0-86e6-ea305380a1ea","resolution":{"observed_at":"2026-08-11T14:29:05.726506Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.710263Z","title":null,"venue":null,"work_id":"b8c61de6-b4ac-40fc-82c2-9dd704965acf","year":2025},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.236565Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:322fb512583406ed403b7044982b5a057fcfe3c483ae88c40afe261c48b7f9f7","observation_id":"cf36e3f8-1cf0-47f0-99b0-e5df23e2bebd","resolution":{"observed_at":"2026-08-11T14:29:05.714406Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.695554Z","title":"Long-tail errors such as rare symbol confusions may be overlooked","venue":null,"work_id":"edcb0c98-34a5-463c-810e-5265d9ae3040","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.241019Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:60d803e0de2a4ca1891b750509ce683aa783e6306796c164d8ea83111a66890b","observation_id":"dc4f8666-a545-44b0-94fa-83dfa2de3e96","resolution":{"observed_at":"2026-08-11T14:29:05.699129Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.00157","last_updated":"2024-09-16T19:20:59Z","snapshot_observed_at":"2026-08-16T14:22:43.713388Z","submitted_at":"2024-01-31T20:26:32Z","title":"Large Language Models for Mathematical Reasoning: Progresses and Challenges","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.00157","snapshot_observed_at":"2026-08-11T14:29:04.986025Z","title":"Janice Ahn, Rishu Verma, Renze Lou, Di Liu, Rui Zhang, and Wenpeng Yin","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":147,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:04.986025Z"},"links":{"cited_paper":"/paper/2402.00157","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:959de1539fe75fd6defb4042beed0d67bb70486f2d217d829330e50434066821","observation_id":"2e35f305-c370-4cec-9ace-74211b0a6271","resolution":{"observed_at":"2026-08-11T14:29:04.986025Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1905.13319","last_updated":"2019-05-30T21:28:12Z","snapshot_observed_at":"2026-08-16T02:16:32.508169Z","submitted_at":"2019-05-30T21:28:12Z","title":"MathQA: Towards Interpretable Math Word Problem Solving with Operation-Based Formalisms","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1905.13319","snapshot_observed_at":"2026-08-11T14:29:04.991803Z","title":"arXiv preprint arXiv:1905.13319","venue":null,"work_id":null,"year":1905},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:04.991803Z"},"links":{"cited_paper":"/paper/1905.13319","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:be1e2fa744d076679e474346b65c3fe215669f615b625a37a9ef3bad25e17d3c","observation_id":"2d6b954e-ddac-4748-a237-4c8a0eb053d6","resolution":{"observed_at":"2026-08-11T14:29:04.991803Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.13435","last_updated":"2022-05-11T16:19:25Z","snapshot_observed_at":"2026-08-16T18:06:53.798045Z","submitted_at":"2021-07-28T15:28:41Z","title":"MWP-BERT: Numeracy-Augmented Pre-training for Math Word Problem Solving","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.13435","snapshot_observed_at":"2026-08-11T14:29:05.040249Z","title":"arXiv preprint arXiv:2107.13435","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.040249Z"},"links":{"cited_paper":"/paper/2107.13435","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:fe276f7effe9a54648e63e038b12430b40cb068a6806ab35384a96d336750c9d","observation_id":"100240ca-376d-4f4f-8463-db2c396533e6","resolution":{"observed_at":"2026-08-11T14:29:05.040249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16265","last_updated":"2024-06-26T14:01:15Z","snapshot_observed_at":"2026-08-16T13:49:27.658035Z","submitted_at":"2024-05-25T15:07:33Z","title":"MindStar: Enhancing Math Reasoning in Pre-trained LLMs at Inference Time","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16265","snapshot_observed_at":"2026-08-11T14:29:05.023302Z","title":"Educational Research Review, 37:100480","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.023302Z"},"links":{"cited_paper":"/paper/2405.16265","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:b9fe2d04e211ec3e09831bd0040559cd0027148a0d7f420fe489cb538b398418","observation_id":"b82d1d28-879f-4225-8e89-516e33005d76","resolution":{"observed_at":"2026-08-11T14:29:05.023302Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:06.084549Z","title":"Advances in Neural Information Processing Systems, 36:5539–5568","venue":null,"work_id":"c985252b-53e6-484a-b2cc-7568266463f8","year":null},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.009221Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:1b7d78b1503e7ee399ff6a631d2261eb96a42df5745f0a59e53f8e4064e803f9","observation_id":"17ade6d4-05a1-499d-bb98-db823689930b","resolution":{"observed_at":"2026-08-11T14:29:06.088222Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.15736","last_updated":"2024-12-05T23:59:06Z","snapshot_observed_at":"2026-08-16T13:40:34.066153Z","submitted_at":"2024-06-22T05:04:39Z","title":"Evaluating Large Vision-and-Language Models on Children's Mathematical Olympiads","version":2},"cited_work":{"arxiv_id":"2406.15736","doi":null,"metadata_source":"pith","pith_arxiv_id":"2406.15736","snapshot_observed_at":"2026-08-11T14:29:05.645627Z","title":"Evaluating Large Vision-and-Language Models on Children's Mathematical Olympiads","venue":"cs.LG","work_id":"8ec2bd89-c55c-47a2-96aa-61d8149fcb44","year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:04.996663Z"},"links":{"cited_paper":"/paper/2406.15736","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:c71ec76ad4028adb68a0c95da044367c9693b11a11f43cf338a3def8ec9d5c41","observation_id":"efd5c1c7-8f9c-4680-b841-34e91896f3d2","resolution":{"observed_at":"2026-08-11T14:29:05.652045Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T14:29:05.057155Z","title":"arXiv preprint arXiv:2501.04686","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.057155Z"},"links":{"citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:61653a3743c2a94c2b1f14912edf5eddbcc4d83528b0d21323573e188d660e72","observation_id":"9aea032e-f4f4-49d0-8fe4-e25da0fb9f47","resolution":{"observed_at":"2026-08-11T14:29:05.057155Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.17564","last_updated":"2023-12-21T06:21:11Z","snapshot_observed_at":"2026-08-14T12:54:48.492396Z","submitted_at":"2023-03-30T17:30:36Z","title":"BloombergGPT: A Large Language Model for Finance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.17564","snapshot_observed_at":"2026-08-11T14:29:05.087342Z","title":"Shijie Wu, Ozan Irsoy, Steven Lu, Vadim Dabravolski, Mark Dredze, Sebastian Gehrmann, Prabhanjan Kam- badur, David Rosenberg, and Gideon Mann","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges","version":3},"reference_index":2256,"source":"pdf_text","source_observed_at":"2026-08-11T14:29:05.087342Z"},"links":{"cited_paper":"/paper/2303.17564","citing_paper":"/paper/2412.11936"},"observation_digest":"sha256:3ef322925e647b900a42d4c0036a63d570149be04befd2a1579b0f5d24c4cef2","observation_id":"e2f9304d-de2f-4d5a-9e88-fc8542099eea","resolution":{"observed_at":"2026-08-11T14:29:05.087342Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2412.11936","last_updated":"2025-05-20T14:45:21Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-18T08:55:39.288248Z","submitted_at":"2024-12-16T16:21:41Z","title":"A Survey of Mathematical Reasoning in the Era of Multimodal Large Language Model: Benchmark, Method & Challenges"},"reference_resolution":{"displayed":49,"state_counts":{"malformed_identifier":0,"metadata_mismatch":3,"parse_uncertain":0,"unresolved":31,"verified_exact":0,"verified_fuzzy":15},"total_outbound_references":49},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 49 of 49 outbound references and 15 inbound Pith citation observations for arXiv:2412.11936."}