{"as_of":"2026-08-13T01:24:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c5e5845bc8256950b08226ccebb30e34bcab9c893df30be9e861726459f4acae","coverage":[{"denominator":87,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":87,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T12:33:13.820971Z","state":"measured"},{"denominator":91,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":91,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":4,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":4,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T10:39:03.028373Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.24238","snapshot_observed_at":"2026-08-05T10:39:03.028373Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.03871","last_updated":"2025-09-04T04:12:31Z","snapshot_observed_at":"2026-08-08T03:54:26.940042Z","submitted_at":"2025-09-04T04:12:31Z","title":"A Comprehensive Survey on Trustworthiness in Reasoning with Large Language Models","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-05T10:39:03.028373Z"},"links":{"cited_paper":"/paper/2505.24238","citing_paper":"/paper/2509.03871"},"observation_digest":"sha256:f2fcac68fc41555a0bc6795833e8c71eb190bb4e7f36b96f22b8d2bb9f933c4a","observation_id":"8f61cce6-5c2d-4d17-8b3f-2c8cc3ad9045","resolution":{"observed_at":"2026-08-05T10:39:03.028373Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"cited_work":{"arxiv_id":"2505.24238","doi":"10.48550/arxiv.2505.24238","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.24238","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Mirage: Assessing hal- lucination in multimodal reasoning chains of mllm","venue":"ArXiv.org","work_id":"e1b2fa10-0b42-496c-a013-661376a69b10","year":2025},"citing_paper":{"arxiv_id":"2604.03179","last_updated":"2026-04-03T16:56:34Z","snapshot_observed_at":"2026-08-12T23:53:53.905795Z","submitted_at":"2026-04-03T16:56:34Z","title":"Understanding the Role of Hallucination in Reinforcement Post-Training of Multimodal Reasoning Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-13T20:48:52.130130Z"},"links":{"cited_paper":"/paper/2505.24238","citing_paper":"/paper/2604.03179"},"observation_digest":"sha256:1ba1799747820aaf3abb2884d22117d1329f376a3fcc8287f268890fd4096066","observation_id":"93d009a1-1467-4909-ae1a-aa45d4e09569","resolution":{"observed_at":"2026-05-13T20:53:16.282414Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"cited_work":{"arxiv_id":"2505.24238","doi":"10.48550/arxiv.2505.24238","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.24238","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Mirage: Assessing hal- lucination in multimodal reasoning chains of mllm","venue":"ArXiv.org","work_id":"e1b2fa10-0b42-496c-a013-661376a69b10","year":2025},"citing_paper":{"arxiv_id":"2604.21027","last_updated":"2026-08-02T16:15:42Z","snapshot_observed_at":"2026-08-10T23:41:30.183421Z","submitted_at":"2026-04-22T19:18:36Z","title":"HypEHR: Hyperbolic Modeling of Electronic Health Records for Efficient Question Answering","version":1},"reference_index":230,"source":"arxiv_source","source_observed_at":"2026-05-09T23:51:47.724033Z"},"links":{"cited_paper":"/paper/2505.24238","citing_paper":"/paper/2604.21027"},"observation_digest":"sha256:322f6fdc9b917ec65c76dceef7b287adf820ff51b4a0d57f72ccf6cbc85c1e0e","observation_id":"8db78eb4-7cbe-4125-83e5-637815967363","resolution":{"observed_at":"2026-05-09T23:54:45.651801Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"cited_work":{"arxiv_id":"2505.24238","doi":"10.48550/arxiv.2505.24238","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.24238","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Mirage: Assessing hal- lucination in multimodal reasoning chains of mllm","venue":"ArXiv.org","work_id":"e1b2fa10-0b42-496c-a013-661376a69b10","year":2025},"citing_paper":{"arxiv_id":"2604.23392","last_updated":"2026-04-25T17:46:59Z","snapshot_observed_at":"2026-08-12T10:54:39.248558Z","submitted_at":"2026-04-25T17:46:59Z","title":"SoccerRef-Agents: Multi-Agent System for Automated Soccer Refereeing","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-08T08:03:57.287297Z"},"links":{"cited_paper":"/paper/2505.24238","citing_paper":"/paper/2604.23392"},"observation_digest":"sha256:52b3084a81bbf891a1e204e763a62a22c7104e92231ac2cf8e31b4f843aa229e","observation_id":"45eb3f94-6839-4143-946f-13d83c6c4310","resolution":{"observed_at":"2026-05-11T20:46:13.723544Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.24238/citation-record","integrity":"/paper/2505.24238/integrity","json":"/paper/2505.24238/citation-record.json","paper":"/paper/2505.24238"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2404.14219","last_updated":"2024-08-30T21:17:17Z","snapshot_observed_at":"2026-08-10T14:07:02.234322Z","submitted_at":"2024-04-22T14:32:33Z","title":"Phi-3 Technical Report: A Highly Capable Language Model Locally on Your Phone","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.14219","snapshot_observed_at":"2026-08-07T12:33:06.080923Z","title":"Phi-3 technical report: A highly capable language model locally on your phone.arXiv preprint arXiv:2404.14219, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.080923Z"},"links":{"cited_paper":"/paper/2404.14219","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:1b6152c0b1f2847bb0d2cd3229a40f119535dd3f355d3a06e65c8ee5d2f4dfa2","observation_id":"cd373f87-c04c-4530-bc32-e4af6da0f6fd","resolution":{"observed_at":"2026-08-07T12:33:06.080923Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-08-12T17:29:41.806995Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-08-07T12:33:06.122153Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.122153Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:27b44883bf800857a4e30ab1425b942bc8f8b6e1c652e83bc4a4756df2ba93da","observation_id":"1f5eb507-8925-4c9a-b2f3-080a49e013e3","resolution":{"observed_at":"2026-08-07T12:33:06.122153Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.03631","last_updated":"2024-10-16T19:35:55Z","snapshot_observed_at":"2026-08-10T08:11:09.870911Z","submitted_at":"2023-12-06T17:28:03Z","title":"Mitigating Open-Vocabulary Caption Hallucinations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.03631","snapshot_observed_at":"2026-08-07T12:33:06.187286Z","title":"Mocha: Multi-objective reinforcement mitigating caption hallucinations.arXiv preprint arXiv:2312.03631, 2, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.187286Z"},"links":{"cited_paper":"/paper/2312.03631","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:4d8e4d247cb8c599f533cbeb19b9ac6d3f9387059cac365dbcd75b5dcef3b2b8","observation_id":"e235ea63-d9dd-4ac5-bd0b-e70324949e95","resolution":{"observed_at":"2026-08-07T12:33:06.187286Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00698","last_updated":"2025-06-04T16:20:49Z","snapshot_observed_at":"2026-08-13T01:19:52.290014Z","submitted_at":"2025-02-02T07:12:03Z","title":"MM-IQ: Benchmarking Human-Like Abstraction and Reasoning in Multimodal Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.00698","snapshot_observed_at":"2026-08-07T12:33:06.268894Z","title":"Mm-iq: Benchmarking human-like abstraction and reasoning in multimodal models.arXiv preprint arXiv:2502.00698, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.268894Z"},"links":{"cited_paper":"/paper/2502.00698","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:05208c91782c3b969f0fe41c15f31b4a33e154440ef850741ed5987259ca9da5","observation_id":"0ac1d8d1-47ac-46b8-a4fa-ba1fb0af3714","resolution":{"observed_at":"2026-08-07T12:33:06.268894Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:06.382734Z","title":"Internvl: Scaling up vision foundation models and aligning for generic visual-linguistic tasks","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.382734Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:7096cc1bde5d70766c9cbcd7a797730a54f1259c3fe4c3c2b40dc4e61f86007c","observation_id":"a02ce98e-e834-4a5d-8358-f84a06c5ca25","resolution":{"observed_at":"2026-08-07T12:33:06.382734Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.16479","last_updated":"2023-11-27T09:30:02Z","snapshot_observed_at":"2026-08-12T12:50:41.983286Z","submitted_at":"2023-11-27T09:30:02Z","title":"Mitigating Hallucination in Visual Language Models with Visual Supervision","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.16479","snapshot_observed_at":"2026-08-07T12:33:06.475621Z","title":"Mitigating hallucination in visual language models with visual supervision.arXiv preprint arXiv:2311.16479, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.475621Z"},"links":{"cited_paper":"/paper/2311.16479","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:1629518acfa3f00be0342655eb6da28fbd3b82185f51625e5959651fa730807e","observation_id":"6b2a2fc9-cbcb-4800-9df7-e62cbf88b301","resolution":{"observed_at":"2026-08-07T12:33:06.475621Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:06.558504Z","title":"Puzzlevqa: Diagnosing multimodal reasoning challenges of language models with abstract visual patterns","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.558504Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:885031bfc39794f286776ffef66fbf0d0477197c3d93ef05ad2c575ea6ca6881","observation_id":"60e44d62-1f95-4794-bbb8-ed07a8806100","resolution":{"observed_at":"2026-08-07T12:33:06.558504Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:06.634033Z","title":"Zero-shot generalizable incremental learning for vision-language object detection","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.634033Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:504f65a7372badcdc7fa0a8d8cbcfc3424d5e033c9b7fca7e46787a658a92a0c","observation_id":"e6c82916-4191-463e-9612-10cbbf6c873d","resolution":{"observed_at":"2026-08-07T12:33:06.634033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15979","last_updated":"2024-12-23T16:55:57Z","snapshot_observed_at":"2026-08-11T20:56:12.400151Z","submitted_at":"2024-12-20T15:22:51Z","title":"MR-GDINO: Efficient Open-World Continual Object Detection","version":2},"cited_work":{"arxiv_id":"2412.15979","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.15979","snapshot_observed_at":"2026-08-07T12:33:14.622898Z","title":"MR-GDINO: Efficient Open-World Continual Object Detection","venue":"cs.CV","work_id":"429f1ade-2c6a-426f-a961-1c3c0845fa18","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.701427Z"},"links":{"cited_paper":"/paper/2412.15979","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:6857a7c22f435a8eb5c08c51b00e7a9592b70fd10897847be3943e2cf94ffc20","observation_id":"d23fb554-b6b2-4c64-80be-3802ef89af64","resolution":{"observed_at":"2026-08-07T12:33:14.676390Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:06.750153Z","title":"Lpt: Long-tailed prompt tuning for image classification","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.750153Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:79aac44d8503d94cfe69b32bd415b443e36904e7cfb10b2019e49afb610df4f2","observation_id":"f50f399f-d566-4461-af0e-1f313c9b77a1","resolution":{"observed_at":"2026-08-07T12:33:06.750153Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.00234","last_updated":"2024-10-05T11:47:02Z","snapshot_observed_at":"2026-07-06T14:36:25.690733Z","submitted_at":"2022-12-31T15:57:09Z","title":"A Survey on In-context Learning","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.00234","snapshot_observed_at":"2026-08-07T12:33:06.816254Z","title":"A survey on in-context learning.arXiv preprint arXiv:2301.00234, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.816254Z"},"links":{"cited_paper":"/paper/2301.00234","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:3c14e11447e632c1969d703e7612de3921772d3f5eefc8818104c7dd1daff90e","observation_id":"40226a05-9e99-486a-8a25-e68cf003dc0d","resolution":{"observed_at":"2026-08-07T12:33:06.816254Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.01904","last_updated":"2025-02-05T09:17:01Z","snapshot_observed_at":"2026-08-12T03:33:59.313788Z","submitted_at":"2025-01-03T17:14:16Z","title":"Virgo: A Preliminary Exploration on Reproducing o1-like MLLM","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.01904","snapshot_observed_at":"2026-08-07T12:33:06.915732Z","title":"Virgo: A preliminary exploration on reproducing o1-like mllm.arXiv preprint arXiv:2501.01904, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:06.915732Z"},"links":{"cited_paper":"/paper/2501.01904","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:c6f88b48cd6139b02ad98becd23882e267f0ce8f788f1a0174822f1f7daa2baf","observation_id":"ea35cb9c-7fd7-4ce0-bc0d-e7d524fc5b9f","resolution":{"observed_at":"2026-08-07T12:33:06.915732Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:07.022394Z","title":"Detecting hallucinations in large language models using semantic entropy.Nature, 630(8017):625–630, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.022394Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:c90177f834f347d78ddb0251c8a6b231fa434ee90b7f7b5f150479fb17f19955","observation_id":"49e1bd80-b7be-43fc-863d-fb2bf75fdb27","resolution":{"observed_at":"2026-08-07T12:33:07.022394Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.13394","last_updated":"2025-10-24T02:45:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-23T09:22:36Z","title":"MME: A Comprehensive Evaluation Benchmark for Multimodal Large Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.13394","snapshot_observed_at":"2026-08-07T12:33:07.090750Z","title":"Mme: A comprehensive evaluation benchmark for multimodal large language models.arXiv preprint arXiv:2306.13394, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.090750Z"},"links":{"cited_paper":"/paper/2306.13394","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:2c3fa7c8c78524ec616abdaf38f25e4aea50e71c3cb5a2b27ef46af9a0bf7b36","observation_id":"b3dc9fb6-d758-4d73-af6e-464ab485d758","resolution":{"observed_at":"2026-08-07T12:33:07.090750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:07.151002Z","title":"GPTScore: Evaluate as you desire","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.151002Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:075b0e85e3b8ea322c33a46da7740ee7e559956600de34012929c2b23a77470e","observation_id":"6cd15c42-a470-4dbe-9f0d-e99b8db63c67","resolution":{"observed_at":"2026-08-07T12:33:07.151002Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:07.271503Z","title":"The capacity for moral self-correction in large language models.Parameters, 109(1010):1011","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.271503Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:f98323568b08981426ca10d278bf9555a00889183a10af8a351720afb5ef723e","observation_id":"85588250-a0aa-44c6-a5fe-3bd9c163fbff","resolution":{"observed_at":"2026-08-07T12:33:07.271503Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.03864","last_updated":"2024-03-13T00:50:05Z","snapshot_observed_at":"2026-08-12T23:07:47.884354Z","submitted_at":"2024-03-06T17:15:04Z","title":"Are Language Models Puzzle Prodigies? Algorithmic Puzzles Unveil Serious Challenges in Multimodal Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.03864","snapshot_observed_at":"2026-08-07T12:33:07.360301Z","title":"Are language models puzzle prodigies? algorithmic puzzles unveil serious challenges in multimodal reasoning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.360301Z"},"links":{"cited_paper":"/paper/2403.03864","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:e94da6ec86e0d9ca7904305123d82f1a445d0c16539b8973363084212d027da1","observation_id":"131f8f24-faff-4e6e-9cc9-ecdc554449ba","resolution":{"observed_at":"2026-08-07T12:33:07.360301Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-10T16:40:37.411115Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-07T12:33:07.475863Z","title":"The llama 3 herd of models.arXiv preprint arXiv:2407.21783, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.475863Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:ccbd9cd65945f1d81f2edacb672e5886c94aad055b0c22811d384208ababcd53","observation_id":"d385e19b-f4e6-4955-89e5-25a3370bbfea","resolution":{"observed_at":"2026-08-07T12:33:07.475863Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15594","last_updated":"2025-10-19T10:32:43Z","snapshot_observed_at":"2026-08-02T10:23:50.881300Z","submitted_at":"2024-11-23T16:03:35Z","title":"A Survey on LLM-as-a-Judge","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15594","snapshot_observed_at":"2026-08-07T12:33:07.571919Z","title":"A survey on llm-as-a-judge.arXiv preprint arXiv:2411.15594, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.571919Z"},"links":{"cited_paper":"/paper/2411.15594","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:d9484689a935fb3c348391c700f01b4e80366e6a66da476237749737b1aaf9ff","observation_id":"82fb3068-1a9f-4e79-a8f4-edb963f00afd","resolution":{"observed_at":"2026-08-07T12:33:07.571919Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:07.664840Z","title":"Hallusionbench: an advanced diagnostic suite for entangled language hallucination and visual illusion in large vision-language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.664840Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:cc3c964c41dfc5a6166a740d2a495b0f8ae2d257ed37fdb06c3c7cf38a048329","observation_id":"a4bcd11f-f250-480f-bf44-546db8c7a328","resolution":{"observed_at":"2026-08-07T12:33:07.664840Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-07T12:33:07.772797Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning.arXiv preprint arXiv:2501.12948, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.772797Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:f00db49fdc250f365d0b2e825d96955557436a777fbc3f536e0466d31eb2256e","observation_id":"aaf1ac63-d878-4cb2-807b-c8f4228efe91","resolution":{"observed_at":"2026-08-07T12:33:07.772797Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.01887","last_updated":"2025-02-27T03:39:10Z","snapshot_observed_at":"2026-08-07T17:44:15.581009Z","submitted_at":"2025-02-27T03:39:10Z","title":"When Continue Learning Meets Multimodal Large Language Model: A Survey","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.01887","snapshot_observed_at":"2026-08-07T12:33:07.869192Z","title":"When continue learning meets multimodal large language model: A survey.arXiv preprint arXiv:2503.01887, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.869192Z"},"links":{"cited_paper":"/paper/2503.01887","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:741004ace5820c74258e195d6085ca0497cd9c1aa6a52e1ce7c2f99388f9d848","observation_id":"e4ea4dc8-3afa-4b0e-a85f-889900e9f05d","resolution":{"observed_at":"2026-08-07T12:33:07.869192Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21276","last_updated":"2024-10-25T17:43:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-25T17:43:01Z","title":"GPT-4o System Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21276","snapshot_observed_at":"2026-08-07T12:33:07.959565Z","title":"Gpt-4o system card.arXiv preprint arXiv:2410.21276, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:07.959565Z"},"links":{"cited_paper":"/paper/2410.21276","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:47708ba6e0afe76ca3d0f4954e2e203323fe1ad22bbdc7a32fdb63659d67e1e3","observation_id":"3c11f0aa-8b18-422f-9198-fd055ab6103c","resolution":{"observed_at":"2026-08-07T12:33:07.959565Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16720","last_updated":"2026-04-30T02:46:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-21T18:04:31Z","title":"OpenAI o1 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16720","snapshot_observed_at":"2026-08-07T12:33:08.049492Z","title":"Openai o1 system card.arXiv preprint arXiv:2412.16720, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:08.049492Z"},"links":{"cited_paper":"/paper/2412.16720","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:db21b4305b4ab4811d9067d7fd9ac36eb91ddb4f17da52333c45707036fc772b","observation_id":"efeed5f0-0c9e-4944-a8b8-e7aa228c618e","resolution":{"observed_at":"2026-08-07T12:33:08.049492Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:18.358252Z","title":"Towards mitigating LLM hallucination via self reflection","venue":null,"work_id":"7faf17ea-8cd0-4434-9b1e-b8663d88d4cc","year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:08.183950Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:f0c921382043fe6b5996f946ee10d9859af572c00b5e016979d1af51fb253cce","observation_id":"ffa0a01f-ae22-4626-935d-4c149f169e29","resolution":{"observed_at":"2026-08-07T12:33:18.503341Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:08.290558Z","title":"Hallucination augmented contrastive learning for multimodal large language model","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:08.290558Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:029b5cb8041ca95fcd3be69778801ace4551279bb652058bdfc40a032f4ac816","observation_id":"62c94a89-a059-4823-bb1d-acd48b4486d5","resolution":{"observed_at":"2026-08-07T12:33:08.290558Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.09621","last_updated":"2025-02-13T18:59:46Z","snapshot_observed_at":"2026-08-09T00:23:29.329507Z","submitted_at":"2025-02-13T18:59:46Z","title":"MME-CoT: Benchmarking Chain-of-Thought in Large Multimodal Models for Reasoning Quality, Robustness, and Efficiency","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.09621","snapshot_observed_at":"2026-08-07T12:33:08.379110Z","title":"Mme-cot: Benchmarking chain-of-thought in large multimodal models for reasoning quality, robustness, and efficiency.arXiv preprint arXiv:2502.09621, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:08.379110Z"},"links":{"cited_paper":"/paper/2502.09621","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:db1ecd76fc6be0961066a10f28e348be7b14f731fdcebd7567d9d0cba45dfaa4","observation_id":"a72d196d-62cc-4209-83be-37aec9055dc7","resolution":{"observed_at":"2026-08-07T12:33:08.379110Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:18.048689Z","title":"Decoupling representation and classifier for long-tailed recognition","venue":null,"work_id":"d729828f-0b8c-4f1e-9914-c9f13d112344","year":null},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:08.480767Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:34d6d8ea23a6bd169e04c25b5e32830126ffea56249ab37fbcb695d8755c6348","observation_id":"d3eabd59-e4da-4a04-a2a9-c88aed94f37e","resolution":{"observed_at":"2026-08-07T12:33:18.166965Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:08.559186Z","title":"Gonzalez, Hao Zhang, and Ion Stoica","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:08.559186Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:337ea8e87ba6ca6a3e77238857e8b97e38092798d81e70451a430c7d515c8e86","observation_id":"2e8ba4f8-f945-4fb9-9e87-57bb85616c2f","resolution":{"observed_at":"2026-08-07T12:33:08.559186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.16125","last_updated":"2023-08-02T08:02:35Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-07-30T04:25:16Z","title":"SEED-Bench: Benchmarking Multimodal LLMs with Generative Comprehension","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.16125","snapshot_observed_at":"2026-08-07T12:33:08.631430Z","title":"Seed- bench: Benchmarking multimodal llms with generative comprehension.arXiv preprint arXiv:2307.16125, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:08.631430Z"},"links":{"cited_paper":"/paper/2307.16125","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:d7e1367d002820f755373ade899ebaff67045e7960b4f03679a7273acc073ab4","observation_id":"c3e2ac17-0813-463f-9f74-946b2dfbbd0c","resolution":{"observed_at":"2026-08-07T12:33:08.631430Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.05579","last_updated":"2024-12-10T05:49:12Z","snapshot_observed_at":"2026-08-11T23:22:45.789386Z","submitted_at":"2024-12-07T08:07:24Z","title":"LLMs-as-Judges: A Comprehensive Survey on LLM-based Evaluation Methods","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.05579","snapshot_observed_at":"2026-08-07T12:33:08.729239Z","title":"Llms-as-judges: a comprehensive survey on llm-based evaluation methods.arXiv preprint arXiv:2412.05579, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:08.729239Z"},"links":{"cited_paper":"/paper/2412.05579","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:45725176814ac0db0056daf307f7f6d696f137bd472faa2c0fb6f1e3c36dd639","observation_id":"5dc004e3-889e-47bb-aa86-7f110d4a3be6","resolution":{"observed_at":"2026-08-07T12:33:08.729239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:08.797773Z","title":"Salad-bench: A hierarchical and comprehensive safety benchmark for large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:08.797773Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:cc6fe6c8e733532e410d3f452b5216e9a92659424357d1dd8cad64d5c8a7600c","observation_id":"3042c996-c970-4899-9c34-7b2afc9e70d9","resolution":{"observed_at":"2026-08-07T12:33:08.797773Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:17.785846Z","title":"Long-tailed visual recognition via gaussian clouded logit adjustment","venue":null,"work_id":"2549a02f-27a7-4c29-8134-5c60307f0945","year":2022},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:08.934739Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:a11ab32404e9a8909d4c86cbc2d0f7fbbe987e0beaa806b470e95ce5857bf6d5","observation_id":"ebe72b6c-d1fb-434d-b5a8-7dd39b936bf4","resolution":{"observed_at":"2026-08-07T12:33:17.903767Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:17.559701Z","title":"Evaluating object hallucination in large vision-language models","venue":null,"work_id":"146320e4-019e-4bce-add9-7b029e26eb5b","year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.006807Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:2d66d2a2417e3670fedce692de591620d1a693f9d8f155f1fb89b1040d865486","observation_id":"4b24e925-a8de-49f7-ae0e-ae79231d2290","resolution":{"observed_at":"2026-08-07T12:33:17.659774Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:09.096630Z","title":"Omnibench: Towards the future of universal omni-language models.arXiv preprint arXiv:2409.15272, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.096630Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:3f39454cbe8c83449b670a7414500204a257e357647e4310df9df204a3dbd865","observation_id":"a46d00ea-67a9-4727-8ff8-6fe63c8f0dda","resolution":{"observed_at":"2026-08-07T12:33:09.096630Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.19437","last_updated":"2025-02-18T17:26:38Z","snapshot_observed_at":"2026-08-11T01:48:59.557045Z","submitted_at":"2024-12-27T04:03:16Z","title":"DeepSeek-V3 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.19437","snapshot_observed_at":"2026-08-07T12:33:09.169761Z","title":"Deepseek-v3 technical report.arXiv preprint arXiv:2412.19437, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.169761Z"},"links":{"cited_paper":"/paper/2412.19437","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:8c7b27230ed8269883fadd099e421e5aa7ee13529269f835cf8e38482b5dc1a7","observation_id":"ef6c1503-b847-4245-861e-fa84899b4c71","resolution":{"observed_at":"2026-08-07T12:33:09.169761Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:17.346063Z","title":"Mitigat- ing hallucination in large multi-modal models via robust instruction tuning","venue":null,"work_id":"27a04a3e-a596-4f14-910c-746754311e0a","year":null},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.244473Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:e0f1be5e744084d35f2d230940108a7a5e88a603806f5967c4b759b1aed81a1c","observation_id":"169dedb4-ccf8-4215-b20e-99bfc5d5e067","resolution":{"observed_at":"2026-08-07T12:33:17.441938Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:17.034810Z","title":"Visual instruction tuning, 2023","venue":null,"work_id":"6977991b-1598-4c03-b8ed-7cba7804aebb","year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.337797Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:84f3712199f89a101b73f57a37db71ef899b585fdaec123d0a5e6522729b6152","observation_id":"a549b860-d5f8-4579-88d9-6e0627e2fb8f","resolution":{"observed_at":"2026-08-07T12:33:17.186401Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.01785","last_updated":"2025-03-03T18:16:32Z","snapshot_observed_at":"2026-08-05T03:13:54.147007Z","submitted_at":"2025-03-03T18:16:32Z","title":"Visual-RFT: Visual Reinforcement Fine-Tuning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.01785","snapshot_observed_at":"2026-08-07T12:33:09.432965Z","title":"Visual-rft: Visual reinforcement fine-tuning.arXiv preprint arXiv:2503.01785, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.432965Z"},"links":{"cited_paper":"/paper/2503.01785","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:95b44e747c41a5b2760b8c0353b57bc43f8391a555f49d841b3c3eb54ed96dea","observation_id":"2551a763-37bd-41ce-931b-e4db63c3bf7d","resolution":{"observed_at":"2026-08-07T12:33:09.432965Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:09.543779Z","title":"Decoupled weight decay regularization","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.543779Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:7bfb8a511ba43294a9ffff9718f19cef8e2c440ad3630b3d66ae50a0b7b685db","observation_id":"7f6f5d91-2270-4ad3-be3f-e631d7dceb6c","resolution":{"observed_at":"2026-08-07T12:33:09.543779Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:09.643603Z","title":"Mathvista: Evaluating mathematical reasoning of foundation models in visual contexts","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.643603Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:e9ad8f5eff8eb1f63dab2727ef3e36b3b450e2cdc95475eb084bfdd3015b91b6","observation_id":"b7779aaa-102d-49ed-8517-7739ee88a52b","resolution":{"observed_at":"2026-08-07T12:33:09.643603Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:09.735339Z","title":"Learn to explain: Multimodal reasoning via thought chains for science question answering","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.735339Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:417dce84396d5f28af4cd3277ff689265d1cc7959281004d0d28bc93d325080f","observation_id":"50b2646c-45b8-47df-92bf-d05cb8e0947f","resolution":{"observed_at":"2026-08-07T12:33:09.735339Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:09.808767Z","title":"Ursa: Understanding and verifying chain-of-thought reasoning in multimodal mathematics.arXiv preprint arXiv:2501.04686, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.808767Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:8abb76db876a5bb69a2140a6ab3dd9b8024003de67215991667bb06869990710","observation_id":"0c0d71be-894d-4563-aa5a-2caaf14733de","resolution":{"observed_at":"2026-08-07T12:33:09.808767Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:09.913245Z","title":"Ok-vqa: A visual question answering benchmark requiring external knowledge","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.913245Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:c4ae30f89956b8953a6ba684c72095b10df69239988a69e643448768d3285492","observation_id":"63567212-d451-45be-8661-10868da2d31b","resolution":{"observed_at":"2026-08-07T12:33:09.913245Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:09.990090Z","title":"Chartqa: A benchmark for question answering about charts with visual and logical reasoning","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:09.990090Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:428687b5385866b60da5733a3759608f7774bab6c017f024bfd6ab2531924de6","observation_id":"d1fc6cd4-025a-416c-a8d9-0a324dbcea08","resolution":{"observed_at":"2026-08-07T12:33:09.990090Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.07365","last_updated":"2025-04-15T14:22:45Z","snapshot_observed_at":"2026-08-09T01:06:58.674891Z","submitted_at":"2025-03-10T14:23:12Z","title":"MM-Eureka: Exploring the Frontiers of Multimodal Reasoning with Rule-based Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.07365","snapshot_observed_at":"2026-08-07T12:33:10.071390Z","title":"Mm-eureka: Exploring visual aha moment with rule-based large-scale reinforcement learning.arXiv preprint arXiv:2503.07365, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:10.071390Z"},"links":{"cited_paper":"/paper/2503.07365","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:a7163f1dfd3a8a7cd2f8e7e32326baf52120c0c0241eae392c22b2844ad7491e","observation_id":"d12ecdd1-663f-4408-a3f2-485ff7baeab6","resolution":{"observed_at":"2026-08-07T12:33:10.071390Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:10.139052Z","title":"Compositional chain-of- thought prompting for large multimodal models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:10.139052Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:7e83a0784f4915630f879ae456159b8b25fd8b7be4f188cefa555b22c9c43823","observation_id":"4d684078-fce3-4628-b438-3871fe71aeab","resolution":{"observed_at":"2026-08-07T12:33:10.139052Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:10.217648Z","title":"Pytorch: An imperative style, high-performance deep learning library","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:10.217648Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:1dd86ad655ce897f65bea628430d500a19c4be4efd72849826ff4e46d7b89cf2","observation_id":"a53719f6-9c58-428f-b724-d963b96de839","resolution":{"observed_at":"2026-08-07T12:33:10.217648Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.07536","last_updated":"2025-03-11T03:32:59Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-10T17:04:14Z","title":"LMM-R1: Empowering 3B LMMs with Strong Reasoning Abilities Through Two-Stage Rule-Based RL","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.07536","snapshot_observed_at":"2026-08-07T12:33:10.310457Z","title":"Lmm-r1: Empowering 3b lmms with strong reasoning abilities through two-stage rule-based rl.arXiv preprint arXiv:2503.07536, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:10.310457Z"},"links":{"cited_paper":"/paper/2503.07536","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:ef73f2de55b6053c2be2f0cf4ac82e4fe16dccc76f892ea3d25e0ab5216fd337","observation_id":"04a56abd-815d-43aa-81bf-37bc82357b95","resolution":{"observed_at":"2026-08-07T12:33:10.310457Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.10703","last_updated":"2023-11-30T23:33:07Z","snapshot_observed_at":"2026-08-04T12:27:12.584527Z","submitted_at":"2023-04-21T02:19:06Z","title":"ReCEval: Evaluating Reasoning Chains via Correctness and Informativeness","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.10703","snapshot_observed_at":"2026-08-07T12:33:10.405095Z","title":"Receval: Evaluating reasoning chains via correctness and informativeness.arXiv preprint arXiv:2304.10703, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:10.405095Z"},"links":{"cited_paper":"/paper/2304.10703","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:6523733c1be5e1bc36d5958d88728b519ac9dc69624de47ff22c0b4b6f83e9f1","observation_id":"9e83b7d7-971b-4220-8f1f-5ce5d603d714","resolution":{"observed_at":"2026-08-07T12:33:10.405095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:16.662088Z","title":"Mitigating object hallucination in mllms via data-augmented phrase-level alignment","venue":null,"work_id":"72fdd5b8-5f81-4f2b-a9a9-e50e1cfa29ed","year":null},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:10.547565Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:1fda476ee46f632b660832de1e28802690be1912ad86bae2e5cb40a2a817aea9","observation_id":"698b7e53-57d5-4cb9-a2ee-331fa319c673","resolution":{"observed_at":"2026-08-07T12:33:16.787632Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-07T12:33:10.633888Z","title":"Proximal policy optimization algorithms.arXiv preprint arXiv:1707.06347, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:10.633888Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:06857306d76b8d0f255a3f5f4e0104e61672b80d78b034e88fce6134e522925c","observation_id":"284f205c-e24d-49d4-98ed-c7382e884264","resolution":{"observed_at":"2026-08-07T12:33:10.633888Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-07T12:33:10.745345Z","title":"Deepseekmath: Pushing the limits of mathematical reasoning in open language models.arXiv preprint arXiv:2402.03300, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:10.745345Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:81821ef82a8be10c40881ffeb803e5103839c5edeb08a35987b9d7dcd43e2ebf","observation_id":"a0bec2fa-ad09-47da-9b3d-dcfc443a20cb","resolution":{"observed_at":"2026-08-07T12:33:10.745345Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:10.856167Z","title":"Monte carlo tree search: A review of recent modifications and applications.Artificial Intelligence Review, 56(3):2497–2562, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:10.856167Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:effbdf9a361b7bde95e33906d37b8ea04dd8589d3b5560098b9f51d0cb5e447f","observation_id":"b2b2aefe-d4c4-4b50-9c4c-dffc83574f91","resolution":{"observed_at":"2026-08-07T12:33:10.856167Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-08-07T12:33:10.924496Z","title":"Gemini: a family of highly capable multimodal models.arXiv preprint arXiv:2312.11805, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:10.924496Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:f8049c4a9ed319a31d91597096d541d22d366f8d3c730437e2d677e1d0a03cb4","observation_id":"dcba0640-8e74-499b-936a-ba7fdccf62b7","resolution":{"observed_at":"2026-08-07T12:33:10.924496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:16.558078Z","title":"QVQ: To See the World with Wisdom","venue":null,"work_id":"3ebd8de2-774f-44af-97d8-c5bad45a67e0","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.048068Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:ff7a527db38fb8eb8e3227e71778646ed573ac4cea4ef71c75d8eff57b492a20","observation_id":"443738d3-a1a4-4e01-8fa7-10e641d3eecb","resolution":{"observed_at":"2026-08-07T12:33:16.606887Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.10960","last_updated":"2024-04-16T23:56:38Z","snapshot_observed_at":"2026-08-13T00:28:43.041015Z","submitted_at":"2024-04-16T23:56:38Z","title":"Uncertainty-Based Abstention in LLMs Improves Safety and Reduces Hallucinations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.10960","snapshot_observed_at":"2026-08-07T12:33:11.153366Z","title":"Uncertainty-based abstention in llms improves safety and reduces hallucinations.arXiv preprint arXiv:2404.10960, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.153366Z"},"links":{"cited_paper":"/paper/2404.10960","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:b70759eda7507c1de6780cb25a4457e638ead7829229ddd39e8ed485016bbce5","observation_id":"61e93b22-71a1-4bf0-957f-2c3634b8c1f2","resolution":{"observed_at":"2026-08-07T12:33:11.153366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:11.217618Z","title":"Eyes wide shut? exploring the visual shortcomings of multimodal llms","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.217618Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:b1d26a397e738d8b41309d3b2b02c4d4261ebcd1e0ca32e7f71d1f617567dd7e","observation_id":"8084a308-e4ac-4805-9d83-a724b7daae11","resolution":{"observed_at":"2026-08-07T12:33:11.217618Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:16.433777Z","title":"Measuring multimodal mathematical reasoning with math-vision dataset","venue":null,"work_id":"830a1c6e-4181-42a9-90b0-7d14b6a89282","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.281006Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:0dd5c12799d09144152bfb4cc003c9247e45b7ae82f33d69026126a03f540e6f","observation_id":"e08b3ec0-020d-434d-8a05-7a02209f9913","resolution":{"observed_at":"2026-08-07T12:33:16.486557Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12191","last_updated":"2024-10-03T15:54:49Z","snapshot_observed_at":"2026-08-06T05:35:29.109022Z","submitted_at":"2024-09-18T17:59:32Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12191","snapshot_observed_at":"2026-08-07T12:33:11.350592Z","title":"Qwen2-vl: Enhancing vision-language model’s perception of the world at any resolution.arXiv preprint arXiv:2409.12191, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.350592Z"},"links":{"cited_paper":"/paper/2409.12191","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:9e2079e3ee561e0dad774a7f1736e0d84f3a140f4000d6b5c2a9f8d978979c05","observation_id":"22da4bca-b3c8-4658-8779-f635c78267e6","resolution":{"observed_at":"2026-08-07T12:33:11.350592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:11.422440Z","title":"Chain-of-thought prompting elicits reasoning in large language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.422440Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:ecfe4117784d545225dae6748c1d264771ca0934c1604a293130d2c184494c38","observation_id":"1772c607-1e3e-4cc8-a011-797bd52d4277","resolution":{"observed_at":"2026-08-07T12:33:11.422440Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.17821","last_updated":"2024-12-16T10:27:35Z","snapshot_observed_at":"2026-08-12T23:56:32.542116Z","submitted_at":"2024-05-28T04:41:02Z","title":"RITUAL: Random Image Transformations as a Universal Anti-hallucination Lever in Large Vision Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.17821","snapshot_observed_at":"2026-08-07T12:33:11.492919Z","title":"Ritual: Ran- dom image transformations as a universal anti-hallucination lever in lvlms.arXiv preprint arXiv:2405.17821, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.492919Z"},"links":{"cited_paper":"/paper/2405.17821","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:7b806d243dd0d4e129f36acc29b9704b53a212d51fb6666ad6066dc09ea11145","observation_id":"edbb8b53-7b56-4357-9147-0fd3e8b09e71","resolution":{"observed_at":"2026-08-07T12:33:11.492919Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:16.236580Z","title":"Autohallusion: Automatic gen- eration of hallucination benchmarks for vision-language models","venue":null,"work_id":"e4da8994-d3af-44bf-bfcb-5506e12b4e00","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.562584Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:1f8dacb70344672f639b7047c54bd34fcad255b6c81ef8aefddf05c846e74285","observation_id":"630b0b44-60ff-4031-b5bd-136c5134da95","resolution":{"observed_at":"2026-08-07T12:33:16.317333Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:16.021405Z","title":"Grok 3 Beta — The Age of Reasoning Agents.https://x.ai/grok, 2025","venue":null,"work_id":"ca22c377-da75-4dee-b958-43d0d20698ab","year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.613404Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:836eee444408e2b7c37b849b2d7dcf53aa0d8da2a6d792b4118bbfdcc5edcccc","observation_id":"2767d4b3-7a1e-44f6-9233-16d9eaf5c5ab","resolution":{"observed_at":"2026-08-07T12:33:16.123781Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:15.892304Z","title":"Mitigating object hallucination via concentric causal attention.Advances in Neural Information Processing Systems, 37:92012–92035, 2024","venue":null,"work_id":"64674a95-2ce7-43ca-9dfa-6d551f3a10be","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.724385Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:abbb00afe67416013507f1ac818e656b95fe362174d58afec4879be73ae77044","observation_id":"9b1c93ed-31f5-4c5e-aef0-419a9196b660","resolution":{"observed_at":"2026-08-07T12:33:15.957279Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:15.753534Z","title":"Llava-cot: Let vision language models reason step-by-step, 2024","venue":null,"work_id":"650c66cd-bfcd-4e14-b193-a06e84aa41de","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.795732Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:66f27fc292cbb7b2bbed96e3c8fbf02b39c0570fde4e0686cad94adea1fc4feb","observation_id":"6ce3eff8-0eda-42fa-bd40-3d76fa9bab4f","resolution":{"observed_at":"2026-08-07T12:33:15.820797Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-08-07T12:33:11.894050Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.894050Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:db0203ebe36e498234ec8691290967bdad21238431315a9e46c6d6fc443fc470","observation_id":"eebbacd3-290d-40b6-9616-cec260b62609","resolution":{"observed_at":"2026-08-07T12:33:11.894050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:15.615351Z","title":"Soft-prompting with graph-of-thought for multi-modal representation learning","venue":null,"work_id":"7853f1e8-e2ac-4384-a2b2-a2381b7f7279","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:11.952400Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:181a583fba6bcd05eef116d9ca42737387ade50cfc2c067c969bf808becdfaef","observation_id":"98b48a0d-4c30-470a-b2d7-380e837c3767","resolution":{"observed_at":"2026-08-07T12:33:15.698313Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:15.512064Z","title":"Mitigating hallucination in large vision- language models via modular attribution and intervention","venue":null,"work_id":"f74ec223-eced-4598-b8dd-6642c33d82be","year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:12.053954Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:8e43d223398a99bf6d2673c2a004ca1b47461a708d3df6ee008d8ce6665edea0","observation_id":"9304223c-8c7d-4574-b644-7bb3046630c9","resolution":{"observed_at":"2026-08-07T12:33:15.543195Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10615","last_updated":"2025-03-18T08:52:34Z","snapshot_observed_at":"2026-08-07T12:50:27.666060Z","submitted_at":"2025-03-13T17:56:05Z","title":"R1-Onevision: Advancing Generalized Multimodal Reasoning through Cross-Modal Formalization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10615","snapshot_observed_at":"2026-08-07T12:33:12.177383Z","title":"R1-onevision: Advancing generalized multimodal reasoning through cross-modal formalization.arXiv preprint arXiv:2503.10615, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:12.177383Z"},"links":{"cited_paper":"/paper/2503.10615","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:915b38f7649f5a986f6f8a55bcfe12dd8ddbfa4a7672af6d109f29bb4d5ef541","observation_id":"73dcefd7-a958-4dd6-bded-72c5e627c82b","resolution":{"observed_at":"2026-08-07T12:33:12.177383Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.18319","last_updated":"2024-12-31T07:41:30Z","snapshot_observed_at":"2026-08-12T11:51:34.032343Z","submitted_at":"2024-12-24T10:07:51Z","title":"Mulberry: Empowering MLLM with o1-like Reasoning and Reflection via Collective Monte Carlo Tree Search","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.18319","snapshot_observed_at":"2026-08-07T12:33:12.258432Z","title":"Mulberry: Empowering mllm with o1-like reasoning and reflection via collective monte carlo tree search.arXiv preprint arXiv:2412.18319, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:12.258432Z"},"links":{"cited_paper":"/paper/2412.18319","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:a06183f9e47b82ef9a8b4c9b20c5a646f57702790b6b6b9a06f33f007324c647","observation_id":"0b48fbce-4c79-415a-aeac-4e3e872724ae","resolution":{"observed_at":"2026-08-07T12:33:12.258432Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:12.383675Z","title":"React: Synergizing reasoning and acting in language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:12.383675Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:e44fb28b5925c9172b0eb7ae3d308d30578af098f4e149893ec1bf6a4e608de3","observation_id":"c618bb1f-9378-4f4e-a8e8-05ca272c661b","resolution":{"observed_at":"2026-08-07T12:33:12.383675Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.03387","last_updated":"2025-07-29T16:23:02Z","snapshot_observed_at":"2026-08-08T04:13:22.884923Z","submitted_at":"2025-02-05T17:23:45Z","title":"LIMO: Less is More for Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.03387","snapshot_observed_at":"2026-08-07T12:33:12.448500Z","title":"Limo: Less is more for reasoning.arXiv preprint arXiv:2502.03387, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:12.448500Z"},"links":{"cited_paper":"/paper/2502.03387","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:1eb6cd7d09555b3b8e230750e88f8a5bc68dd8e6fc57e2f4dda4bbc8ce319327","observation_id":"ed50d571-49f4-4efd-a3c0-5fc29315c7a7","resolution":{"observed_at":"2026-08-07T12:33:12.448500Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:12.557605Z","title":"Hallucidoctor: Mitigating hallucinatory toxicity in visual instruction data","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:12.557605Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:2937072ee948ec7a69a72902bb85bcce17d4a2062a153cf5ca6ca32566b7c7b3","observation_id":"3cb4a8e8-4452-4317-ac9e-074e4f911ef1","resolution":{"observed_at":"2026-08-07T12:33:12.557605Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.14476","last_updated":"2025-05-20T01:37:34Z","snapshot_observed_at":"2026-08-02T01:40:54.187278Z","submitted_at":"2025-03-18T17:49:06Z","title":"DAPO: An Open-Source LLM Reinforcement Learning System at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.14476","snapshot_observed_at":"2026-08-07T12:33:12.653152Z","title":"Dapo: An open-source llm reinforcement learning system at scale.arXiv preprint arXiv:2503.14476, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:12.653152Z"},"links":{"cited_paper":"/paper/2503.14476","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:d989f0660e9ee2f4f5e52fc9494a97730059cfb15e66af6be6b36825c5dc1e66","observation_id":"9d8616b4-cc9f-478c-83fe-fa1433f61ea9","resolution":{"observed_at":"2026-08-07T12:33:12.653152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:12.750492Z","title":"Rlhf-v: Towards trustworthy mllms via behavior alignment from fine-grained correctional human feedback","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:12.750492Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:8475f6644031e0547e31660e1ca8785dd08942d25b3b75098408014a24047bc4","observation_id":"0f2fb295-f912-4fbd-8059-15d7e34eb919","resolution":{"observed_at":"2026-08-07T12:33:12.750492Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02884","last_updated":"2024-11-21T07:07:59Z","snapshot_observed_at":"2026-08-12T22:31:19.889900Z","submitted_at":"2024-10-03T18:12:29Z","title":"LLaMA-Berry: Pairwise Optimization for O1-like Olympiad-Level Mathematical Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.02884","snapshot_observed_at":"2026-08-07T12:33:12.857251Z","title":"Llama-berry: Pairwise optimization for o1-like olympiad-level mathematical reasoning.arXiv preprint arXiv:2410.02884, 2024","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:12.857251Z"},"links":{"cited_paper":"/paper/2410.02884","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:c039af7f568ddfdb186d09127b05b3eb1f0a37c81afe8d29df4daadfe6ef68e9","observation_id":"9ebaadc9-87c8-41a0-ba69-f23f61aa0c81","resolution":{"observed_at":"2026-08-07T12:33:12.857251Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:15.318574Z","title":"Reflective instruction tuning: Mitigating hallucinations in large vision-language models","venue":null,"work_id":"74323fb2-9c31-4f60-9fce-18cb4852c6a9","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:12.980715Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:a115fd45edbda46f53e4370471e5923f188a57870fdadd09030c7bff5e003dbe","observation_id":"5e0dc656-f057-48a7-a72b-70435f63512d","resolution":{"observed_at":"2026-08-07T12:33:15.380205Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:13.073898Z","title":"Mathverse: Does your multi-modal llm truly see the diagrams in visual math problems? InEuropean Conference on Computer Vision, pages 169–186","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:13.073898Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:d342c5867fa42888a2219af9ae9253f04a4b005e31994b7e72b79ba1097c1962","observation_id":"5a298254-9648-4782-abb3-752e69622b8f","resolution":{"observed_at":"2026-08-07T12:33:13.073898Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.11919","last_updated":"2024-11-28T13:35:56Z","snapshot_observed_at":"2026-08-12T18:41:17.892508Z","submitted_at":"2024-11-18T04:06:04Z","title":"VL-Uncertainty: Detecting Hallucination in Large Vision-Language Model via Uncertainty Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.11919","snapshot_observed_at":"2026-08-07T12:33:13.190579Z","title":"Vl-uncertainty: Detecting hallucination in large vision-language model via uncertainty estimation.arXiv preprint arXiv:2411.11919, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:13.190579Z"},"links":{"cited_paper":"/paper/2411.11919","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:f55c86dfec7a6956e36f0bfb40ed52eddfd25c87eee46e90ed33a8d2b980f70e","observation_id":"140db751-7cb8-4809-9e30-2e832334feb1","resolution":{"observed_at":"2026-08-07T12:33:13.190579Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07301","last_updated":"2025-06-05T16:34:24Z","snapshot_observed_at":"2026-08-03T11:11:25.359494Z","submitted_at":"2025-01-13T13:10:16Z","title":"The Lessons of Developing Process Reward Models in Mathematical Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07301","snapshot_observed_at":"2026-08-07T12:33:13.281932Z","title":"The lessons of developing process reward models in mathematical reasoning.arXiv preprint arXiv:2501.07301, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:13.281932Z"},"links":{"cited_paper":"/paper/2501.07301","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:0db9fd4fec977bb7756f946984c6af8fb3cc8f749a3e14aa1a791ddf513a7003","observation_id":"334b295f-408d-44bd-8026-aa20d981a015","resolution":{"observed_at":"2026-08-07T12:33:13.281932Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:13.378434Z","title":"Multimodal chain-of-thought reasoning in language models.Transactions on Machine Learning Research","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:13.378434Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:9c4705c14760d16d26e67eb289e63e6929a0a8e91a681af845b242bace35e52a","observation_id":"5fd38f61-18d5-478f-b822-1218dce4fd63","resolution":{"observed_at":"2026-08-07T12:33:13.378434Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.16839","last_updated":"2024-02-06T16:43:31Z","snapshot_observed_at":"2026-08-08T04:17:23.797697Z","submitted_at":"2023-11-28T14:54:37Z","title":"Beyond Hallucinations: Enhancing LVLMs through Hallucination-Aware Direct Preference Optimization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.16839","snapshot_observed_at":"2026-08-07T12:33:13.470889Z","title":"Beyond hallucinations: Enhancing lvlms through hallucination-aware direct preference optimization","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:13.470889Z"},"links":{"cited_paper":"/paper/2311.16839","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:ca20950eeb47a523a4c1802007991ea3e4c359895827e4fefd162414cd3f14cd","observation_id":"daedcae0-b3a7-437e-b406-b2b3c87dc327","resolution":{"observed_at":"2026-08-07T12:33:13.470889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:15.136333Z","title":"A picture is worth a graph: A blueprint debate paradigm for multimodal reasoning","venue":null,"work_id":"8d2a7e63-2ecb-41ba-a1d9-0413839af48f","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:13.563966Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:0e14d99a8ef9711f9e51beb9245f7adc178e865301808e488af1da2402683b0b","observation_id":"cfacc5c1-dfc2-4639-8b13-4af198dd8327","resolution":{"observed_at":"2026-08-07T12:33:15.196268Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.12591","last_updated":"2024-11-15T21:01:37Z","snapshot_observed_at":"2026-08-12T19:30:30.766081Z","submitted_at":"2024-11-15T21:01:37Z","title":"Thinking Before Looking: Improving Multimodal LLM Reasoning via Mitigating Visual Hallucination","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.12591","snapshot_observed_at":"2026-08-07T12:33:13.663322Z","title":"Thinking before looking: Improving multimodal llm reasoning via mitigating visual hallucination.arXiv preprint arXiv:2411.12591, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:13.663322Z"},"links":{"cited_paper":"/paper/2411.12591","citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:943703faf8974561ca52bea125f411505978e863c04ba01e797fd92e0b634b88","observation_id":"7f35f66e-4cbe-463e-a511-e8f97f085ad1","resolution":{"observed_at":"2026-08-07T12:33:13.663322Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:14.975232Z","title":"Relying on the unreliable: The impact of language models’ reluctance to express uncertainty","venue":null,"work_id":"67b086cf-5d1d-4470-a1d2-b1ea5100d000","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:13.747662Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:6d53a589fa39d09aaa3c16cbdf7c73485e264d349c88b7d94a7d8f9858c7bd66","observation_id":"0db258ac-a6f2-4765-b849-f5e6ecfc0fd7","resolution":{"observed_at":"2026-08-07T12:33:15.038180Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:33:14.828375Z","title":"the answer is [answer in the input]","venue":null,"work_id":"8972c6af-36e2-4cac-9bb5-d82166ec43a4","year":2024},"citing_paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM","version":2},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-07T12:33:13.820971Z"},"links":{"citing_paper":"/paper/2505.24238"},"observation_digest":"sha256:34059d918180104963f12d3138a880cc259523b6edb01b12b886b01a1ddd6c63","observation_id":"0d179a1e-f3fd-41b8-a842-23e070afc6e7","resolution":{"observed_at":"2026-08-07T12:33:14.881658Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.24238","last_updated":"2025-06-02T04:16:04Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-10T15:39:33.509863Z","submitted_at":"2025-05-30T05:54:36Z","title":"MIRAGE: Assessing Hallucination in Multimodal Reasoning Chains of MLLM"},"reference_resolution":{"displayed":87,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":67,"verified_exact":1,"verified_fuzzy":18},"total_outbound_references":87},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 87 of 87 outbound references and 4 inbound Pith citation observations for arXiv:2505.24238."}