{"as_of":"2026-08-10T09:58:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2b6b9ec889fdda2ed475ba832e179d747321cb188ab3d3776c0e3ba584e39167","coverage":[{"denominator":86,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":86,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T12:35:46.811009Z","state":"measured"},{"denominator":97,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":97,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":11,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":11,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T00:24:14.990061Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T02:56:30.004384Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":"2505.24863","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-07-02T02:56:30.004384Z","title":"arXiv preprint arXiv:2505.24863 , year=","venue":null,"work_id":"d82f4102-c19f-46d0-bece-b8c1eeb68013","year":2024},"citing_paper":{"arxiv_id":"2503.16419","last_updated":"2025-08-21T19:14:40Z","snapshot_observed_at":"2026-08-07T04:27:23.738927Z","submitted_at":"2025-03-20T17:59:38Z","title":"Stop Overthinking: A Survey on Efficient Reasoning for Large Language Models","version":4},"reference_index":233,"source":"pdf_text","source_observed_at":"2026-05-14T01:29:56.480020Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2503.16419"},"observation_digest":"sha256:7b767738c7a36cea18f06c52bcc548acd097c2c466dc5a172ff69880f0543abc","observation_id":"bcfad6db-13f4-4948-aeeb-73cc2887acc1","resolution":{"observed_at":"2026-05-14T01:29:57.372131Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-08-07T00:24:14.990061Z","title":"Alphaone: Reasoning models thinking slow and fast at test time","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.14234","last_updated":"2025-06-17T06:47:19Z","snapshot_observed_at":"2026-08-10T00:27:45.764398Z","submitted_at":"2025-06-17T06:47:19Z","title":"Xolver: Multi-Agent Reasoning with Holistic Experience Learning Just Like an Olympiad Team","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-07T00:24:14.990061Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2506.14234"},"observation_digest":"sha256:c95d8e299607b6b09ed14f538970d0d3663ddbfdec8e318d1918dde8b230e8e5","observation_id":"875c14e5-9da1-4761-810f-29e2a32e649b","resolution":{"observed_at":"2026-08-07T00:24:14.990061Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-08-06T18:59:43.011814Z","title":"Alphaone: Reasoning models thinking slow and fast at test time","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.06829","last_updated":"2025-07-09T13:28:35Z","snapshot_observed_at":"2026-08-07T18:20:33.639424Z","submitted_at":"2025-07-09T13:28:35Z","title":"Adaptive Termination for Multi-round Parallel Reasoning: An Universal Semantic Entropy-Guided Framework","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-06T18:59:43.011814Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2507.06829"},"observation_digest":"sha256:ea0ed57f999257583715aa24de00767908a2a45a36c21f6cb9609178df326d82","observation_id":"1e6f008f-b77d-4be7-8773-21f03ad50616","resolution":{"observed_at":"2026-08-06T18:59:43.011814Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-08-06T17:54:18.072944Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.09662","last_updated":"2025-07-13T14:51:59Z","snapshot_observed_at":"2026-08-07T01:15:50.475193Z","submitted_at":"2025-07-13T14:51:59Z","title":"Towards Concise and Adaptive Thinking in Large Reasoning Models: A Survey","version":1},"reference_index":241,"source":"arxiv_source","source_observed_at":"2026-08-06T17:54:18.072944Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2507.09662"},"observation_digest":"sha256:fc05a34b5aefd4c30e2005664019fe836d96df81e3168e99f58d0dd80d264dbd","observation_id":"464cfba2-3cac-4ea1-a149-16e04da590c4","resolution":{"observed_at":"2026-08-06T17:54:18.072944Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-08-04T16:07:46.801109Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.16679","last_updated":"2025-09-20T13:11:28Z","snapshot_observed_at":"2026-08-04T16:07:24.699834Z","submitted_at":"2025-09-20T13:11:28Z","title":"Reinforcement Learning Meets Large Language Models: A Survey of Advancements and Applications Across the LLM Lifecycle","version":1},"reference_index":221,"source":"pdf_text","source_observed_at":"2026-08-04T16:07:46.801109Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2509.16679"},"observation_digest":"sha256:41259b7507063f4482ad77694583f65c7e64624f7805113dae087e07440c5447","observation_id":"3edafc1f-a652-4ff2-90c9-d081221d1493","resolution":{"observed_at":"2026-08-04T16:07:46.801109Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":"2505.24863","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-07-02T02:56:30.004384Z","title":"arXiv preprint arXiv:2505.24863 , year=","venue":null,"work_id":"d82f4102-c19f-46d0-bece-b8c1eeb68013","year":2024},"citing_paper":{"arxiv_id":"2509.21743","last_updated":"2026-03-31T22:32:08Z","snapshot_observed_at":"2026-08-10T01:07:51.498001Z","submitted_at":"2025-09-26T01:17:35Z","title":"Retrieval-of-Thought: Efficient Reasoning via Reusing Thoughts","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-18T13:42:07.883909Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2509.21743"},"observation_digest":"sha256:6f1e856b1bd5e971129b5764f1c70b91fc96f0ea31f8706b770d10aa9ed3adcf","observation_id":"04a6adae-24d6-484e-821a-7517b539e69a","resolution":{"observed_at":"2026-05-18T13:42:38.577012Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":"2505.24863","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-07-02T02:56:30.004384Z","title":"arXiv preprint arXiv:2505.24863 , year=","venue":null,"work_id":"d82f4102-c19f-46d0-bece-b8c1eeb68013","year":2024},"citing_paper":{"arxiv_id":"2510.19669","last_updated":"2026-05-08T14:25:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-10-22T15:16:06Z","title":"DiffAdapt: Difficulty-Adaptive Reasoning for Token-Efficient LLM Inference","version":5},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-18T04:31:40.360872Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2510.19669"},"observation_digest":"sha256:176736c681568360b526e85c18b31b6432749f0ceb047e704abea519b21479c2","observation_id":"09caae37-e7bf-4ff7-b14d-d3b50eff72d5","resolution":{"observed_at":"2026-05-18T04:32:23.065413Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":"2505.24863","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-07-02T02:56:30.004384Z","title":"arXiv preprint arXiv:2505.24863 , year=","venue":null,"work_id":"d82f4102-c19f-46d0-bece-b8c1eeb68013","year":2024},"citing_paper":{"arxiv_id":"2605.06165","last_updated":"2026-05-07T12:51:49Z","snapshot_observed_at":"2026-07-06T23:18:41.400741Z","submitted_at":"2026-05-07T12:51:49Z","title":"Post Reasoning: Improving the Performance of Non-Thinking Models at No Cost","version":1},"reference_index":226,"source":"arxiv_source","source_observed_at":"2026-05-08T10:19:08.451445Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2605.06165"},"observation_digest":"sha256:679d5c2743c924e8691a218e68ebfb3d18f3edcfaaf19a47ae817ea9832adb7a","observation_id":"f759f154-7d45-47c3-9da3-3401b56d78b2","resolution":{"observed_at":"2026-05-11T20:06:09.141579Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":"2505.24863","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-07-02T02:56:30.004384Z","title":"arXiv preprint arXiv:2505.24863 , year=","venue":null,"work_id":"d82f4102-c19f-46d0-bece-b8c1eeb68013","year":2024},"citing_paper":{"arxiv_id":"2605.15205","last_updated":"2026-04-28T15:38:31Z","snapshot_observed_at":"2026-08-08T12:32:02.265510Z","submitted_at":"2026-04-28T15:38:31Z","title":"Does Theory of Mind Improvement Really Benefit Human-AI Interactions? Empirical Findings from Interactive Evaluations","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-05-19T17:55:04.458343Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2605.15205"},"observation_digest":"sha256:21d03432a8a87ecf434d9a09ba374aa730a375bedb617af7226fb5646f1699e4","observation_id":"16629ace-445e-46bd-828e-972456121a75","resolution":{"observed_at":"2026-05-19T17:57:42.615891Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":"2505.24863","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-07-02T02:56:30.004384Z","title":"arXiv preprint arXiv:2505.24863 , year=","venue":null,"work_id":"d82f4102-c19f-46d0-bece-b8c1eeb68013","year":2024},"citing_paper":{"arxiv_id":"2606.03102","last_updated":"2026-06-02T03:42:04Z","snapshot_observed_at":"2026-08-03T04:21:57.609394Z","submitted_at":"2026-06-02T03:42:04Z","title":"Small RL Controller, Large Language Model: RL-Guided Adaptive Sampling for Test-Time Scaling","version":1},"reference_index":106,"source":"arxiv_source","source_observed_at":"2026-06-28T10:25:10.559953Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2606.03102"},"observation_digest":"sha256:337101f8507c53ba4c913626e247d9ce4a183cd480271a41b4c6c372f23ef83c","observation_id":"a1204c21-632d-414a-bba3-c360f13b7e96","resolution":{"observed_at":"2026-07-02T02:56:30.005998Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.24863","snapshot_observed_at":"2026-08-05T13:40:59.007507Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.03745","last_updated":"2026-08-04T14:38:06Z","snapshot_observed_at":"2026-08-08T09:35:52.621563Z","submitted_at":"2026-08-04T14:38:06Z","title":"Risky Business: Measuring The Faithfulness-Safety Tension","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-05T13:40:59.007507Z"},"links":{"cited_paper":"/paper/2505.24863","citing_paper":"/paper/2608.03745"},"observation_digest":"sha256:196ec4c77483c3c84d3be7fada30d57eaf1e3d8f3bd9c869bf52be7c2b1b0fef","observation_id":"e218206b-38d0-4b11-a59c-d26d7dd57f99","resolution":{"observed_at":"2026-08-05T13:40:59.007507Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2505.24863/citation-record","integrity":"/paper/2505.24863/integrity","json":"/paper/2505.24863/citation-record.json","paper":"/paper/2505.24863"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:38.415046Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:38.415046Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:0a344ddd48457bd5d3f87f4df1c0759b2d1c94038a0c315f1a26c058d5f3b148","observation_id":"446c5930-59d5-4b7f-9dde-966703889210","resolution":{"observed_at":"2026-08-07T12:35:38.415046Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:38.472967Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:38.472967Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:44ae434d5a3b1195da9b3489f94658f45e0accec338727e4279805d794c2371c","observation_id":"955a208e-5e40-43fb-bb03-504f22f3ea19","resolution":{"observed_at":"2026-08-07T12:35:38.472967Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.08905","last_updated":"2024-12-12T03:37:41Z","snapshot_observed_at":"2026-08-05T04:04:21.846023Z","submitted_at":"2024-12-12T03:37:41Z","title":"Phi-4 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.08905","snapshot_observed_at":"2026-08-07T12:35:38.633160Z","title":"Hewett, Mojan Javaheripi, Piero Kauffmann, James R","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:38.633160Z"},"links":{"cited_paper":"/paper/2412.08905","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:956e96dafb2159f0dcb4c16ef42437bbfb236f6736fb58d4ed7c82f29cd28cf4","observation_id":"52c84fcd-d8ca-496d-a2bc-f45cb725d69a","resolution":{"observed_at":"2026-08-07T12:35:38.633160Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:38.724420Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:38.724420Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:824a72430937fa1970059c95023b304b5af24f6b267a398650500a6f547bae94","observation_id":"655a5abd-529f-4c07-ac47-a30223ab8458","resolution":{"observed_at":"2026-08-07T12:35:38.724420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:38.789544Z","title":"Menick, Sebastian Borgeaud, and 8 others","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:38.789544Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:15e3f608934067e815bd9d79c5135eec62f25b71aa002940675e58c0970ead9b","observation_id":"c8eed412-c836-4f0d-8718-ddadf2a0a29a","resolution":{"observed_at":"2026-08-07T12:35:38.789544Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:38.906640Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:38.906640Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:19e5eef785b9f10ec0d1e97ae5f26eb2126157017d2a0af08e8eebcbc99d182d","observation_id":"56575740-3d15-4b18-bf46-548a760da44b","resolution":{"observed_at":"2026-08-07T12:35:38.906640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:38.981718Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:38.981718Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:9ed691d3c461c64d93813107c796b41959aaa2946b44af8f7efa1b3c6dde9437","observation_id":"b639731b-a3e7-426f-b108-4524486d4707","resolution":{"observed_at":"2026-08-07T12:35:38.981718Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2108.07258","last_updated":"2022-07-12T23:45:14Z","snapshot_observed_at":"2026-08-02T09:20:40.804790Z","submitted_at":"2021-08-16T17:50:08Z","title":"On the Opportunities and Risks of Foundation Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2108.07258","snapshot_observed_at":"2026-08-07T12:35:39.053745Z","title":"Hudson, Ehsan Adeli, Russ Altman, Simran Arora, Sydney von Arx, Michael S","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.053745Z"},"links":{"cited_paper":"/paper/2108.07258","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:3419634d84190259bb0b48d620f32729b320841218f062b21fdd61511e21739d","observation_id":"9c3019d7-2402-4a41-a077-b3459d8220bf","resolution":{"observed_at":"2026-08-07T12:35:39.053745Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:39.125626Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.125626Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:d7e20bc3818b7bf6d86d2ae14aaacb036f2d72d58287b59371762233fdd26f00","observation_id":"bf46384b-e2e6-4279-a681-7d571b36b7ce","resolution":{"observed_at":"2026-08-07T12:35:39.125626Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.21187","last_updated":"2025-02-01T07:57:37Z","snapshot_observed_at":"2026-08-01T16:43:44.704797Z","submitted_at":"2024-12-30T18:55:12Z","title":"Do NOT Think That Much for 2+3=? On the Overthinking of o1-Like LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.21187","snapshot_observed_at":"2026-08-07T12:35:39.203261Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.203261Z"},"links":{"cited_paper":"/paper/2412.21187","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:268ad22784df63ac9943040d6f0ffd888f8a382f103a7d5b5effe6b491801934","observation_id":"1bbfc691-591a-44ce-bf1b-149f5c43c65d","resolution":{"observed_at":"2026-08-07T12:35:39.203261Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:55.404727Z","title":null,"venue":null,"work_id":"e1a775e7-2080-452c-b25f-a4090bebb225","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.281262Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:33cc9a0890b6d5b1053152349f4f40aa5363e64ac5104a3c3726ef5d7389987a","observation_id":"c524ac07-3d7f-44e9-9400-b24bd86b0d3d","resolution":{"observed_at":"2026-08-07T12:35:55.467565Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:55.134813Z","title":null,"venue":null,"work_id":"e6bf7fcb-99c6-45df-86f9-35f1a8c975a3","year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.360251Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:a0a43f47a754e61949cc6ab5fba2b65d3b2d9386017e237b3eb43b5b4c9f64de","observation_id":"0ce27cc8-1a8c-47fc-8ffe-a85bb581e82f","resolution":{"observed_at":"2026-08-07T12:35:55.286596Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:54.804199Z","title":"Christiano, Jan Leike, Tom B","venue":null,"work_id":"71e1fe5a-4163-40b4-9e49-96101225cc49","year":2017},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.458997Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:b3ec7eec49ba948932105246e1617e02abcabb6e1990d61474f5fbd88925905e","observation_id":"0b6ce5dc-b7fa-44ef-b1ae-454286563397","resolution":{"observed_at":"2026-08-07T12:35:54.947424Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.14168","last_updated":"2021-11-18T00:23:45Z","snapshot_observed_at":"2026-08-07T01:45:38.840969Z","submitted_at":"2021-10-27T04:49:45Z","title":"Training Verifiers to Solve Math Word Problems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.14168","snapshot_observed_at":"2026-08-07T12:35:39.537836Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.537836Z"},"links":{"cited_paper":"/paper/2110.14168","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:109d91a7a5d726cb4540a091b840b5a74b675630fa98ec94062fd993c6e3e90c","observation_id":"072608b7-4c26-4ce1-9bc0-5887153eff0d","resolution":{"observed_at":"2026-08-07T12:35:39.537836Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-07T12:35:39.653397Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.653397Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:f9642577fe57984703bf5f1c7b6e4400c4caffd69d684c7d4a3da3f359110ff6","observation_id":"db5b4ff2-c681-4c83-a20e-a618104e51eb","resolution":{"observed_at":"2026-08-07T12:35:39.653397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:54.625348Z","title":null,"venue":null,"work_id":"2a4df57a-afa6-4685-bf74-c9e1ad85ef1c","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.721445Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:0498412d82d0b6a566bce2ca667c85535ab26e0dfd92a3a818fdc934fee62a67","observation_id":"bf9f418e-d37a-4642-a3fd-9bf6bb304f6f","resolution":{"observed_at":"2026-08-07T12:35:54.641437Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:54.266975Z","title":null,"venue":null,"work_id":"5f46ec6a-b9ae-4af0-bbb0-bc91ac402177","year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.849173Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:87000fc5037de12c045e0f54b775674f7548d681ba9d596ed4a212fd7d816cb5","observation_id":"8ce38e14-673f-4f13-9526-eb914063ecf5","resolution":{"observed_at":"2026-08-07T12:35:54.417872Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:53.978249Z","title":null,"venue":null,"work_id":"491accbb-011b-4719-8c86-df7ab04e2a69","year":2023},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:39.923165Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:c90fa7712326d2e07ed90b7914e5a119c62796c20531fca8154dc48fd4711817","observation_id":"4d77c6b2-6942-4d59-bc97-4d7378c908d8","resolution":{"observed_at":"2026-08-07T12:35:54.146058Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:53.670364Z","title":null,"venue":null,"work_id":"50b0fc8c-20dd-4d29-ac84-cb3076320808","year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:40.031171Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:f80c97619a10cf9cb45ecf08a68e02a3c29e1f545431ed8ec49107266116c4b3","observation_id":"062ec658-7f61-42dc-9676-2f04ef7663be","resolution":{"observed_at":"2026-08-07T12:35:53.779964Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.01707","last_updated":"2024-12-25T13:32:54Z","snapshot_observed_at":"2026-07-06T19:26:22.646942Z","submitted_at":"2024-10-02T16:15:31Z","title":"Interpretable Contrastive Monte Carlo Tree Search Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.01707","snapshot_observed_at":"2026-08-07T12:35:40.169069Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:40.169069Z"},"links":{"cited_paper":"/paper/2410.01707","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:f82aac0babf36c80115ca51d000abdd9a8e5afb3c5992c299b9bcee5f7abd78a","observation_id":"defa0594-d3b8-4891-94e8-9aabfcea5a7d","resolution":{"observed_at":"2026-08-07T12:35:40.169069Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:53.452204Z","title":null,"venue":null,"work_id":"ae010fd0-c49a-4038-a838-61a72b709ff0","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:40.358111Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:57b6929f8406e4bc0536dad40cd0e15385cfc79bc61e14d2d1ace25447946784","observation_id":"570cfa71-248c-4369-80e1-71693283e9c9","resolution":{"observed_at":"2026-08-07T12:35:53.549630Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:53.173157Z","title":null,"venue":null,"work_id":"2afa62da-3762-4402-91a1-215751a9a308","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:40.498792Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:60f1266298ebed3c7dd7769fab51a51af795b78a2ac6932f9392303f86b305b1","observation_id":"bbca503e-5214-4914-8521-7b4e2d86dc75","resolution":{"observed_at":"2026-08-07T12:35:53.288057Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:40.642261Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:40.642261Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:119caa03fc4e1d26e8c73cda9bcfa2f730c0df620b29b68e1caa13199a275d46","observation_id":"98a66aef-624b-475a-a3d1-7224008881dd","resolution":{"observed_at":"2026-08-07T12:35:40.642261Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.06769","last_updated":"2025-11-03T00:53:34Z","snapshot_observed_at":"2026-08-07T06:05:27.895209Z","submitted_at":"2024-12-09T18:55:56Z","title":"Training Large Language Models to Reason in a Continuous Latent Space","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.06769","snapshot_observed_at":"2026-08-07T12:35:40.840693Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:40.840693Z"},"links":{"cited_paper":"/paper/2412.06769","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:ca00a6c601475194b627fb980d9a3dea4f39be2e881c8ba978847bbdaa6043e3","observation_id":"2a602427-8d05-44f4-aa0e-8ea0aaf2eb7f","resolution":{"observed_at":"2026-08-07T12:35:40.840693Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:41.024029Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:41.024029Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:06d76c6af9a4496069b22ff20f9885603b826e36a25445b98f756636535e1c3a","observation_id":"315f201c-c8a8-4e7c-8e56-15da14f362a1","resolution":{"observed_at":"2026-08-07T12:35:41.024029Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:52.978800Z","title":null,"venue":null,"work_id":"a4fb55b3-ad6e-4777-9619-216057a585f3","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:41.159285Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:cef722ed06d25516e066e192e2cc26df6bf1c1942e70ec2b71b459f53c8852e2","observation_id":"10ac0aee-2c18-4190-8544-1bf329c82465","resolution":{"observed_at":"2026-08-07T12:35:53.042872Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.16489","last_updated":"2024-11-25T15:31:27Z","snapshot_observed_at":"2026-07-06T19:56:39.115630Z","submitted_at":"2024-11-25T15:31:27Z","title":"O1 Replication Journey -- Part 2: Surpassing O1-preview through Simple Distillation, Big Progress or Bitter Lesson?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.16489","snapshot_observed_at":"2026-08-07T12:35:41.259399Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:41.259399Z"},"links":{"cited_paper":"/paper/2411.16489","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:d52bfa43cf683d72ca2dda0c7a92ddf3f0937d6b82eb14d6188797fa550a6a30","observation_id":"8ce8cbb9-e0b3-4779-90bd-79f1d93cbdd2","resolution":{"observed_at":"2026-08-07T12:35:41.259399Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16720","last_updated":"2026-04-30T02:46:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-21T18:04:31Z","title":"OpenAI o1 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16720","snapshot_observed_at":"2026-08-07T12:35:41.369633Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:41.369633Z"},"links":{"cited_paper":"/paper/2412.16720","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:d4751b769a797bbf3ac9c51e13c055c32b1c1946f63f515387cdad5ee1072ce3","observation_id":"3d05acff-7719-45f9-af41-66427d7c4438","resolution":{"observed_at":"2026-08-07T12:35:41.369633Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:41.519698Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:41.519698Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:25ec8cba60ce93b31c5bb8be4c297f7d9ddda6edc76caa6fb6afe7abe25dde2e","observation_id":"f33e9d89-d918-4083-ac5c-037526f114d3","resolution":{"observed_at":"2026-08-07T12:35:41.519698Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.00703","last_updated":"2025-07-01T11:46:57Z","snapshot_observed_at":"2026-08-07T20:35:00.460457Z","submitted_at":"2025-05-01T17:59:46Z","title":"T2I-R1: Reinforcing Image Generation with Collaborative Semantic-level and Token-level CoT","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.00703","snapshot_observed_at":"2026-08-07T12:35:41.646136Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:41.646136Z"},"links":{"cited_paper":"/paper/2505.00703","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:fb8647d8ae1446cfe6bea9c29b71edee75b397bb9435a8a2f10dd2b65c0ac6fd","observation_id":"f3bd3219-bc58-467c-8528-8baebc2508fa","resolution":{"observed_at":"2026-08-07T12:35:41.646136Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:52.800192Z","title":null,"venue":null,"work_id":"0566580a-c3c9-4bca-8ccc-adec78d6c945","year":2011},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:41.735649Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:b041aa1562b5370635e3593237e3ec6579c53ba67fbb74c73b4406c718acba48","observation_id":"3c84f132-5401-4c90-920b-ae78a630eaa4","resolution":{"observed_at":"2026-08-07T12:35:52.860572Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:41.841093Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:41.841093Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:b940ef1fb534631406a798876708517ec410bcefd583b23ebf63b7c206c3efa4","observation_id":"35e2c96a-be73-4291-bc46-c0c37800fe74","resolution":{"observed_at":"2026-08-07T12:35:41.841093Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:41.962370Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:41.962370Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:d1a6d9687924f8e39183d322e6a8af96f1f24856362f9eac173680eda4db14d5","observation_id":"6a527d3c-60ea-4423-8698-3b98b22c7730","resolution":{"observed_at":"2026-08-07T12:35:41.962370Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:52.630474Z","title":null,"venue":null,"work_id":"763af7f6-81a2-4a45-a735-6b14c892c55a","year":2019},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:42.057317Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:ab78ff4d81d9f43fc8a8a7cb4315fa5c6fc3f3abc04ffb4d1b3527d073c0fcf1","observation_id":"ac58c141-10d3-4d77-95ad-f2e6a1592408","resolution":{"observed_at":"2026-08-07T12:35:52.706324Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:42.145814Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:42.145814Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:b187897efca664cdec1475192bf4126c6808745708a6192a9418c8ffbc5fcb7e","observation_id":"795e7652-2f7f-4abb-b3c8-7831448700be","resolution":{"observed_at":"2026-08-07T12:35:42.145814Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:52.492021Z","title":"Ramasesh, Ambrose Slone, Cem Anil, Imanol Schlag, Theo Gutman - Solo, Yuhuai Wu, Behnam Neyshabur, Guy Gur - Ari, and Vedant Misra","venue":null,"work_id":"444cd681-66ed-49e9-b038-58f0801a6133","year":2022},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:42.210918Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:5106d549fec1e2844af9a00474983d57158f9093b7c1e0ac73b23ac9b43fee9c","observation_id":"cab79f45-d389-41a8-8a80-687af203b081","resolution":{"observed_at":"2026-08-07T12:35:52.531814Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.14382","last_updated":"2025-02-20T09:18:53Z","snapshot_observed_at":"2026-08-07T18:02:33.909462Z","submitted_at":"2025-02-20T09:18:53Z","title":"S*: Test Time Scaling for Code Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.14382","snapshot_observed_at":"2026-08-07T12:35:42.308005Z","title":"Gonzalez, and Ion Stoica","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:42.308005Z"},"links":{"cited_paper":"/paper/2502.14382","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:c40da5cf954b36d5e523378ed1e70c9c036767a77e9a53854da406383952aa3d","observation_id":"4479646c-d535-474b-a217-90628cfbcd97","resolution":{"observed_at":"2026-08-07T12:35:42.308005Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:42.400638Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:42.400638Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:a5603e39c596d88c715df473538f8b64f6684d6051493324984c162aa933228d","observation_id":"b0113fdd-e786-46ce-8dc0-b6dc635a2423","resolution":{"observed_at":"2026-08-07T12:35:42.400638Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:52.283785Z","title":null,"venue":null,"work_id":"3102674c-3282-4413-8bfc-1d56add8880c","year":2023},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:42.554587Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:40c9ebbe9506dba538f7ad4b4bf3dbb6222ca4f438d98bc1ae940d2d0da24012","observation_id":"b428acd3-fbf5-40ca-a227-97c81500c7aa","resolution":{"observed_at":"2026-08-07T12:35:52.357938Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.09858","last_updated":"2025-04-14T04:08:16Z","snapshot_observed_at":"2026-08-08T01:06:59.135762Z","submitted_at":"2025-04-14T04:08:16Z","title":"Reasoning Models Can Be Effective Without Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.09858","snapshot_observed_at":"2026-08-07T12:35:42.651354Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:42.651354Z"},"links":{"cited_paper":"/paper/2504.09858","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:e9d053884db5345934ed63e697cbec342daf3f0ab33a11578a6c292e58eff5ca","observation_id":"263c395e-1ac1-487e-969c-2663dfa30520","resolution":{"observed_at":"2026-08-07T12:35:42.651354Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:42.805021Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:42.805021Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:c1095b107e76b211407e85c1dd25d9dd7b89de8a18b4c5e4d5a3a4ebb9f18dc3","observation_id":"bcecd26b-746a-45ff-a307-6bd3d0231f70","resolution":{"observed_at":"2026-08-07T12:35:42.805021Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:52.085827Z","title":null,"venue":null,"work_id":"c9221661-9bd6-4f7d-95e9-7e1c45766ef3","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:42.898673Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:2cc5736a83ad956a15bd589c3856e7e1af53c5525b262b369454b13e7674fdf7","observation_id":"29b9b933-2026-4d1a-bbf5-daf5a69a3694","resolution":{"observed_at":"2026-08-07T12:35:52.163271Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.19393","last_updated":"2025-03-01T06:07:39Z","snapshot_observed_at":"2026-07-06T20:29:11.710285Z","submitted_at":"2025-01-31T18:48:08Z","title":"s1: Simple test-time scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.19393","snapshot_observed_at":"2026-08-07T12:35:42.998263Z","title":"Cand \\` e s, and Tatsunori Hashimoto","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:42.998263Z"},"links":{"cited_paper":"/paper/2501.19393","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:151731a1e6a53f0481fe27d880f1d0b5810610ca1c807fdea9eab5c3be3a05bd","observation_id":"20423e41-95da-43c5-9ae8-1aabc36ca944","resolution":{"observed_at":"2026-08-07T12:35:42.998263Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:51.891365Z","title":null,"venue":null,"work_id":"fd5a3057-ece6-462d-8d76-2d81ffe3ab59","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:43.088852Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:ba6895deaa1d2396fa3b126ea0713f6c1e9fa6db907713b2f413b14d08e1c53c","observation_id":"2457bc5c-ba33-480a-ae1f-1ba8c22e01d6","resolution":{"observed_at":"2026-08-07T12:35:52.000124Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:51.684739Z","title":null,"venue":null,"work_id":"ed720a78-1fcb-4084-a607-a23aa6a04a33","year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:43.215053Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:484ddf22d0753fbaa4f5adb3563373d46246ce04c58397485b04cc295c9aa926","observation_id":"a64be304-b120-400f-aac8-ed44db0f0cd4","resolution":{"observed_at":"2026-08-07T12:35:51.794963Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.13367","last_updated":"2025-04-17T22:16:30Z","snapshot_observed_at":"2026-08-07T16:01:17.031141Z","submitted_at":"2025-04-17T22:16:30Z","title":"THOUGHTTERMINATOR: Benchmarking, Calibrating, and Mitigating Overthinking in Reasoning Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.13367","snapshot_observed_at":"2026-08-07T12:35:43.292037Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:43.292037Z"},"links":{"cited_paper":"/paper/2504.13367","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:e83cb4d56dccb51d0ce6e19dc788aea4472369222e0c242f550c4c3b71fa291d","observation_id":"61691de3-89af-489a-b581-36ef3b0a6fab","resolution":{"observed_at":"2026-08-07T12:35:43.292037Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1007/978-3-031-72775-7","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T00:03:45.533247Z","title":null,"venue":"Lecture notes in computer science","work_id":"13907969-3d52-4279-8e3a-188e34568644","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:43.398623Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:f390f1adf4e6753dd41b621ce825332cca99bdc3f1b3e2934a45eedd29e4cb7c","observation_id":"2474d358-c620-4d73-af7e-cc73081e9058","resolution":{"observed_at":"2026-08-07T12:35:47.834851Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:43.491241Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:43.491241Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:99b462f04a25c0d5d63e4ba8b0feeb2f5aa2d351db6c62defaf6cd69ee53f093","observation_id":"463a7494-a520-402a-aff9-d6005d60cd7c","resolution":{"observed_at":"2026-08-07T12:35:43.491241Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.16033","last_updated":"2025-09-03T03:41:08Z","snapshot_observed_at":"2026-08-10T08:34:28.320325Z","submitted_at":"2024-10-18T04:38:21Z","title":"TreeBoN: Enhancing Inference-Time Alignment with Speculative Tree-Search and Best-of-N Sampling","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.16033","snapshot_observed_at":"2026-08-07T12:35:43.551478Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:43.551478Z"},"links":{"cited_paper":"/paper/2410.16033","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:b1d1dd73ec15debb80901a082f732555584d55a791d6f360f1f477295e17b25b","observation_id":"98ceab6f-63fb-4f54-b4a1-9b16b5104084","resolution":{"observed_at":"2026-08-07T12:35:43.551478Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.07572","last_updated":"2025-03-10T17:40:43Z","snapshot_observed_at":"2026-08-07T17:16:00.148793Z","submitted_at":"2025-03-10T17:40:43Z","title":"Optimizing Test-Time Compute via Meta Reinforcement Fine-Tuning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.07572","snapshot_observed_at":"2026-08-07T12:35:43.664836Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:43.664836Z"},"links":{"cited_paper":"/paper/2503.07572","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:27fec4b50f4f6b6efafde667ef72cc8b7d5f1663153e3c2db38d5535db36f2b7","observation_id":"95308d0d-230e-4fb2-9a55-11b1e9187da4","resolution":{"observed_at":"2026-08-07T12:35:43.664836Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:51.470423Z","title":null,"venue":null,"work_id":"b4ba6d83-fb60-4ddb-901f-d72bf8589e90","year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:43.761716Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:263e8d602bfdff86c5c0bde268a115b8e4586eccfb8dcbd714838772d4615d0b","observation_id":"92a12974-8258-4b92-9390-039215357227","resolution":{"observed_at":"2026-08-07T12:35:51.569576Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-07T12:35:43.849088Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:43.849088Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:aec91d1ffaf9393e11d0f651de61d820e41698812145a2116c71e95de2a8d105","observation_id":"7321962b-6c29-4edc-9ef3-f96841678463","resolution":{"observed_at":"2026-08-07T12:35:43.849088Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:51.129925Z","title":null,"venue":null,"work_id":"37356232-9129-4611-960c-7d89867f84b1","year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:43.950798Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:bc925e26b5b0b6866583ea537f473ea9fa86f7bd9f30f0e5d7071c6cd4021de1","observation_id":"711e9cea-9f23-4659-a416-b063b4b0362c","resolution":{"observed_at":"2026-08-07T12:35:51.252969Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:50.870712Z","title":null,"venue":null,"work_id":"5b29069f-b562-4759-849a-4a57568ae2a0","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.088639Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:82c74bb2ad2db0855952b16cb1e9414aea9d53a41af82532ccb07e412cfbaf27","observation_id":"4124fbf7-f3b5-4065-b9fd-d54c00b222a7","resolution":{"observed_at":"2026-08-07T12:35:51.020048Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-07T12:35:44.172805Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.172805Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:6573483f91a229e616b6e4970bd3e67cdf6ec8bb88154858f3c08a6f8d498be5","observation_id":"094b410d-d8e2-435e-a7bd-f0a73d31e1f4","resolution":{"observed_at":"2026-08-07T12:35:44.172805Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.00127","last_updated":"2025-04-30T18:48:06Z","snapshot_observed_at":"2026-08-07T17:34:13.816265Z","submitted_at":"2025-04-30T18:48:06Z","title":"Between Underthinking and Overthinking: An Empirical Study of Reasoning Length and correctness in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.00127","snapshot_observed_at":"2026-08-07T12:35:44.268960Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.268960Z"},"links":{"cited_paper":"/paper/2505.00127","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:c8aa997fbf7b62b368372109438b03451df372b58b703b0e3e0db637512c5701","observation_id":"8e091db0-109b-4b9f-9aa7-5bbab5644f17","resolution":{"observed_at":"2026-08-07T12:35:44.268960Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.16419","last_updated":"2025-08-21T19:14:40Z","snapshot_observed_at":"2026-08-07T04:27:23.738927Z","submitted_at":"2025-03-20T17:59:38Z","title":"Stop Overthinking: A Survey on Efficient Reasoning for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.16419","snapshot_observed_at":"2026-08-07T12:35:44.375801Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.375801Z"},"links":{"cited_paper":"/paper/2503.16419","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:db7ffcf8a6ad23481662e16c727d74f5c3dae95b0f1ff0668a0ed5c96c52691a","observation_id":"859f24d2-079a-4aeb-b75b-bbe597c4893d","resolution":{"observed_at":"2026-08-07T12:35:44.375801Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:44.436737Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.436737Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:b67f4154fedb53adfbb7492b55fdf8c75e96bdda6e010accaf799ba4e8494067","observation_id":"00f7280f-778d-4cd4-a139-d4bab3d8d8ec","resolution":{"observed_at":"2026-08-07T12:35:44.436737Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.09818","last_updated":"2025-03-21T05:54:00Z","snapshot_observed_at":"2026-08-09T20:05:31.409634Z","submitted_at":"2024-05-16T05:23:41Z","title":"Chameleon: Mixed-Modal Early-Fusion Foundation Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.09818","snapshot_observed_at":"2026-08-07T12:35:44.521287Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.521287Z"},"links":{"cited_paper":"/paper/2405.09818","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:e06afedf521dcbfa8273e9cc5180c0e7d8745c01ecbded16767953b6c692e492","observation_id":"b5a3a1fc-c613-4596-af87-3b98847edaf4","resolution":{"observed_at":"2026-08-07T12:35:44.521287Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.14275","last_updated":"2022-11-25T18:19:44Z","snapshot_observed_at":"2026-08-01T02:16:43.109337Z","submitted_at":"2022-11-25T18:19:44Z","title":"Solving math word problems with process- and outcome-based feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.14275","snapshot_observed_at":"2026-08-07T12:35:44.602461Z","title":"Francis Song, Noah Y","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.602461Z"},"links":{"cited_paper":"/paper/2211.14275","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:abfa70663b9d369d0ac826291be439c34af5ba363067e08f7cb0d9faa5c8f471","observation_id":"5ba2f181-830d-4e3f-9a4d-e8fe411fd44f","resolution":{"observed_at":"2026-08-07T12:35:44.602461Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.17017","last_updated":"2025-02-04T03:59:49Z","snapshot_observed_at":"2026-08-05T15:40:45.171501Z","submitted_at":"2024-08-30T05:14:59Z","title":"Reasoning Aware Self-Consistency: Leveraging Reasoning Paths for Efficient LLM Sampling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.17017","snapshot_observed_at":"2026-08-07T12:35:44.685044Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.685044Z"},"links":{"cited_paper":"/paper/2408.17017","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:a3f93e7c8086b315b770611bb98a3b47462f73156aa1734cc468a091dcd7f921","observation_id":"1528cdce-a3db-4ea2-b691-7424cdd57358","resolution":{"observed_at":"2026-08-07T12:35:44.685044Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:50.421100Z","title":null,"venue":null,"work_id":"79a0bafe-68cf-4dd0-91f7-3fb60576429f","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.751691Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:a82def9ffc4e42e090492d47b4bd92e4beca9e33a4ec697db5f0a9c3fd5f28ce","observation_id":"e6f58211-698d-4e2c-83d1-497c01071647","resolution":{"observed_at":"2026-08-07T12:35:50.617379Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.16721","last_updated":"2025-05-01T22:22:43Z","snapshot_observed_at":"2026-07-06T19:56:48.086175Z","submitted_at":"2024-11-23T02:17:17Z","title":"Steering Away from Harm: An Adaptive Approach to Defending Vision Language Model Against Jailbreaks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.16721","snapshot_observed_at":"2026-08-07T12:35:44.854189Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.854189Z"},"links":{"cited_paper":"/paper/2411.16721","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:fe57559d4deacb703c6466ee80acbcd573a5c0a8859adbc9f8e300687fdf3191","observation_id":"63b2911c-5e93-4a7a-a278-77d7ffac9b30","resolution":{"observed_at":"2026-08-07T12:35:44.854189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:44.947344Z","title":"Le, Ed H","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:44.947344Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:e27b18bfcbca322ec9f21e0a6345129f66d1b2c061ca20c2855dd6a9b2582f0a","observation_id":"ef0a90fb-3e58-470e-8b94-4e2cbfda6dd2","resolution":{"observed_at":"2026-08-07T12:35:44.947344Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.18585","last_updated":"2025-02-18T16:51:53Z","snapshot_observed_at":"2026-08-09T22:51:48.255703Z","submitted_at":"2025-01-30T18:58:18Z","title":"Thoughts Are All Over the Place: On the Underthinking of o1-Like LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.18585","snapshot_observed_at":"2026-08-07T12:35:45.024969Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.024969Z"},"links":{"cited_paper":"/paper/2501.18585","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:b00b972cd54f2ca3be0ad7a45f94eb53664be0ee4a716b9f01f6a00a1e5c1083","observation_id":"62524ee9-31ba-4b0c-8ff7-55363b19f075","resolution":{"observed_at":"2026-08-07T12:35:45.024969Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.20631","last_updated":"2025-01-26T23:16:36Z","snapshot_observed_at":"2026-07-06T20:14:27.082070Z","submitted_at":"2024-12-30T00:40:35Z","title":"Slow Perception: Let's Perceive Geometric Figures Step-by-step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.20631","snapshot_observed_at":"2026-08-07T12:35:45.102203Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.102203Z"},"links":{"cited_paper":"/paper/2412.20631","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:91c2864b148883b5baa2bb617e092f701e4025923ae9d72c14ee2c4e8baf2e6e","observation_id":"7c6372c0-c03e-4d01-a43d-ac18f1d08388","resolution":{"observed_at":"2026-08-07T12:35:45.102203Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:45.171425Z","title":"Chi, Quoc V","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.171425Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:fb4413ec0480dfb057e415254d48c475f520c68565783735b6cf07daf9359ed9","observation_id":"facc8f97-a9ac-4559-90ee-2325613362c9","resolution":{"observed_at":"2026-08-07T12:35:45.171425Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.07165","last_updated":"2025-04-09T17:59:02Z","snapshot_observed_at":"2026-08-07T16:07:06.340888Z","submitted_at":"2025-04-09T17:59:02Z","title":"Perception in Reflection","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.07165","snapshot_observed_at":"2026-08-07T12:35:45.255651Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.255651Z"},"links":{"cited_paper":"/paper/2504.07165","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:6aaf1c3f92436000ad9bf851969d092bfe350654f46c6b1ff3547956f3425365","observation_id":"7df7f307-3356-4a8c-893d-30388abdab7d","resolution":{"observed_at":"2026-08-07T12:35:45.255651Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:45.320189Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.320189Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:e5146fcb77127a69911f9a406651efb02c36b3b027d1a7d180302492342a282d","observation_id":"c79f1a58-3c44-47f8-a907-5b3a3aaac30a","resolution":{"observed_at":"2026-08-07T12:35:45.320189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:50.197049Z","title":null,"venue":null,"work_id":"7e122350-96d3-4ac1-8cd6-61c09f1836c6","year":2023},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.428439Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:416eed9bf03ba44e145e3203ecc5640f9e34d10845bd0c31170de8b0142b39f4","observation_id":"d04086fe-e7b7-4a85-ba66-a8109b7b35ed","resolution":{"observed_at":"2026-08-07T12:35:50.258215Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02712","last_updated":"2025-03-04T00:49:07Z","snapshot_observed_at":"2026-08-09T22:43:34.534202Z","submitted_at":"2024-10-03T17:36:33Z","title":"LLaVA-Critic: Learning to Evaluate Multimodal Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.02712","snapshot_observed_at":"2026-08-07T12:35:45.496701Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.496701Z"},"links":{"cited_paper":"/paper/2410.02712","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:385237cfdb9edb64660f5c503bbf3a926da2980889b09d7de1154d0fa9235879","observation_id":"1308ad25-3d16-4e2b-b6bf-4c6b1eff3632","resolution":{"observed_at":"2026-08-07T12:35:45.496701Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.18600","last_updated":"2025-03-03T17:08:21Z","snapshot_observed_at":"2026-08-07T17:49:06.764625Z","submitted_at":"2025-02-25T19:36:06Z","title":"Chain of Draft: Thinking Faster by Writing Less","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.18600","snapshot_observed_at":"2026-08-07T12:35:45.581504Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.581504Z"},"links":{"cited_paper":"/paper/2502.18600","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:b2d39c3746cd49b18018e061d4ca9897ceb674798cfc85ec4899ed717634aec2","observation_id":"53a62993-42ea-44a3-ab43-32915fe4e47f","resolution":{"observed_at":"2026-08-07T12:35:45.581504Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:45.655639Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.655639Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:9695271985f216ffb3c87a9ffc619a4bb6b6de6623645f513aa98d290fb65092","observation_id":"042f0b57-b1cb-4aeb-bb57-8155c7d129de","resolution":{"observed_at":"2026-08-07T12:35:45.655639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.09560","last_updated":"2025-06-05T07:22:50Z","snapshot_observed_at":"2026-07-06T20:36:12.548343Z","submitted_at":"2025-02-13T18:11:34Z","title":"EmbodiedBench: Comprehensive Benchmarking Multi-modal Large Language Models for Vision-Driven Embodied Agents","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.09560","snapshot_observed_at":"2026-08-07T12:35:45.743237Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.743237Z"},"links":{"cited_paper":"/paper/2502.09560","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:559192866c3ba1be9e58d4c7389dc8f1ab7a71cdd8a3371e072d7327cff75944","observation_id":"59a88a82-5b9f-4f4a-82d4-8966f718dea5","resolution":{"observed_at":"2026-08-07T12:35:45.743237Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.12329","last_updated":"2026-06-02T22:05:11Z","snapshot_observed_at":"2026-08-08T01:16:34.338605Z","submitted_at":"2025-04-12T21:25:32Z","title":"Speculative Thinking: Enhancing Small-Model Reasoning with Large Model Guidance at Inference Time","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.12329","snapshot_observed_at":"2026-08-07T12:35:45.833987Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.833987Z"},"links":{"cited_paper":"/paper/2504.12329","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:7ca0ee3d2afe6ba1531328724dbdeefeb6beb37ead0d81a4e66bb980b1bfcd1a","observation_id":"0949a9d6-4506-41a2-8c4a-8c506d1afe5e","resolution":{"observed_at":"2026-08-07T12:35:45.833987Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:45.909816Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:45.909816Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:b0b3dd51a3162b22ad11772666f0ef4001ab39de44c2f3dc0742aaa2ec6e5017","observation_id":"607b2cd3-573f-441d-8fbc-c84f9ed2d53d","resolution":{"observed_at":"2026-08-07T12:35:45.909816Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:46.007790Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:46.007790Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:1a86be811bd6c81a4a30869b191525556a8369b683ab949a5f9d304a775a4fcf","observation_id":"901c11b5-abaf-4f11-90d4-e3d215c32ec5","resolution":{"observed_at":"2026-08-07T12:35:46.007790Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.07954","last_updated":"2025-04-10T17:58:27Z","snapshot_observed_at":"2026-08-07T16:06:45.172290Z","submitted_at":"2025-04-10T17:58:27Z","title":"Perception-R1: Pioneering Perception Policy with Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.07954","snapshot_observed_at":"2026-08-07T12:35:46.108172Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:46.108172Z"},"links":{"cited_paper":"/paper/2504.07954","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:c989345285097bf434c91a5bfb7a2568d0efd102caafcbd4c7c5c8c1beafe9f4","observation_id":"3a2ae015-0241-454b-8d78-6ab933a81cb4","resolution":{"observed_at":"2026-08-07T12:35:46.108172Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:46.217788Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:46.217788Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:41855a7b96eaa1903180acc55f8633aa6beb105fb16a37db9c8ea5786ce6d486","observation_id":"0300e877-9f4c-4e77-ab3f-61f26e6939b1","resolution":{"observed_at":"2026-08-07T12:35:46.217788Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:50.021789Z","title":null,"venue":null,"work_id":"6487f404-423e-485e-9292-467870045efc","year":2022},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:46.305865Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:d417a2779efe2bacbb4b902eb358eeffe0f038f7b0e14a30cfaf2b6f74e8a263","observation_id":"e9d99be9-908d-473e-ba87-f94288714a22","resolution":{"observed_at":"2026-08-07T12:35:50.090633Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.12215","last_updated":"2025-03-03T15:29:43Z","snapshot_observed_at":"2026-08-07T18:13:29.325301Z","submitted_at":"2025-02-17T07:21:11Z","title":"Revisiting the Test-Time Scaling of o1-like Models: Do they Truly Possess Test-Time Scaling Capabilities?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.12215","snapshot_observed_at":"2026-08-07T12:35:46.401699Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:46.401699Z"},"links":{"cited_paper":"/paper/2502.12215","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:27d6c2f209c1adc5fd92cc2bbf70df1fef46a0f0c009c4e97c45ffc016ccb8de","observation_id":"12b8ad7e-216d-4768-ae81-02c9df415f4b","resolution":{"observed_at":"2026-08-07T12:35:46.401699Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:49.867781Z","title":"Tenenbaum, and Chuang Gan","venue":null,"work_id":"de2ecfe8-fe2f-4c46-a29e-4d3d2758fdd7","year":2023},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:46.484766Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:35c29a2951e409c0e1d4ebc81e27778b89f4263a16eff40c15c302c1c5b9e470","observation_id":"f4ccc554-06d2-4856-9bba-f106a226aae0","resolution":{"observed_at":"2026-08-07T12:35:49.923682Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:49.476074Z","title":"Nguyen, Jun Sun, and Tat - Seng Chua","venue":null,"work_id":"700594ed-9f1d-4cbe-b82e-ab0276576393","year":2024},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:46.575363Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:f405dd769b057c797ceea063f44df8a12a99fdf1c8006ca8cda6ef3766375eda","observation_id":"b24861f7-9105-43fc-823c-257f46497add","resolution":{"observed_at":"2026-08-07T12:35:49.707604Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00511","last_updated":"2025-02-13T07:35:08Z","snapshot_observed_at":"2026-08-10T01:19:12.603891Z","submitted_at":"2025-02-01T18:09:49Z","title":"Bridging Internal Probability and Self-Consistency for Effective and Efficient LLM Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.00511","snapshot_observed_at":"2026-08-07T12:35:46.674252Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:46.674252Z"},"links":{"cited_paper":"/paper/2502.00511","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:4de9dc53ec96fbb504871d6902e2e40a7347ad226a5a6e5f366bd8dfa708481e","observation_id":"fcfcb776-138f-471a-abff-533d179ebdbb","resolution":{"observed_at":"2026-08-07T12:35:46.674252Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T12:35:48.962035Z","title":null,"venue":null,"work_id":"118c97d5-31f1-4b9c-84a8-b00b466b5e0f","year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:46.740725Z"},"links":{"citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:77982943cfcd1725844310e180f9b39a1958db1e6e9f37c46391faad0d3160a1","observation_id":"aecaf179-4011-4d7c-b0d9-87b9b5ed2b02","resolution":{"observed_at":"2026-08-07T12:35:49.228690Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16084","last_updated":"2025-06-30T15:59:26Z","snapshot_observed_at":"2026-07-06T21:13:13.686703Z","submitted_at":"2025-04-22T17:59:56Z","title":"TTRL: Test-Time Reinforcement Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.16084","snapshot_observed_at":"2026-08-07T12:35:46.811009Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:46.811009Z"},"links":{"cited_paper":"/paper/2504.16084","citing_paper":"/paper/2505.24863"},"observation_digest":"sha256:4dd792cba46dd7621810b12e92dbe48fedb3d5a2935b428c9adcad9ea7d00272","observation_id":"bd90f597-4e8b-4684-bda7-e04f38162630","resolution":{"observed_at":"2026-08-07T12:35:46.811009Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.24863","last_updated":"2025-05-30T17:58:36Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T20:38:56.380571Z","submitted_at":"2025-05-30T17:58:36Z","title":"AlphaOne: Reasoning Models Thinking Slow and Fast at Test Time"},"reference_resolution":{"displayed":86,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":80,"verified_exact":1,"verified_fuzzy":4},"total_outbound_references":86},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 86 of 86 outbound references and 11 inbound Pith citation observations for arXiv:2505.24863."}