{"as_of":"2026-08-14T17:30:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3daf3baf655a0d6957fc1f65a3d0f4ef3c82786d7b137907bacba4b06d093624","coverage":[{"denominator":43,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":43,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T23:10:27.460894Z","state":"measured"},{"denominator":44,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":44,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T17:54:17.954851Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-06T17:54:19.431591Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"cited_work":{"arxiv_id":"2506.19492","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.19492","snapshot_observed_at":"2026-08-06T17:54:19.431591Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","venue":"cs.CL","work_id":"25f70f04-0a19-4012-851c-ca91ed56f8f4","year":2025},"citing_paper":{"arxiv_id":"2507.09662","last_updated":"2025-07-13T14:51:59Z","snapshot_observed_at":"2026-08-14T02:30:04.338310Z","submitted_at":"2025-07-13T14:51:59Z","title":"Towards Concise and Adaptive Thinking in Large Reasoning Models: A Survey","version":1},"reference_index":220,"source":"arxiv_source","source_observed_at":"2026-08-06T17:54:17.954851Z"},"links":{"cited_paper":"/paper/2506.19492","citing_paper":"/paper/2507.09662"},"observation_digest":"sha256:6a1961d1e646929409a9839c977665dba3b1c0753cc0f5febc3875731722bf59","observation_id":"ba371586-8389-4de2-b0ec-df63a0b3ba18","resolution":{"observed_at":"2026-08-06T17:54:19.436865Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.19492/citation-record","integrity":"/paper/2506.19492/integrity","json":"/paper/2506.19492/citation-record.json","paper":"/paper/2506.19492"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"1606.06565","last_updated":"2016-07-25T17:23:29Z","snapshot_observed_at":"2026-07-06T05:00:46.434335Z","submitted_at":"2016-06-21T13:37:05Z","title":"Concrete Problems in AI Safety","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.06565","snapshot_observed_at":"2026-08-06T23:10:24.360508Z","title":"Concrete problems in ai safety.arXiv preprint arXiv:1606.06565, 2016","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.360508Z"},"links":{"cited_paper":"/paper/1606.06565","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:e9b250a10d8d483674e7ade7b713a98e54e55f5acab13e3025aaece524acfbc2","observation_id":"548782f0-c7a1-4232-8917-7d7783246534","resolution":{"observed_at":"2026-08-06T23:10:24.360508Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:24.429267Z","title":"Chain-of-thought reasoning in the wild is not always faithful","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.429267Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:28497a8853845a0316564a6fc214e2d17a526034323853694ff30f326910e325","observation_id":"1d265062-f11e-46eb-b92c-36406210a084","resolution":{"observed_at":"2026-08-06T23:10:24.429267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.08679","last_updated":"2026-06-16T17:36:22Z","snapshot_observed_at":"2026-08-14T07:47:45.666218Z","submitted_at":"2025-03-11T17:56:30Z","title":"Chain-of-Thought Reasoning In The Wild Is Not Always Faithful","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.08679","snapshot_observed_at":"2026-08-06T23:10:24.498510Z","title":"Chain-of-thought reasoning in the wild is not always faithful, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.498510Z"},"links":{"cited_paper":"/paper/2503.08679","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:d0fbe3751c772ed18a095a89e985ea0892c299e802540cab39590dd479cbcbc3","observation_id":"e4047ab3-1937-40f7-8b9c-f670e5675eb9","resolution":{"observed_at":"2026-08-06T23:10:24.498510Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:24.581791Z","title":"Training language models to reason efficiently.arXiv preprint arXiv:2502.04463, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.581791Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:3ca12f34fba14c55b0f40bc7b2b8b97b75d4eb21b9af173c40655755e6b80cc7","observation_id":"6bc124e3-dd59-4762-b3e3-a751462704b5","resolution":{"observed_at":"2026-08-06T23:10:24.581791Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.11926","last_updated":"2025-03-14T23:50:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-14T23:50:34Z","title":"Monitoring Reasoning Models for Misbehavior and the Risks of Promoting Obfuscation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.11926","snapshot_observed_at":"2026-08-06T23:10:24.679284Z","title":"Guan, Aleksander Madry, Wojciech Zaremba, Jakub Pachocki, and David Farhi","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.679284Z"},"links":{"cited_paper":"/paper/2503.11926","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:cb1fb868f15303fc31d1967c932718cb19c704fa302ba1dcefbe852f21cc5d72","observation_id":"639c2dd4-23c0-422a-b5cf-7eb820033e8d","resolution":{"observed_at":"2026-08-06T23:10:24.679284Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:31.097257Z","title":"Monitoring reasoning models for misbehavior and the risks of promoting obfuscation","venue":null,"work_id":"c400e572-937f-42ab-a403-ef65c5041465","year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.755948Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:86e9c7b411b7d74e64904903f70c330a9aa9cc3c90273166b68a7d0d0daffc7a","observation_id":"7c43f894-d01d-4e97-b5bc-acd855f5aac8","resolution":{"observed_at":"2026-08-06T23:10:31.101313Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:31.084242Z","title":"Weak-to-strong generalization: Eliciting strong capabilities with weak supervision, 2023","venue":null,"work_id":"03a7ca2e-780c-4b90-ac24-775a872c10cb","year":2023},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.826795Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:3636c155b0e2c345655f5e856c931c485ed42a4b554f558df9ee79c7e69286c7","observation_id":"0183fd02-3508-4fa5-8ae8-1d3ce12980d9","resolution":{"observed_at":"2026-08-06T23:10:31.088807Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.21187","last_updated":"2025-02-01T07:57:37Z","snapshot_observed_at":"2026-08-01T16:43:44.704797Z","submitted_at":"2024-12-30T18:55:12Z","title":"Do NOT Think That Much for 2+3=? On the Overthinking of o1-Like LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.21187","snapshot_observed_at":"2026-08-06T23:10:24.922177Z","title":"Do not think that much for 2+ 3=? on the overthinking of o1-like llms.arXiv preprint arXiv:2412.21187, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.922177Z"},"links":{"cited_paper":"/paper/2412.21187","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:4a1c57dbc1111bdf4ffd1d9919000eb4e7d4478eab5372aa1ad417fa68d103b4","observation_id":"e96086b5-7dee-4f0d-96e2-315d89c06a51","resolution":{"observed_at":"2026-08-06T23:10:24.922177Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:31.071117Z","title":"Reasoning models don’t always say what they think","venue":null,"work_id":"13f4883b-bc8e-40f3-b0f0-5de8260c7c87","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.009936Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:9a0b66da6c12f3883922ecbbe1eeebfc0c9b8ad25ab21ac0db5230f673ada8e1","observation_id":"60b1c90b-26cc-4740-b5d6-3888a514e631","resolution":{"observed_at":"2026-08-06T23:10:31.074962Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-08-13T15:58:13.809876Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-06T23:10:25.135228Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.135228Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:4291e1161ce6a23a2c4b0c38cb72b756548118d94ab4bf0143a759f3f5c76140","observation_id":"07ca73a6-efda-4613-bfbc-57ee614e6c13","resolution":{"observed_at":"2026-08-06T23:10:25.135228Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10162","last_updated":"2024-06-29T00:28:47Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-14T16:26:20Z","title":"Sycophancy to Subterfuge: Investigating Reward-Tampering in Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10162","snapshot_observed_at":"2026-08-06T23:10:25.217290Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.217290Z"},"links":{"cited_paper":"/paper/2406.10162","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:dc974a6d454a12086f26481a84873fe6673568cce071fa802af6d98c273c083c","observation_id":"70f9d6c7-0bb0-43a5-8290-02d958728470","resolution":{"observed_at":"2026-08-06T23:10:25.217290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:31.058411Z","title":"A wolf in sheep’s clothing: Generalized nested jailbreak prompts can fool large language models easily","venue":null,"work_id":"facc5e9d-99de-4aea-89d7-6ce60aca94a5","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.285390Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:56888774addbd2d6bd23f307409177c73599265d2cea01ccb11c6bb3cf1b63a2","observation_id":"697d417d-401a-4482-a561-b06c548eb481","resolution":{"observed_at":"2026-08-06T23:10:31.062985Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:31.044850Z","title":"Reward tampering problems and solutions in reinforcement learning: A causal influence diagram perspective.Synthese, 198(Suppl 27): 6435–6467, 2021","venue":null,"work_id":"7ab4cbbe-ea82-4e3b-99c5-c053f04a9338","year":2021},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.380243Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:9693a7464e633115a8ccc75ef6bf319595efc3f1209fbfccfe72ff47fe8f62a6","observation_id":"a0f602a2-9bc3-4d3f-8c74-1fa6d18c06b4","resolution":{"observed_at":"2026-08-06T23:10:31.049640Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:25.467968Z","title":"Syceval: Evaluating llm sycophancy.arXiv preprint arXiv:2502.08177, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.467968Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:1850ae8d1c999779af49959c9565b064b6ca0146e8aed190b51764e5e0251e45","observation_id":"851d42d2-bced-48be-af82-31191296ee6e","resolution":{"observed_at":"2026-08-06T23:10:25.467968Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.910505Z","title":"Who’s asking? user personas and the mechanics of latent misalignment.Advances in Neural Information Processing Systems, 37:125967–126003, 2024","venue":null,"work_id":"0056ff8b-d755-4a71-93d9-e3bf6abb194e","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.541246Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:d41cd2d879c58055c16d6ced1262ee7a36eb90291c899d9f96adb20419b3dbd9","observation_id":"60f88d12-c1a4-4506-98c6-e08dc9d13f76","resolution":{"observed_at":"2026-08-06T23:10:31.005881Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.792289Z","title":"Alignment faking in large language models.CoRR, 2024","venue":null,"work_id":"9a850253-bf6b-4ab3-bc5f-a8ce7c7e9617","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.613807Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:082bdb6beec7d1a6464d8f27d2fa37cd3b973fa8862f5089fed2167c6ec54d68","observation_id":"8a313f76-fa7e-49ca-9b82-9c58476ae472","resolution":{"observed_at":"2026-08-06T23:10:30.816339Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:25.729917Z","title":"Olympiadbench: A challenging benchmark for promoting agi with olympiad-level bilingual multimodal scientific problems","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.729917Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:89a8a27d3194115d0fe9eb94324400bb7f1a60ecf80cffbdf5fec3b593286929","observation_id":"5c466036-a97c-4154-87d2-9e9a58e59176","resolution":{"observed_at":"2026-08-06T23:10:25.729917Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:25.812905Z","title":"C3ot: Generating shorter chain-of-thought without compromising effectiveness","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.812905Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:a88bc1d55e3af6d38285d00f35130e7481f5279d1b6120838226bb41eddd0312","observation_id":"36329b67-aed0-429e-95b3-0c8af85eec41","resolution":{"observed_at":"2026-08-06T23:10:25.812905Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.01141","last_updated":"2025-04-01T00:41:36Z","snapshot_observed_at":"2026-08-08T04:02:43.828729Z","submitted_at":"2025-03-03T03:48:20Z","title":"How Well do LLMs Compress Their Own Chain-of-Thought? A Token Complexity Approach","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.01141","snapshot_observed_at":"2026-08-06T23:10:25.880397Z","title":"How well do llms compress their own chain-of-thought? a token complexity approach, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.880397Z"},"links":{"cited_paper":"/paper/2503.01141","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:e51e318e5c74f692bfb4bbcb4d052372ec548265503b97d87c7b67e9fd11a3a4","observation_id":"91f29c4a-bcc5-4e15-8794-2a82d943533f","resolution":{"observed_at":"2026-08-06T23:10:25.880397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13860","last_updated":"2024-03-10T13:58:08Z","snapshot_observed_at":"2026-07-06T15:31:18.144952Z","submitted_at":"2023-05-23T09:33:38Z","title":"Jailbreaking ChatGPT via Prompt Engineering: An Empirical Study","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.13860","snapshot_observed_at":"2026-08-06T23:10:25.963095Z","title":"Jailbreaking chatgpt via prompt engineering: An empirical study.arXiv preprint arXiv:2305.13860, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.963095Z"},"links":{"cited_paper":"/paper/2305.13860","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:3246562fbe1de3577e81e66a96cfbf02f07826182684198858b29a1cfde475ce","observation_id":"dabfeedd-0961-4400-934d-aab56adf40b8","resolution":{"observed_at":"2026-08-06T23:10:25.963095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12570","last_updated":"2025-01-29T03:11:03Z","snapshot_observed_at":"2026-08-10T17:00:34.555363Z","submitted_at":"2025-01-22T01:35:11Z","title":"O1-Pruner: Length-Harmonizing Fine-Tuning for O1-Like Reasoning Pruning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12570","snapshot_observed_at":"2026-08-06T23:10:26.017771Z","title":"O1-pruner: Length-harmonizing fine-tuning for o1-like reasoning pruning.arXiv preprint arXiv:2501.12570, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.017771Z"},"links":{"cited_paper":"/paper/2501.12570","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:497a19e680d5ec4d31e772b4cebc34531a6022835986a887185cf2f77b315afb","observation_id":"23da9ed6-0f8d-41f2-bffa-d936c0ace837","resolution":{"observed_at":"2026-08-06T23:10:26.017771Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.09858","last_updated":"2025-04-14T04:08:16Z","snapshot_observed_at":"2026-08-08T01:06:59.135762Z","submitted_at":"2025-04-14T04:08:16Z","title":"Reasoning Models Can Be Effective Without Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.09858","snapshot_observed_at":"2026-08-06T23:10:26.064315Z","title":"Reasoning models can be effective without thinking, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.064315Z"},"links":{"cited_paper":"/paper/2504.09858","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:761b45bfee5fc20501c47c48718ffd8290e3d24b9c73fc39ef46e8ee2c3ac6ba","observation_id":"157ca23f-801b-430f-8ecb-79ca607520ae","resolution":{"observed_at":"2026-08-06T23:10:26.064315Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.631840Z","title":"Are self-explanations from large language models faithful? InFindings of the Association for Computational Linguistics ACL 2024, pages 295–337, 2024","venue":null,"work_id":"5fa3de32-a4b2-423a-85cc-2baa44ffd2e1","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.111750Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:e426067c6dfef37bc810a63a3b1e444800543f87e4410fa1fe27070756f21188","observation_id":"302d724a-88a7-427d-8cea-bfac70a48a75","resolution":{"observed_at":"2026-08-06T23:10:30.727087Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.04984","last_updated":"2025-01-14T20:16:01Z","snapshot_observed_at":"2026-08-14T10:15:12.633739Z","submitted_at":"2024-12-06T12:09:50Z","title":"Frontier Models are Capable of In-context Scheming","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.04984","snapshot_observed_at":"2026-08-06T23:10:26.170863Z","title":"Frontier models are capable of in-context scheming.arXiv preprint arXiv:2412.04984, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.170863Z"},"links":{"cited_paper":"/paper/2412.04984","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:c53d46fab06215deaa008a6501b5dfe69e964b29a9de06a69c9fee0755ea1cc1","observation_id":"8d016c3c-1c87-4295-98dd-c8eb93774d26","resolution":{"observed_at":"2026-08-06T23:10:26.170863Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.503889Z","title":"Self-training elicits concise reasoning in large language models.CoRR, 2025","venue":null,"work_id":"a87297d3-be05-4573-bf70-bbfb5ca1f649","year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.236575Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:8d111ce8084cc133e97b0b355437336c584013a06eebec7df9ae737d5d757369","observation_id":"28c01fdc-bcfb-41dd-9cbc-175f53360a1a","resolution":{"observed_at":"2026-08-06T23:10:30.583304Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.315614Z","title":"Show your work: Scratchpads for intermediate computation with language models","venue":null,"work_id":"daaf1774-d428-4d62-ac48-f1f3aea741d5","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.289426Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:de26e8e42cce35fe6a823178cc42707e88ceaa5974c32e1d98fbce4a90be9599","observation_id":"50a0b1e3-21f2-4cd2-91b4-da88d7ff332f","resolution":{"observed_at":"2026-08-06T23:10:30.395731Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.061662Z","title":"Learning to reason with llms","venue":null,"work_id":"03964fca-9905-4d2b-b516-e97c1cf1dc02","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.375978Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:5d5af6fccb77c1cefdb2722a9735c371aa27f20c4872b21424248f2c760db8ea","observation_id":"50b54a60-db45-4499-b960-bc72af29f821","resolution":{"observed_at":"2026-08-06T23:10:30.195407Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:26.439725Z","title":"Discovering language model behaviors with model-written evaluations","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.439725Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:f5443c5734cb71d625bf2b77210ee4129388f893222e6d95081681c7615d47c1","observation_id":"41429ca3-8eaf-4e6b-b9b1-ce061caa908c","resolution":{"observed_at":"2026-08-06T23:10:26.439725Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:29.837598Z","title":"Do Anything Now","venue":null,"work_id":"445b6b54-e098-4fb9-ac06-fb4f024c1f8a","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.505238Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:6fef90ad7a69ac1da13874e0c9f9a17d1e8554aaf5d57dd68e6575c34a13344e","observation_id":"55b4560e-fccb-42ea-9770-3ae2d9078c4c","resolution":{"observed_at":"2026-08-06T23:10:29.918948Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:26.574076Z","title":"Defining and characterizing reward gaming.Advances in Neural Information Processing Systems, 35:9460–9471, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.574076Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:0d3285b7eeee91a58ceb8c4369870efe4b586dc7beabfcf520a000f36b86c649","observation_id":"6293ab38-fd35-4d91-aa75-d87d80608643","resolution":{"observed_at":"2026-08-06T23:10:26.574076Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.16419","last_updated":"2025-08-21T19:14:40Z","snapshot_observed_at":"2026-08-11T13:10:23.709172Z","submitted_at":"2025-03-20T17:59:38Z","title":"Stop Overthinking: A Survey on Efficient Reasoning for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.16419","snapshot_observed_at":"2026-08-06T23:10:26.642653Z","title":"Stop overthinking: A survey on efficient reasoning for large language models.arXiv preprint arXiv:2503.16419, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.642653Z"},"links":{"cited_paper":"/paper/2503.16419","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:4cc90c71f0267e0fad550b0f26e3c1ae90c8137692add0b4325a94759c6f5571","observation_id":"a15feb78-3f2c-49ae-835a-4615d4cebdbd","resolution":{"observed_at":"2026-08-06T23:10:26.642653Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12599","last_updated":"2025-06-03T02:14:54Z","snapshot_observed_at":"2026-08-11T01:31:26.242281Z","submitted_at":"2025-01-22T02:48:14Z","title":"Kimi k1.5: Scaling Reinforcement Learning with LLMs","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12599","snapshot_observed_at":"2026-08-06T23:10:26.694880Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.694880Z"},"links":{"cited_paper":"/paper/2501.12599","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:695d6900bd4403c27c088f13c3efae8150306a92e0a85aa9dd9ab3632e4777f1","observation_id":"b04d6c25-3c44-4b35-8bad-021012d32e03","resolution":{"observed_at":"2026-08-06T23:10:26.694880Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:26.755752Z","title":"Qwen3, April 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.755752Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:5b76dfb5bafbbeff8c55dfee451c0479d4ca79dd5040293ec1a79ecd4ed29b63","observation_id":"b33198e8-c487-45a2-a93a-dbb6066d5cb3","resolution":{"observed_at":"2026-08-06T23:10:26.755752Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:26.845600Z","title":"Language models don’t always say what they think: Unfaithful explanations in chain-of-thought prompting.Advances in Neural Information Processing Systems, 36:74952–74965, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.845600Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:2da84c2a84e1db28f7c2bd7b6f2fcf6b4fcc8cfa974f5d6d1e77575d06843baa","observation_id":"838e7af5-614b-4797-ba56-5c4ae7712324","resolution":{"observed_at":"2026-08-06T23:10:26.845600Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:29.495544Z","title":null,"venue":null,"work_id":"166b0865-7a90-4f6f-854a-5092e9dd5bab","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.912149Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:ac7b871792645f3c6e566abc4452288c5c13e4a14617830146feb4cf029b7953","observation_id":"f6c5e37f-e127-415b-9a2c-0fb8a1b965af","resolution":{"observed_at":"2026-08-06T23:10:29.642751Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:29.214889Z","title":"Large language models often say one thing and do another","venue":null,"work_id":"7cc95cbc-e60e-4ddd-93f3-15b005a47483","year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.973650Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:58b1db01e670facd05e15f187abb56a22fe9b3c7b372b6870f6c575f30259bbe","observation_id":"887d404b-d7e8-4e37-9f48-572700bbd90f","resolution":{"observed_at":"2026-08-06T23:10:29.323877Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-08-12T09:06:50.363435Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15043","snapshot_observed_at":"2026-08-06T23:10:27.042013Z","title":"no free lunch","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.042013Z"},"links":{"cited_paper":"/paper/2307.15043","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:4c8794ce1ca08b179751d003969fb7ba7b01d3ef681b34f32f1c545cc3a0a3cb","observation_id":"653ad9f5-5790-4dd2-a177-8a3e4883d017","resolution":{"observed_at":"2026-08-06T23:10:27.042013Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:28.981803Z","title":null,"venue":null,"work_id":"487b18f5-0c96-4652-abff-834448203783","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.100826Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:0d89564ae9fd2e1693e0e0d6d41707463629a135242dd5e081f8c2fe2fc1db63","observation_id":"f57189ff-5941-47fe-9c36-df4163c9c159","resolution":{"observed_at":"2026-08-06T23:10:29.089201Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:28.770920Z","title":"IC” if the two responses show an inconsistency (refusal vs. answer, or different final results). - “CO","venue":null,"work_id":"0e4fa3b7-ef8c-4e29-9481-e000255667e2","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.144731Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:477336c809e6c62925cbb7db1b3c689ace86f56044411364d26403254910dee6","observation_id":"471e33ee-960e-48f7-9c1a-1ec2a0a71a2f","resolution":{"observed_at":"2026-08-06T23:10:28.869126Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:28.559443Z","title":null,"venue":null,"work_id":"25269b88-cc67-47aa-ad54-4cfa020bf629","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.213843Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:a1a484055b9b0681a3a3fefb6e106600466fc15f3ba7f85ede5b15e573459cec","observation_id":"9b4843ca-055d-464b-bd7b-e6d392f0c876","resolution":{"observed_at":"2026-08-06T23:10:28.670354Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:28.335035Z","title":null,"venue":null,"work_id":"07d205a9-2760-4098-8a80-8f032a8cf968","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.293021Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:525a4e6cfb72707beffcfe04708e6a5f074cff53be071e926df88d54662403e4","observation_id":"f0f8efbe-7847-4426-a1dd-a6774df564c8","resolution":{"observed_at":"2026-08-06T23:10:28.465354Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:28.119098Z","title":null,"venue":null,"work_id":"29879c81-d06b-4ba8-85ad-31504422820d","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.372720Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:9a3e4531cc2fa60c3f8f118083ef183ac13046879b0ba1f2065b8146603ffc27","observation_id":"ba0e5744-defa-4b10-aa7e-f23fd86b589b","resolution":{"observed_at":"2026-08-06T23:10:28.200050Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:27.929374Z","title":"Deploymode","venue":null,"work_id":"89104294-8e9d-4852-bee4-7fc7971e2152","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.460894Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:0e6fddc4dbcc93562a52c0b5cab945d0aa4a476b4e54503f911d748d4faa8868","observation_id":"fa8d8b32-6988-4364-92fe-122977d98eef","resolution":{"observed_at":"2026-08-06T23:10:28.026131Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-14T11:12:47.256853Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs"},"reference_resolution":{"displayed":43,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":28,"verified_exact":0,"verified_fuzzy":15},"total_outbound_references":43},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 43 of 43 outbound references and 1 inbound Pith citation observation for arXiv:2506.19492."}