{"as_of":"2026-08-16T23:08:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c689c411a3bfeef59b3195263109e67ea7b9958b3cf798d79a3a020e27ee0d04","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":14,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":14,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-16T06:30:59.297886+00:00","state":"measured"},{"denominator":14,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":14,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T04:39:55.101972Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-09T17:37:40.787506Z","title":"From sparse dependence to sparse attention: unveiling how chain-of-thought enhances transformer sample efficiency","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01694","last_updated":"2025-03-01T10:27:24Z","snapshot_observed_at":"2026-08-15T11:22:52.429751Z","submitted_at":"2025-02-02T18:19:14Z","title":"Metastable Dynamics of Chain-of-Thought Reasoning: Provable Benefits of Search, RL and Distillation","version":2},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-09T17:37:40.787506Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2502.01694"},"observation_digest":"sha256:7ef2ecfca9a28587555967e08460b5ccbe90572c2d7a51a825663719546e0d6f","observation_id":"7dda2a9b-e9a1-410b-b2db-4868a9765058","resolution":{"observed_at":"2026-08-09T17:37:40.787506Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-09T05:22:33.468853Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03275","last_updated":"2025-09-01T18:25:38Z","snapshot_observed_at":"2026-08-13T02:33:47.410604Z","submitted_at":"2025-02-05T15:33:00Z","title":"Token Assorted: Mixing Latent and Text Tokens for Improved Language Model Reasoning","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-09T05:22:33.468853Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2502.03275"},"observation_digest":"sha256:fcad17b98aa378f03180467278a2f37fc011f8a0b008f05f235a1171e82620c2","observation_id":"d3743882-f94d-451a-bbc8-96d8964191d8","resolution":{"observed_at":"2026-08-09T05:22:33.468853Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-16T04:39:55.101972Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of- thought enhances transformer sample efficiency","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.00926","last_updated":"2025-05-28T23:17:46Z","snapshot_observed_at":"2026-08-16T04:29:11.072998Z","submitted_at":"2025-05-02T00:07:35Z","title":"How Transformers Learn Regular Language Recognition: A Theoretical Study on Training Dynamics and Implicit Bias","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-16T04:39:55.101972Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2505.00926"},"observation_digest":"sha256:1de2e34740d69cdf7bda53a55f4ad33fa5cb8e80049151ea8ccc678dde69fe1c","observation_id":"471940bf-60e2-48b6-bc36-428bcb39696a","resolution":{"observed_at":"2026-08-16T04:39:55.101972Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-07T14:21:01.863974Z","title":"From sparse de- pendence to sparse attention: unveiling how chain-of-thought enhances transformer sample efficiency.arXiv preprint arXiv:2410.05459,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.19531","last_updated":"2025-05-26T05:33:26Z","snapshot_observed_at":"2026-08-14T23:57:20.901346Z","submitted_at":"2025-05-26T05:33:26Z","title":"Minimalist Softmax Attention Provably Learns Constrained Boolean Functions","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T14:21:01.863974Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2505.19531"},"observation_digest":"sha256:df89e8a942bc2feb644fc7df310acd9ee3bac8578c43d01ade0320fe0ae35cf3","observation_id":"939fe7ff-cdb3-4f8f-88fc-e807cdd348f8","resolution":{"observed_at":"2026-08-07T14:21:01.863974Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-07T12:46:39.840714Z","title":"From sparse dependence to sparse attention: unveiling how chain-of-thought enhances transformer sample efficiency","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.23683","last_updated":"2025-05-29T17:22:00Z","snapshot_observed_at":"2026-08-16T11:25:00.586713Z","submitted_at":"2025-05-29T17:22:00Z","title":"Learning Compositional Functions with Transformers from Easy-to-Hard Data","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-07T12:46:39.840714Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2505.23683"},"observation_digest":"sha256:50fa073fd141f712880ea42da8182c5bbcdbd552ed947d360fd7c66692d72117","observation_id":"4c59a0ee-c65c-4fe3-8215-7a17f3099a43","resolution":{"observed_at":"2026-08-07T12:46:39.840714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T22:10:17.592712Z","title":", Zhang, H","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.07571","last_updated":"2025-08-19T11:54:07Z","snapshot_observed_at":"2026-08-15T16:46:43.172924Z","submitted_at":"2025-08-11T03:05:36Z","title":"Towards Theoretical Understanding of Transformer Test-Time Computing: Investigation on In-Context Linear Regression","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-05T22:10:17.592712Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2508.07571"},"observation_digest":"sha256:1c44059bb9e7b63998a61d250b931cd1e8ebcd6f72e2224a2d942bdf65182be1","observation_id":"7ae8667b-f9ac-43d8-b442-f3e6c35a4cdc","resolution":{"observed_at":"2026-08-05T22:10:17.592712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-03T05:27:44.529687Z","title":"From sparse dependence to sparse attention: unveiling how chain-of-thought enhances transformer sample efficiency.arXiv preprint arXiv:2410.05459,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.02470","last_updated":"2026-05-31T23:31:59Z","snapshot_observed_at":"2026-08-15T18:11:00.951498Z","submitted_at":"2026-02-02T18:50:57Z","title":"Breaking the Reversal Curse in Autoregressive Language Models via Identity Bridge","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-03T05:27:44.529687Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2602.02470"},"observation_digest":"sha256:cbc33b5080f0c8cfba514d6fda594d65930ca4057fd85402b814569a9f7765e8","observation_id":"42f0197e-30c9-46bb-a399-aea76b88e17d","resolution":{"observed_at":"2026-08-03T05:27:44.529687Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-02T23:11:56.681936Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.14872","last_updated":"2026-06-29T01:04:09Z","snapshot_observed_at":"2026-08-15T22:28:45.285943Z","submitted_at":"2026-02-16T16:03:08Z","title":"On the Emergence of Implicit Curriculum in RLVR Learning Dynamics","version":3},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-02T23:11:56.681936Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2602.14872"},"observation_digest":"sha256:13a214808a1d1aa417a3f961d5f52896e0c4b43bbebc4494cd634baaf2b63a8c","observation_id":"4f8fbc10-55d4-40ce-82d7-6c4d511cdeff","resolution":{"observed_at":"2026-08-02T23:11:56.681936Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":"2410.05459","doi":"10.48550/arxiv.2410.05459","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency.ArXiv, abs/2410.05459","venue":"arXiv (Cornell University)","work_id":"12dec1cc-a66a-4ef5-bfe7-061e33762461","year":2024},"citing_paper":{"arxiv_id":"2604.22951","last_updated":"2026-07-08T19:29:03Z","snapshot_observed_at":"2026-08-15T23:07:26.295273Z","submitted_at":"2026-04-24T18:49:08Z","title":"The Power of Power Law: Asymmetry Enables Compositional Reasoning","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-05-08T11:49:49.787123Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2604.22951"},"observation_digest":"sha256:fd11a03f1fd8b7ecd0b97c9477121afd6a57afd5f372ba6a409a03e214a0af70","observation_id":"34b2affc-a800-43ae-86ad-2d78e95b882a","resolution":{"observed_at":"2026-05-11T19:31:08.381738Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-07-12T18:26:05.728364Z","title":"From sparse dependence to sparse attention: unveiling how chain-of-thought enhances transformer sample efficiency","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2604.22951","last_updated":"2026-07-08T19:29:03Z","snapshot_observed_at":"2026-08-15T23:07:26.295273Z","submitted_at":"2026-04-24T18:49:08Z","title":"The Power of Power Law: Asymmetry Enables Compositional Reasoning","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-07-12T18:26:05.728364Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2604.22951"},"observation_digest":"sha256:55e54e8608396574166bd0dd9bf93a1a58afe3ffb04a94e1973889f4a0a26489","observation_id":"8a1e30ca-2190-41fe-84d1-9b468517f84d","resolution":{"observed_at":"2026-07-12T18:26:05.728364Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":"2410.05459","doi":"10.48550/arxiv.2410.05459","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency.ArXiv, abs/2410.05459","venue":"arXiv (Cornell University)","work_id":"12dec1cc-a66a-4ef5-bfe7-061e33762461","year":2024},"citing_paper":{"arxiv_id":"2605.10019","last_updated":"2026-05-11T05:44:18Z","snapshot_observed_at":"2026-08-16T08:48:37.559402Z","submitted_at":"2026-05-11T05:44:18Z","title":"The two clocks and the innovation window: When and how generative models learn rules","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-12T03:15:45.257213Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2605.10019"},"observation_digest":"sha256:e308bd470ce696fd2541c4769d2f1541373d40d3c6e337352085d82e259f1df3","observation_id":"fe5c8a4d-2b15-42e1-a141-fe0270cbd952","resolution":{"observed_at":"2026-05-12T03:16:18.110627Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":"2410.05459","doi":"10.48550/arxiv.2410.05459","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency.ArXiv, abs/2410.05459","venue":"arXiv (Cornell University)","work_id":"12dec1cc-a66a-4ef5-bfe7-061e33762461","year":2024},"citing_paper":{"arxiv_id":"2605.23040","last_updated":"2026-05-21T21:13:14Z","snapshot_observed_at":"2026-08-02T04:36:31.374601Z","submitted_at":"2026-05-21T21:13:14Z","title":"Steered Generation via Gradient-Based Optimization on Sparse Query Features","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-25T05:31:29.510639Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2605.23040"},"observation_digest":"sha256:51c389a5530daab58cc98560b27acba93e5abb1edc2c9eed0c6099a7249b8190","observation_id":"4d011ce6-1534-47bc-9fb4-589d17a05c06","resolution":{"observed_at":"2026-05-25T05:36:40.299551Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":"2410.05459","doi":"10.48550/arxiv.2410.05459","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency.ArXiv, abs/2410.05459","venue":"arXiv (Cornell University)","work_id":"12dec1cc-a66a-4ef5-bfe7-061e33762461","year":2024},"citing_paper":{"arxiv_id":"2605.28600","last_updated":"2026-05-27T15:17:06Z","snapshot_observed_at":"2026-08-14T20:14:15.569517Z","submitted_at":"2026-05-27T15:17:06Z","title":"Transformers Provably Learn to Internalize Chain-of-Thought","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-06-29T14:29:10.010212Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2605.28600"},"observation_digest":"sha256:78c24e53b777f2711064dc371241369f514cc4e300e5b68f8100a63e57efeb50","observation_id":"eb03379a-a679-44aa-9a4b-782d9255ef4b","resolution":{"observed_at":"2026-06-29T14:33:30.594670Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":"2410.05459","doi":"10.48550/arxiv.2410.05459","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency.ArXiv, abs/2410.05459","venue":"arXiv (Cornell University)","work_id":"12dec1cc-a66a-4ef5-bfe7-061e33762461","year":2024},"citing_paper":{"arxiv_id":"2606.00183","last_updated":"2026-05-29T14:58:03Z","snapshot_observed_at":"2026-08-13T00:33:18.445193Z","submitted_at":"2026-05-29T14:58:03Z","title":"Agentic Transformers Provably Learn to Search via Reinforcement Learning","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-06-28T23:26:28.158991Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2606.00183"},"observation_digest":"sha256:236baf59c83e5645c24ba710b3e15e57e9f5a77de813b7a746170bb474938878","observation_id":"9d47d62b-7dbe-4f2f-90bf-83035e95c9f1","resolution":{"observed_at":"2026-06-28T23:42:49.906343Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2410.05459/citation-record","integrity":"/paper/2410.05459/integrity","json":"/paper/2410.05459/citation-record.json","paper":"/paper/2410.05459"},"outbound":[],"paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-16T13:11:45.819514Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 14 inbound Pith citation observations for arXiv:2410.05459."}