{"as_of":"2026-08-09T01:09:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7b686f3b00570e6d227c7e7142fbd9e665bb86076e3c4e97bf617e8488beb8c4","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":8,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":8,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":8,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":8,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T00:14:40.968384Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T03:46:32.820784Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2402.00559","last_updated":"2024-05-21T09:55:07Z","snapshot_observed_at":"2026-08-06T18:38:04.978620Z","submitted_at":"2024-02-01T12:46:45Z","title":"A Chain-of-Thought Is as Strong as Its Weakest Link: A Benchmark for Verifiers of Reasoning Chains","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.00559","snapshot_observed_at":"2026-08-07T00:14:40.968384Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14948","last_updated":"2025-06-17T19:59:44Z","snapshot_observed_at":"2026-08-08T14:39:17.578977Z","submitted_at":"2025-06-17T19:59:44Z","title":"Structured Moral Reasoning in Language Models: A Value-Grounded Evaluation Framework","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T00:14:40.968384Z"},"links":{"cited_paper":"/paper/2402.00559","citing_paper":"/paper/2506.14948"},"observation_digest":"sha256:0fb485b007e4704f9e36c6fce3e7114aa4403000968713f926d0977e597dc7e3","observation_id":"8c102ed7-f62d-4f85-8c40-b5e3defca966","resolution":{"observed_at":"2026-08-07T00:14:40.968384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.00559","last_updated":"2024-05-21T09:55:07Z","snapshot_observed_at":"2026-08-06T18:38:04.978620Z","submitted_at":"2024-02-01T12:46:45Z","title":"A Chain-of-Thought Is as Strong as Its Weakest Link: A Benchmark for Verifiers of Reasoning Chains","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.00559","snapshot_observed_at":"2026-08-06T21:58:19.835890Z","title":"A chain-of-thought is as strong as its weakest link: A benchmark for verifiers of reasoning chains,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.23014","last_updated":"2025-06-28T20:55:21Z","snapshot_observed_at":"2026-08-08T23:47:54.278328Z","submitted_at":"2025-06-28T20:55:21Z","title":"Generating Privacy Stories From Software Documentation","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T21:58:19.835890Z"},"links":{"cited_paper":"/paper/2402.00559","citing_paper":"/paper/2506.23014"},"observation_digest":"sha256:79211ce1be1c9d0e512184fddcac1817605848094eb104f0f23dd797b6c00131","observation_id":"6d744290-3de8-4d6c-ae28-5e8e8cb3e1ab","resolution":{"observed_at":"2026-08-06T21:58:19.835890Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.00559","last_updated":"2024-05-21T09:55:07Z","snapshot_observed_at":"2026-08-06T18:38:04.978620Z","submitted_at":"2024-02-01T12:46:45Z","title":"A Chain-of-Thought Is as Strong as Its Weakest Link: A Benchmark for Verifiers of Reasoning Chains","version":4},"cited_work":{"arxiv_id":"2402.00559","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.00559","snapshot_observed_at":"2026-07-02T03:46:32.820784Z","title":"A chain-of-thought is as strong as its weakest link: A bench- mark for verifiers of reasoning chains.arXiv preprint arXiv:2402.00559","venue":null,"work_id":"4c393dd1-8cb5-425d-ad5e-2c040cb9b4f6","year":2024},"citing_paper":{"arxiv_id":"2602.08324","last_updated":"2026-05-25T15:41:17Z","snapshot_observed_at":"2026-08-03T03:25:21.079471Z","submitted_at":"2026-02-09T06:57:15Z","title":"Towards Efficient Large Language Reasoning Models via Extreme-Ratio Chain-of-Thought Compression","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-21T13:43:51.127429Z"},"links":{"cited_paper":"/paper/2402.00559","citing_paper":"/paper/2602.08324"},"observation_digest":"sha256:11f00906f23a02906e50db093b6ae2b9c9c360a3cae88273ee3490c54a87c286","observation_id":"b90000b6-deb4-4a87-b4e1-cf31a9450b97","resolution":{"observed_at":"2026-05-21T13:44:11.400939Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.00559","last_updated":"2024-05-21T09:55:07Z","snapshot_observed_at":"2026-08-06T18:38:04.978620Z","submitted_at":"2024-02-01T12:46:45Z","title":"A Chain-of-Thought Is as Strong as Its Weakest Link: A Benchmark for Verifiers of Reasoning Chains","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.00559","snapshot_observed_at":"2026-08-03T03:25:23.069996Z","title":"A chain-of-thought is as strong as its weakest link: A bench- mark for verifiers of reasoning chains.arXiv preprint arXiv:2402.00559,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.08324","last_updated":"2026-05-25T15:41:17Z","snapshot_observed_at":"2026-08-03T03:25:21.079471Z","submitted_at":"2026-02-09T06:57:15Z","title":"Towards Efficient Large Language Reasoning Models via Extreme-Ratio Chain-of-Thought Compression","version":4},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-03T03:25:23.069996Z"},"links":{"cited_paper":"/paper/2402.00559","citing_paper":"/paper/2602.08324"},"observation_digest":"sha256:2ccd8333719c2a40f074fbee1d571478420e1869dd8c134bc77b6a3f152a03c3","observation_id":"2ab171ee-c6a5-40d3-bd4d-5fb41b370f7b","resolution":{"observed_at":"2026-08-03T03:25:23.069996Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.00559","last_updated":"2024-05-21T09:55:07Z","snapshot_observed_at":"2026-08-06T18:38:04.978620Z","submitted_at":"2024-02-01T12:46:45Z","title":"A Chain-of-Thought Is as Strong as Its Weakest Link: A Benchmark for Verifiers of Reasoning Chains","version":4},"cited_work":{"arxiv_id":"2402.00559","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.00559","snapshot_observed_at":"2026-07-02T03:46:32.820784Z","title":"A chain-of-thought is as strong as its weakest link: A bench- mark for verifiers of reasoning chains.arXiv preprint arXiv:2402.00559","venue":null,"work_id":"4c393dd1-8cb5-425d-ad5e-2c040cb9b4f6","year":2024},"citing_paper":{"arxiv_id":"2604.04942","last_updated":"2026-03-13T13:01:01Z","snapshot_observed_at":"2026-07-06T22:53:46.999911Z","submitted_at":"2026-03-13T13:01:01Z","title":"TDA-RC: Task-Driven Alignment for Knowledge-Based Reasoning Chains in Large Language Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-15T12:21:39.237267Z"},"links":{"cited_paper":"/paper/2402.00559","citing_paper":"/paper/2604.04942"},"observation_digest":"sha256:58c2a6f552adabc35f24927622f7ce03f5ed502f77fa339d2716a2c01a030add","observation_id":"a706d2d4-bdc5-4061-b62f-b4fd279a98bc","resolution":{"observed_at":"2026-05-15T12:25:35.986984Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.00559","last_updated":"2024-05-21T09:55:07Z","snapshot_observed_at":"2026-08-06T18:38:04.978620Z","submitted_at":"2024-02-01T12:46:45Z","title":"A Chain-of-Thought Is as Strong as Its Weakest Link: A Benchmark for Verifiers of Reasoning Chains","version":4},"cited_work":{"arxiv_id":"2402.00559","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.00559","snapshot_observed_at":"2026-07-02T03:46:32.820784Z","title":"A chain-of-thought is as strong as its weakest link: A bench- mark for verifiers of reasoning chains.arXiv preprint arXiv:2402.00559","venue":null,"work_id":"4c393dd1-8cb5-425d-ad5e-2c040cb9b4f6","year":2024},"citing_paper":{"arxiv_id":"2604.15727","last_updated":"2026-04-17T05:59:16Z","snapshot_observed_at":"2026-07-06T23:03:12.000381Z","submitted_at":"2026-04-17T05:59:16Z","title":"Structured Abductive-Deductive-Inductive Reasoning for LLMs via Algebraic Invariants","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T08:45:32.431012Z"},"links":{"cited_paper":"/paper/2402.00559","citing_paper":"/paper/2604.15727"},"observation_digest":"sha256:323a2e883731009c379e210608caaaafb2d9fa458b1e24467eee4a384d72a77b","observation_id":"bdd455e0-3dd4-4f35-a9c6-b568421ebf69","resolution":{"observed_at":"2026-05-10T08:48:01.897400Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.00559","last_updated":"2024-05-21T09:55:07Z","snapshot_observed_at":"2026-08-06T18:38:04.978620Z","submitted_at":"2024-02-01T12:46:45Z","title":"A Chain-of-Thought Is as Strong as Its Weakest Link: A Benchmark for Verifiers of Reasoning Chains","version":4},"cited_work":{"arxiv_id":"2402.00559","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.00559","snapshot_observed_at":"2026-07-02T03:46:32.820784Z","title":"A chain-of-thought is as strong as its weakest link: A bench- mark for verifiers of reasoning chains.arXiv preprint arXiv:2402.00559","venue":null,"work_id":"4c393dd1-8cb5-425d-ad5e-2c040cb9b4f6","year":2024},"citing_paper":{"arxiv_id":"2605.19228","last_updated":"2026-06-07T18:34:11Z","snapshot_observed_at":"2026-07-31T16:56:03.737552Z","submitted_at":"2026-05-19T00:57:51Z","title":"Diagnosing Multi-step Reasoning Failures in Black-box LLMs via Stepwise Confidence Attribution","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-05-20T06:43:23.870127Z"},"links":{"cited_paper":"/paper/2402.00559","citing_paper":"/paper/2605.19228"},"observation_digest":"sha256:6e7f62c9a70344d98167438382df23b0d8dcf6237dc6d4414eea14571fab5be0","observation_id":"77b81608-d64f-4a38-b57a-df8cde6bda58","resolution":{"observed_at":"2026-05-20T06:48:06.036759Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.00559","last_updated":"2024-05-21T09:55:07Z","snapshot_observed_at":"2026-08-06T18:38:04.978620Z","submitted_at":"2024-02-01T12:46:45Z","title":"A Chain-of-Thought Is as Strong as Its Weakest Link: A Benchmark for Verifiers of Reasoning Chains","version":4},"cited_work":{"arxiv_id":"2402.00559","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.00559","snapshot_observed_at":"2026-07-02T03:46:32.820784Z","title":"A chain-of-thought is as strong as its weakest link: A bench- mark for verifiers of reasoning chains.arXiv preprint arXiv:2402.00559","venue":null,"work_id":"4c393dd1-8cb5-425d-ad5e-2c040cb9b4f6","year":2024},"citing_paper":{"arxiv_id":"2606.03660","last_updated":"2026-06-03T14:05:21Z","snapshot_observed_at":"2026-07-06T23:43:55.632508Z","submitted_at":"2026-06-02T13:47:19Z","title":"From Answers to States: Verifiable Process-Level Evaluation of Chemical Reasoning in Large Language Models","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-06-28T09:42:02.298729Z"},"links":{"cited_paper":"/paper/2402.00559","citing_paper":"/paper/2606.03660"},"observation_digest":"sha256:53432378d5b2e63e1d7737d7a03e695d737f97d5aeaab97bd62a8e3926220c3b","observation_id":"2595aa8d-1fe3-4b4a-ac66-c017c36b2a80","resolution":{"observed_at":"2026-07-02T03:46:32.822252Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2402.00559/citation-record","integrity":"/paper/2402.00559/integrity","json":"/paper/2402.00559/citation-record.json","paper":"/paper/2402.00559"},"outbound":[],"paper":{"arxiv_id":"2402.00559","last_updated":"2024-05-21T09:55:07Z","latest_version":4,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-06T18:38:04.978620Z","submitted_at":"2024-02-01T12:46:45Z","title":"A Chain-of-Thought Is as Strong as Its Weakest Link: A Benchmark for Verifiers of Reasoning Chains"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 8 inbound Pith citation observations for arXiv:2402.00559."}