{"as_of":"2026-08-08T22:54:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:01d25360d53c7281e1f3e74821eb6a0fdbaeae0fc1f1413393ce1a75a189a5d1","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":34,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":34,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":34,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":34,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T19:12:42.327878Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":1,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2406.11794","last_updated":"2025-04-21T17:48:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-17T17:42:57Z","title":"DataComp-LM: In search of the next generation of training sets for language models","version":4},"reference_index":108,"source":"pdf_text","source_observed_at":"2026-05-17T22:58:16.523267Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2406.11794"},"observation_digest":"sha256:79bd7adf567a35339715a48c869116d9f6745ebcdc10178b37bff36a6c438d33","observation_id":"4145d778-536d-425b-96ee-0ee01afc7c07","resolution":{"observed_at":"2026-05-17T22:58:17.066499Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2406.18629","last_updated":"2024-06-26T17:43:06Z","snapshot_observed_at":"2026-08-06T00:24:52.274888Z","submitted_at":"2024-06-26T17:43:06Z","title":"Step-DPO: Step-wise Preference Optimization for Long-chain Reasoning of LLMs","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-18T23:58:29.040819Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2406.18629"},"observation_digest":"sha256:6ea0c0835b5b6c487a8c1088fb2e140746a32db9418f8f03531786f40322cd1d","observation_id":"4b075053-92a5-4a37-8f89-3f69aefd8a91","resolution":{"observed_at":"2026-05-18T23:58:29.135044Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-08T19:12:42.327878Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.05478","last_updated":"2025-02-08T07:38:45Z","snapshot_observed_at":"2026-08-08T19:07:07.183771Z","submitted_at":"2025-02-08T07:38:45Z","title":"OntoTune: Ontology-Driven Self-training for Aligning Large Language Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-08T19:12:42.327878Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2502.05478"},"observation_digest":"sha256:40a480bb0a600d6a1be419751c3c8c2ad9c72023e3fa55e9da0c9c12d9970910","observation_id":"eb7ab7cb-cc31-4da1-897a-1e135540ea71","resolution":{"observed_at":"2026-08-08T19:12:42.327878Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-08T11:01:03.397522Z","title":"Rho-1: Not all tokens are what you need","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10454","last_updated":"2025-08-22T07:58:49Z","snapshot_observed_at":"2026-08-08T10:55:47.232000Z","submitted_at":"2025-02-12T02:01:10Z","title":"One Example Shown, Many Concepts Known! Counterexample-Driven Conceptual Reasoning in Mathematical LLMs","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-08T11:01:03.397522Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2502.10454"},"observation_digest":"sha256:cee426b4e73d0760ff39d6096897db81d80cc725b044d8347f8a696d7c3c85e7","observation_id":"f0233435-99e1-408e-a4d0-e29fa6e04b4d","resolution":{"observed_at":"2026-08-08T11:01:03.397522Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2502.12272","last_updated":"2026-04-30T11:57:30Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-17T19:16:37Z","title":"Learning to Reason at the Frontier of Learnability","version":6},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-23T02:41:21.571824Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2502.12272"},"observation_digest":"sha256:de1374f55b6e498f825d1aa7815b6ad57cefa0ea2bd0c02dda8640c98ac8b232","observation_id":"95ee720f-f476-4f9e-b569-0f3e85e7c30c","resolution":{"observed_at":"2026-05-23T02:42:26.218182Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-07T14:09:47.959331Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19893","last_updated":"2025-05-26T12:23:26Z","snapshot_observed_at":"2026-08-07T14:01:57.359586Z","submitted_at":"2025-05-26T12:23:26Z","title":"ESLM: Risk-Averse Selective Language Modeling for Efficient Pretraining","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T14:09:47.959331Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2505.19893"},"observation_digest":"sha256:ba57596057370c4b7d9e0128a2d1624b91359606fd6f430e2814e0f4cb0cc5d9","observation_id":"e59e25b1-d230-4505-b4dd-5b381d0f9aad","resolution":{"observed_at":"2026-08-07T14:09:47.959331Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-07T12:48:50.753591Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.23540","last_updated":"2025-05-29T15:20:44Z","snapshot_observed_at":"2026-08-07T14:34:56.721276Z","submitted_at":"2025-05-29T15:20:44Z","title":"Probability-Consistent Preference Optimization for Enhanced LLM Reasoning","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T12:48:50.753591Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2505.23540"},"observation_digest":"sha256:13e1c028ddd9fbdbb0968595ce8d156d9feb5a108544a723adeb0c734776bfc2","observation_id":"4088bc24-c3d9-4787-8e90-48c51ab95d93","resolution":{"observed_at":"2026-08-07T12:48:50.753591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-07T12:03:34.041761Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.00876","last_updated":"2025-06-01T07:36:45Z","snapshot_observed_at":"2026-08-07T11:53:38.742479Z","submitted_at":"2025-06-01T07:36:45Z","title":"Not Every Token Needs Forgetting: Selective Unlearning to Limit Change in Utility in Large Language Model Unlearning","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T12:03:34.041761Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2506.00876"},"observation_digest":"sha256:ff1918937f504a8a9a3eb97134e0fe446e84a37289d5d3db754c40642bc46fb7","observation_id":"877563c6-5bb8-43ae-b7ed-0d986c0eaca0","resolution":{"observed_at":"2026-08-07T12:03:34.041761Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-07T11:19:46.493246Z","title":"Rho-1: Not all tokens are what you need.arXiv preprint arXiv:2404.07965, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.02867","last_updated":"2025-06-04T15:00:58Z","snapshot_observed_at":"2026-08-08T18:19:35.851376Z","submitted_at":"2025-06-03T13:31:10Z","title":"Demystifying Reasoning Dynamics with Mutual Information: Thinking Tokens are Information Peaks in LLM Reasoning","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:19:46.493246Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2506.02867"},"observation_digest":"sha256:90a7514c728d260752a60097000824bd7e34cbe5bdd5d7ed970821093db69a59","observation_id":"91758f7d-7bab-446c-a2b6-6158350f00ec","resolution":{"observed_at":"2026-08-07T11:19:46.493246Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-07T05:14:47.351148Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.08446","last_updated":"2025-06-10T04:44:28Z","snapshot_observed_at":"2026-08-07T09:15:20.058485Z","submitted_at":"2025-06-10T04:44:28Z","title":"A Survey on Large Language Models for Mathematical Reasoning","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T05:14:47.351148Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2506.08446"},"observation_digest":"sha256:78d3c53c0b01c4b9780a8f642d8c0da7c645529df6806308e8750deb8ec64687","observation_id":"403d8a32-ff29-4590-a22e-5aee1a14e224","resolution":{"observed_at":"2026-08-07T05:14:47.351148Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-07T00:42:00.976666Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13216","last_updated":"2025-06-16T08:16:03Z","snapshot_observed_at":"2026-08-07T00:34:23.901032Z","submitted_at":"2025-06-16T08:16:03Z","title":"Capability Salience Vector: Fine-grained Alignment of Loss and Capabilities for Downstream Task Scaling Law","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T00:42:00.976666Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2506.13216"},"observation_digest":"sha256:497ca7c45bd49d46352581ec89c05fa20dbecc14e06a2ada5ea93ddac327a3f0","observation_id":"7bba15c1-ce11-4a95-b04b-30fe549d7e98","resolution":{"observed_at":"2026-08-07T00:42:00.976666Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-07T00:15:34.227513Z","title":"Rho-1: Not all tokens are what you need.CoRR, abs/2404.07965, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.15021","last_updated":"2025-06-17T23:12:28Z","snapshot_observed_at":"2026-08-08T17:58:45.055043Z","submitted_at":"2025-06-17T23:12:28Z","title":"SFT-GO: Supervised Fine-Tuning with Group Optimization for Large Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T00:15:34.227513Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2506.15021"},"observation_digest":"sha256:c2800b4408e2615ca56d57f301d976cbc898b0ac8b4c395073f25e6caffdd2b6","observation_id":"97c019cf-d445-45f6-b6bb-9652cfd16282","resolution":{"observed_at":"2026-08-07T00:15:34.227513Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-06T20:21:51.821393Z","title":"Rho-1: Not all tokens are what you need","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.03253","last_updated":"2025-07-08T18:15:09Z","snapshot_observed_at":"2026-08-07T10:43:54.316543Z","submitted_at":"2025-07-04T02:19:58Z","title":"RefineX: Learning to Refine Pre-training Data at Scale from Expert-Guided Programs","version":2},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-06T20:21:51.821393Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2507.03253"},"observation_digest":"sha256:0ff509162a2550de88c4323e8d5d22aa8d67b0f5542bae78015d962df5b9bcff","observation_id":"be8874e2-b0a5-42da-9f9e-cffd68cd07d8","resolution":{"observed_at":"2026-08-06T20:21:51.821393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-06T20:18:15.688795Z","title":"Rho-1: Not all tokens are what you need, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.03327","last_updated":"2025-07-04T06:23:06Z","snapshot_observed_at":"2026-08-07T11:24:12.298323Z","submitted_at":"2025-07-04T06:23:06Z","title":"Read Quietly, Think Aloud: Decoupling Comprehension and Reasoning in LLMs","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-06T20:18:15.688795Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2507.03327"},"observation_digest":"sha256:a28a4fb71f182381f4d127f8978f233bfd0c9480c825dec9b59200ee332aed7e","observation_id":"06679556-0779-4f82-915d-41f4319bf107","resolution":{"observed_at":"2026-08-06T20:18:15.688795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-06T18:38:10.027182Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.07725","last_updated":"2025-07-10T12:58:45Z","snapshot_observed_at":"2026-08-07T16:33:04.277399Z","submitted_at":"2025-07-10T12:58:45Z","title":"Not All Preferences are What You Need for Post-Training: Selective Alignment Strategy for Preference Optimization","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-06T18:38:10.027182Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2507.07725"},"observation_digest":"sha256:2d714416ab47fac0cc1d16b104deb701812e9a2cbe264bc94a60913a433eaa3d","observation_id":"0f567fc3-240d-42c3-9de7-06984d22f8d7","resolution":{"observed_at":"2026-08-06T18:38:10.027182Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-06T16:53:11.243914Z","title":"Rho-1: Not all tokens are what you need","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.12466","last_updated":"2025-07-16T17:59:45Z","snapshot_observed_at":"2026-08-07T10:43:47.278118Z","submitted_at":"2025-07-16T17:59:45Z","title":"Language Models Improve When Pretraining Data Matches Target Tasks","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-06T16:53:11.243914Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2507.12466"},"observation_digest":"sha256:795c9619212d00a67dd5d32e8897deefc4e515803b65a98a3c04be7484497a87","observation_id":"a5c318d5-05a2-429f-a5d3-c4ee9a981f3b","resolution":{"observed_at":"2026-08-06T16:53:11.243914Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2507.21420","last_updated":"2026-04-28T19:54:45Z","snapshot_observed_at":"2026-08-06T19:44:38.919643Z","submitted_at":"2025-07-29T01:07:09Z","title":"ReGATE: Learning Faster and Better with Fewer Tokens in MLLMs","version":3},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-19T03:18:11.993413Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2507.21420"},"observation_digest":"sha256:f7b6be221a5bc888b5795a1cf049aa08aaa689fb0d4cbb528b43b50780d0dfc3","observation_id":"bbe78087-5486-4b00-9712-d0ca4cdbea02","resolution":{"observed_at":"2026-05-19T03:22:01.226237Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T22:36:43.649644Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.06800","last_updated":"2025-08-14T16:06:55Z","snapshot_observed_at":"2026-08-07T20:24:29.542648Z","submitted_at":"2025-08-09T03:10:56Z","title":"Hardness-Aware Dynamic Curriculum Learning for Robust Multimodal Emotion Recognition with Missing Modalities","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T22:36:43.649644Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2508.06800"},"observation_digest":"sha256:b44cc56a6cb86d85d91bc950a59b112196baa21736039b414f10276db154883c","observation_id":"6e4daea9-925e-4962-ac34-7143ff219f13","resolution":{"observed_at":"2026-08-05T22:36:43.649644Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2508.15229","last_updated":"2026-04-18T14:51:42Z","snapshot_observed_at":"2026-08-06T13:10:04.037804Z","submitted_at":"2025-08-21T04:32:13Z","title":"VocabTailor: Dynamic Vocabulary Selection for Downstream Tasks in Small Language Models","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-18T22:37:01.014698Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2508.15229"},"observation_digest":"sha256:1d18df629d439842d9a1579e8be2db7c9d5f3a7f5dfd7e568b500b116d22b13a","observation_id":"03963849-4861-4932-88bd-632a4a9077c3","resolution":{"observed_at":"2026-05-18T22:41:54.056264Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T05:46:28.436696Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.05056","last_updated":"2025-09-05T12:35:06Z","snapshot_observed_at":"2026-08-08T00:56:27.101486Z","submitted_at":"2025-09-05T12:35:06Z","title":"Masked Diffusion Language Models with Frequency-Informed Training","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-05T05:46:28.436696Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2509.05056"},"observation_digest":"sha256:913c7d593cdd26bda9d6d9f60133a3122832f1c27f0ec7008032a9858f3e4fdd","observation_id":"74600b74-d345-4bf1-970a-03f436b15788","resolution":{"observed_at":"2026-08-05T05:46:28.436696Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2509.12089","last_updated":"2026-04-30T05:20:05Z","snapshot_observed_at":"2026-07-06T22:29:58.066459Z","submitted_at":"2025-09-15T16:16:57Z","title":"RadarPLM: Adapting Pre-trained Language Models for Marine Radar Target Detection by Selective Fine-tuning","version":6},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-18T16:00:10.105362Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2509.12089"},"observation_digest":"sha256:c024f5744a150ffe83d4bde91de714245f4796b6a2808fe7fd7bb63afee62b42","observation_id":"65ede234-2401-4b0b-9e61-1067258d1c64","resolution":{"observed_at":"2026-05-18T16:01:34.637767Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2509.20863","last_updated":"2026-05-11T17:24:47Z","snapshot_observed_at":"2026-07-31T08:03:41.164839Z","submitted_at":"2025-09-25T07:55:58Z","title":"GIFT: Guided Importance-Aware Fine-Tuning for Diffusion Language Models","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-18T14:49:08.081740Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2509.20863"},"observation_digest":"sha256:75611a244f70f6292be1e13f96df5c6be3b0fb818a346c189992070d40d6f738","observation_id":"4e99f72d-dced-45aa-bcd4-20c497716d84","resolution":{"observed_at":"2026-05-18T14:51:30.219045Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2509.20863","last_updated":"2026-05-11T17:24:47Z","snapshot_observed_at":"2026-07-31T08:03:41.164839Z","submitted_at":"2025-09-25T07:55:58Z","title":"GIFT: Guided Importance-Aware Fine-Tuning for Diffusion Language Models","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-18T14:49:08.081740Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2509.20863"},"observation_digest":"sha256:fc7d9f8b49f937c150bbb582f45def34c9decc65d652784a12d54b0edc1d812b","observation_id":"c01e431c-1ded-4800-baaa-2d5664e4efa6","resolution":{"observed_at":"2026-05-18T14:51:29.829990Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2510.07962","last_updated":"2026-05-21T01:02:13Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-10-09T08:55:12Z","title":"LightReasoner: Can Small Language Models Teach Large Language Models Reasoning?","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-22T12:59:31.011283Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2510.07962"},"observation_digest":"sha256:9b8614c5f251b88fbe506a3207d9f3cfeaf1ad4e8b8aa0c3e0d12da79a1e7ec3","observation_id":"1b65ce9e-cae7-45ab-a332-37dfa695f3b2","resolution":{"observed_at":"2026-05-22T13:01:34.009777Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2601.10348","last_updated":"2026-05-21T06:29:24Z","snapshot_observed_at":"2026-08-03T01:57:39.427752Z","submitted_at":"2026-01-15T12:45:05Z","title":"Training-Trajectory-Aware Token Selection","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-22T11:41:21.275802Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2601.10348"},"observation_digest":"sha256:116bf2606cb859e021a34069e84822cba236aefaa55e8efdc6ee252f1c114b3e","observation_id":"73e7b422-da06-4a95-a442-7bb59e8b46fd","resolution":{"observed_at":"2026-05-22T11:41:29.587727Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2602.20816","last_updated":"2026-05-08T01:10:04Z","snapshot_observed_at":"2026-07-06T22:46:53.725781Z","submitted_at":"2026-02-24T11:54:06Z","title":"Don't Ignore the Tail: Decoupling top-K Probabilities for Efficient Language Model Distillation","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-15T20:09:47.931263Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2602.20816"},"observation_digest":"sha256:8b3d4b9674c606279b453db6742f37efd4e3c8d67a38003a880f6136da225325","observation_id":"2d28eb3c-3d01-473e-b40d-45259048bf83","resolution":{"observed_at":"2026-05-15T20:10:17.997841Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-02T21:05:25.711054Z","title":"Rho-1: Not all tokens are what you need.ArXiv, abs/2404.07965,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.21492","last_updated":"2026-07-18T04:56:29Z","snapshot_observed_at":"2026-08-04T06:06:14.689327Z","submitted_at":"2026-02-25T01:54:50Z","title":"GradAlign: Gradient-Aligned Data Selection for LLM Reinforcement Learning","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-02T21:05:25.711054Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2602.21492"},"observation_digest":"sha256:7c939cb9bb4c5fa6a6da3449761c68dfae766af20bd0197c3fdbbe74c2c97218","observation_id":"82a5af6b-56d8-4f88-a079-ff1dbdc67dd8","resolution":{"observed_at":"2026-08-02T21:05:25.711054Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2603.01097","last_updated":"2026-07-29T13:53:39Z","snapshot_observed_at":"2026-08-02T19:49:42.383713Z","submitted_at":"2026-03-01T13:28:57Z","title":"Understanding LoRA as Knowledge Memory: An Empirical Analysis","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-15T18:09:54.899039Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2603.01097"},"observation_digest":"sha256:a7db42e81867877e179d6d60531a998959cda6c3b6dd2382bd73d594df067cb5","observation_id":"241424ad-a14b-4432-8710-db8a3876f968","resolution":{"observed_at":"2026-05-15T18:10:13.256835Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-02T19:49:45.151728Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2603.01097","last_updated":"2026-07-29T13:53:39Z","snapshot_observed_at":"2026-08-02T19:49:42.383713Z","submitted_at":"2026-03-01T13:28:57Z","title":"Understanding LoRA as Knowledge Memory: An Empirical Analysis","version":5},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-02T19:49:45.151728Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2603.01097"},"observation_digest":"sha256:b546f253be9f1a697de1bcffc74462fda7bf97a824b71228115b24888be2cbf5","observation_id":"d7ca72a4-9fa5-4005-94ab-5957cc4f4893","resolution":{"observed_at":"2026-08-02T19:49:45.151728Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2604.17884","last_updated":"2026-04-20T06:55:26Z","snapshot_observed_at":"2026-07-06T23:04:55.190916Z","submitted_at":"2026-04-20T06:55:26Z","title":"SPREG: Structured Plan Repair with Entropy-Guided Test-Time Intervention for Large Language Model Reasoning","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-10T04:33:14.807721Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2604.17884"},"observation_digest":"sha256:265a869741148accc68f46a5c0d089696775ea20854b54a438334ee5080294b0","observation_id":"0494852a-730f-4f7d-bb1a-ea94a7b6420f","resolution":{"observed_at":"2026-05-11T11:51:02.797848Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2604.22374","last_updated":"2026-04-24T09:08:45Z","snapshot_observed_at":"2026-07-06T23:08:47.078992Z","submitted_at":"2026-04-24T09:08:45Z","title":"Selective Contrastive Learning For Gloss Free Sign Language Translation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-08T11:42:20.114297Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2604.22374"},"observation_digest":"sha256:c4aacddd1543ddf8ba1b733a29f7165f0d24b88205de9b13a17db3266ee00770","observation_id":"e0d41d1c-65d1-4bad-aa7c-9277cac97cd9","resolution":{"observed_at":"2026-05-11T19:31:10.178344Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":"2404.07965","doi":"10.48550/arxiv.2404.07965","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rho-1: Not all tokens are what you need","venue":"arXiv (Cornell University)","work_id":"74e3da93-f1b3-41df-bc52-f8092041fb58","year":2025},"citing_paper":{"arxiv_id":"2606.18286","last_updated":"2026-06-10T04:46:04Z","snapshot_observed_at":"2026-08-07T16:29:18.257991Z","submitted_at":"2026-06-10T04:46:04Z","title":"CODEBLOCK: Learning to Supervise Code at the Right Granularity","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-27T10:52:00.147391Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2606.18286"},"observation_digest":"sha256:005030555cdfef3cfb855700ef22e2cd676693ee7a3631593d5248c81240b95e","observation_id":"eb1877d1-0aec-459b-80dd-3a71fab9ea28","resolution":{"observed_at":"2026-07-03T08:17:45.472573Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-01T06:19:32.644746Z","title":"Rho-1: Not all tokens are what you need.arXiv preprint arXiv:2404.07965, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.22769","last_updated":"2026-07-24T03:54:58Z","snapshot_observed_at":"2026-08-05T05:05:03.937121Z","submitted_at":"2026-07-24T03:54:58Z","title":"DomainPilot: Domain-Level Loss-Guided Two-Stage Data Mixture Optimization for Efficient Language Model Fine-Tuning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-01T06:19:32.644746Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2607.22769"},"observation_digest":"sha256:231db3ad10bd4cabfd21b54d1c872ad6fd4c9e917616a620c5a0425cd019eb94","observation_id":"eef586f3-1521-4934-b262-7cb4f560ebc5","resolution":{"observed_at":"2026-08-01T06:19:32.644746Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07965","snapshot_observed_at":"2026-08-01T03:02:03.996810Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.25271","last_updated":"2026-07-28T04:18:49Z","snapshot_observed_at":"2026-08-08T08:44:07.330246Z","submitted_at":"2026-07-28T04:18:49Z","title":"Bridging Compute- and Data-Optimal Pretraining","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-01T03:02:03.996810Z"},"links":{"cited_paper":"/paper/2404.07965","citing_paper":"/paper/2607.25271"},"observation_digest":"sha256:a3c005acc9357c91d8966dc3b0f402e17bbd08b506f88eb327e3f9f4f56deaae","observation_id":"0de724b4-9a24-409a-8695-89c258a522de","resolution":{"observed_at":"2026-08-01T03:02:03.996810Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2404.07965/citation-record","integrity":"/paper/2404.07965/integrity","json":"/paper/2404.07965/citation-record.json","paper":"/paper/2404.07965"},"outbound":[],"paper":{"arxiv_id":"2404.07965","last_updated":"2025-01-08T09:07:54Z","latest_version":4,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T08:00:00.868748Z","submitted_at":"2024-04-11T17:52:01Z","title":"Rho-1: Not All Tokens Are What You Need"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 34 inbound Pith citation observations for arXiv:2404.07965."}