{"as_of":"2026-08-10T14:22:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:38d7e9ca75be3780e380089e92468878c62c80308cdd419b5a0125d8a2e5bd26","coverage":[{"denominator":17,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":17,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T04:52:56.078724Z","state":"measured"},{"denominator":20,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":20,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-01T01:59:49.794527Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-18T12:46:24.236751Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"cited_work":{"arxiv_id":"2506.09428","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.09428","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Improved supervised fine-tuning for large language models to mitigate catastrophic forgetting","venue":null,"work_id":"2134b0cf-13f3-4707-b051-6cc3168088ab","year":2025},"citing_paper":{"arxiv_id":"2509.23629","last_updated":"2026-05-07T02:29:03Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-09-28T04:10:37Z","title":"Emergent Slow Thinking in LLMs as Inverse Tree Freezing","version":3},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-18T12:43:49.628082Z"},"links":{"cited_paper":"/paper/2506.09428","citing_paper":"/paper/2509.23629"},"observation_digest":"sha256:72abcaac1997df3afd0ff7dd988cd859be6bf45f3299b4adb1f8a05235ecf7dc","observation_id":"b7b5b330-6b44-4911-81f9-2d2ca43417c1","resolution":{"observed_at":"2026-05-18T12:46:24.240283Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"cited_work":{"arxiv_id":"2506.09428","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.09428","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Improved supervised fine-tuning for large language models to mitigate catastrophic forgetting","venue":null,"work_id":"2134b0cf-13f3-4707-b051-6cc3168088ab","year":2025},"citing_paper":{"arxiv_id":"2605.06632","last_updated":"2026-05-07T17:44:07Z","snapshot_observed_at":"2026-07-30T07:26:22.311737Z","submitted_at":"2026-05-07T17:44:07Z","title":"Crafting Reversible SFT Behaviors in Large Language Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-08T12:16:16.104176Z"},"links":{"cited_paper":"/paper/2506.09428","citing_paper":"/paper/2605.06632"},"observation_digest":"sha256:647c1b275ad3357a1b78942a346e726421fbdc64919744679371939a86c3c768","observation_id":"1991c15c-c5da-44ed-bfc2-23ccbf076010","resolution":{"observed_at":"2026-05-11T19:21:07.320848Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.09428","snapshot_observed_at":"2026-08-01T01:59:49.794527Z","title":"Improved supervised fine-tuning for large language models to mitigate catastrophic forgetting.arXiv preprint arXiv:2506.09428,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.25614","last_updated":"2026-07-28T11:45:34Z","snapshot_observed_at":"2026-08-09T01:33:49.293526Z","submitted_at":"2026-07-28T11:45:34Z","title":"MemSFT: Mitigating Alignment Tax with an External Parametric Memory","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-01T01:59:49.794527Z"},"links":{"cited_paper":"/paper/2506.09428","citing_paper":"/paper/2607.25614"},"observation_digest":"sha256:40cd5984363b498abd4ee26d5f64a0b46e4d2eeaf5898c3920790400cb8d5349","observation_id":"66302a0f-e1f9-48af-8a1f-3809a69c10bb","resolution":{"observed_at":"2026-08-01T01:59:49.794527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2506.09428/citation-record","integrity":"/paper/2506.09428/integrity","json":"/paper/2506.09428/citation-record.json","paper":"/paper/2506.09428"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2311.10702","last_updated":"2023-11-20T02:01:33Z","snapshot_observed_at":"2026-08-09T03:30:59.643714Z","submitted_at":"2023-11-17T18:45:45Z","title":"Camels in a Changing Climate: Enhancing LM Adaptation with Tulu 2","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.10702","snapshot_observed_at":"2026-08-07T04:52:56.022809Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.022809Z"},"links":{"cited_paper":"/paper/2311.10702","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:b21f7f286cef9e8bb1cf9b32e4d709f1243c6ec6ff2d80019783c6ae2244afbd","observation_id":"d9983265-7526-4dd2-a7ca-a2e51a661b51","resolution":{"observed_at":"2026-08-07T04:52:56.022809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2103.03874","last_updated":"2021-11-08T21:30:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-03-05T18:59:39Z","title":"Measuring Mathematical Problem Solving With the MATH Dataset","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.03874","snapshot_observed_at":"2026-08-07T04:52:56.012378Z","title":"arXiv preprint arXiv:2103.03874","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.012378Z"},"links":{"cited_paper":"/paper/2103.03874","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:7b5d04853d003c0b8dcc39b5b3c296c3f4a622adaf7d1785e720a90d19bf1b36","observation_id":"be9a094f-9d7d-4544-bad3-c095b5a69b51","resolution":{"observed_at":"2026-08-07T04:52:56.012378Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.01244","last_updated":"2024-05-25T12:17:29Z","snapshot_observed_at":"2026-08-10T06:04:30.032752Z","submitted_at":"2024-03-02T16:11:23Z","title":"Mitigating Catastrophic Forgetting in Large Language Models with Self-Synthesized Rehearsal","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.01244","snapshot_observed_at":"2026-08-07T04:52:56.017505Z","title":"arXiv preprint arXiv:2403.01244","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.017505Z"},"links":{"cited_paper":"/paper/2403.01244","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:d92c55cf18b1604de469d8b5c614f5c77a025e7347baae495060d6b876835d97","observation_id":"e336d970-80ff-45f5-bbbb-71b308a37085","resolution":{"observed_at":"2026-08-07T04:52:56.017505Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06825","last_updated":"2023-10-10T17:54:58Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-10T17:54:58Z","title":"Mistral 7B","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06825","snapshot_observed_at":"2026-08-07T04:52:56.028628Z","title":"arXiv preprint arXiv:2310.06825","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.028628Z"},"links":{"cited_paper":"/paper/2310.06825","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:9e9033384119267cefb7bd899e2c515c6699235a09b47227737a3e7feca3ce68","observation_id":"f86c8349-dc2f-4ae2-8ffb-112ed76d5de9","resolution":{"observed_at":"2026-08-07T04:52:56.028628Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2110.02600","last_updated":"2022-02-28T10:18:17Z","snapshot_observed_at":"2026-08-10T04:07:26.683436Z","submitted_at":"2021-10-06T09:10:10Z","title":"Sequential Reptile: Inter-Task Gradient Alignment for Multilingual Learning","version":3},"cited_work":{"arxiv_id":"2110.02600","doi":null,"metadata_source":"pith","pith_arxiv_id":"2110.02600","snapshot_observed_at":"2026-08-07T04:52:56.261583Z","title":"Sequential Reptile: Inter-Task Gradient Alignment for Multilingual Learning","venue":"cs.CL","work_id":"b5dffa63-e85a-4480-9ab8-77fc1c32475d","year":2021},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.034222Z"},"links":{"cited_paper":"/paper/2110.02600","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:0cd08461f3751d703abfc4df90d73e20dd6f05f0edfb4a5c23be3e02075d1921","observation_id":"982c1fff-25c3-4e19-9947-68e556012a58","resolution":{"observed_at":"2026-08-07T04:52:56.269150Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14734","last_updated":"2024-11-01T20:05:19Z","snapshot_observed_at":"2026-08-03T02:02:15.325394Z","submitted_at":"2024-05-23T16:01:46Z","title":"SimPO: Simple Preference Optimization with a Reference-Free Reward","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.14734","snapshot_observed_at":"2026-08-07T04:52:56.039373Z","title":"arXiv preprint arXiv:2405.14734","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.039373Z"},"links":{"cited_paper":"/paper/2405.14734","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:a2966b21e55c3695d871b1e223aa6165ad2c7f71a35a717c4801a7d5e2d61645","observation_id":"50ebd32b-b888-4368-9177-dd1c1602f3a5","resolution":{"observed_at":"2026-08-07T04:52:56.039373Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.01116","last_updated":"2023-06-01T20:03:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-01T20:03:56Z","title":"The RefinedWeb Dataset for Falcon LLM: Outperforming Curated Corpora with Web Data, and Web Data Only","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.01116","snapshot_observed_at":"2026-08-07T04:52:56.049249Z","title":"arXiv preprint arXiv:2306.01116","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.049249Z"},"links":{"cited_paper":"/paper/2306.01116","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:d57dfd8e3e7a40f6e86133c58986548384c528837d38df120b2eb2183c747d95","observation_id":"9ee831b8-9e8a-432b-80bb-144a89467036","resolution":{"observed_at":"2026-08-07T04:52:56.049249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12022","last_updated":"2023-11-20T18:57:34Z","snapshot_observed_at":"2026-08-10T12:02:35.919497Z","submitted_at":"2023-11-20T18:57:34Z","title":"GPQA: A Graduate-Level Google-Proof Q&A Benchmark","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.12022","snapshot_observed_at":"2026-08-07T04:52:56.054436Z","title":"arXiv preprint arXiv:2311.12022","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.054436Z"},"links":{"cited_paper":"/paper/2311.12022","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:726874212b28d5c9d459e1844f104b504404b114136a305132ed6f9e3ee1d5d4","observation_id":"7fe85f42-63fb-43ba-a000-27e83d22f8a1","resolution":{"observed_at":"2026-08-07T04:52:56.054436Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.18865","last_updated":"2024-02-29T05:27:45Z","snapshot_observed_at":"2026-08-04T00:58:13.854277Z","submitted_at":"2024-02-29T05:27:45Z","title":"Analyzing and Reducing Catastrophic Forgetting in Parameter Efficient Tuning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.18865","snapshot_observed_at":"2026-08-07T04:52:56.059943Z","title":"arXiv preprint arXiv:2402.18865","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.059943Z"},"links":{"cited_paper":"/paper/2402.18865","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:931adb438c2bb4efad4dbdd14cf01a42ada4657084e6e8e28ccb79f650c3fbd4","observation_id":"c6c7f2a2-3a73-44ca-950d-0ce514475618","resolution":{"observed_at":"2026-08-07T04:52:56.059943Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01574","last_updated":"2024-11-06T02:54:00Z","snapshot_observed_at":"2026-08-06T00:29:17.674418Z","submitted_at":"2024-06-03T17:53:00Z","title":"MMLU-Pro: A More Robust and Challenging Multi-Task Language Understanding Benchmark","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01574","snapshot_observed_at":"2026-08-07T04:52:56.069336Z","title":"arXiv preprint arXiv:2406.01574","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.069336Z"},"links":{"cited_paper":"/paper/2406.01574","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:5a7377607869be3eff48c6c9a4a3f8d1b554f836fdd4b978c52e3501cb190bd0","observation_id":"6e830037-89fb-456b-986d-2cec45d0eb6c","resolution":{"observed_at":"2026-08-07T04:52:56.069336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.12244","last_updated":"2025-05-27T06:49:09Z","snapshot_observed_at":"2026-08-10T07:58:15.209421Z","submitted_at":"2023-04-24T16:31:06Z","title":"WizardLM: Empowering large pre-trained language models to follow complex instructions","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.12244","snapshot_observed_at":"2026-08-07T04:52:56.073990Z","title":"arXiv preprint arXiv:2304.12244","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.073990Z"},"links":{"cited_paper":"/paper/2304.12244","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:025e2c05420595ebd38b507d18af0d0109f56c1607a70336bf58edc605ee61a3","observation_id":"dbc8cea8-3982-45ee-979b-2caf821b3639","resolution":{"observed_at":"2026-08-07T04:52:56.073990Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.07911","last_updated":"2023-11-14T05:13:55Z","snapshot_observed_at":"2026-07-06T16:47:08.877195Z","submitted_at":"2023-11-14T05:13:55Z","title":"Instruction-Following Evaluation for Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.07911","snapshot_observed_at":"2026-08-07T04:52:56.078724Z","title":"arXiv preprint arXiv:2311.07911","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.078724Z"},"links":{"cited_paper":"/paper/2311.07911","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:ebb129a1e2e3290179d09c88d2f4f683982ffdf9b6832f1205716f8ef51da6a6","observation_id":"5ac1d499-a016-4a3f-a3db-5a51a661d15e","resolution":{"observed_at":"2026-08-07T04:52:56.078724Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-07T04:52:56.064629Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.064629Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:ae7262a1760a629d399641ab752b0127423589d4877c72a33354e0b83aabb196","observation_id":"c2a3de47-eacf-4f3a-ada4-20aefe873105","resolution":{"observed_at":"2026-08-07T04:52:56.064629Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.00910","last_updated":"2020-10-02T10:32:29Z","snapshot_observed_at":"2026-08-02T22:56:05.822679Z","submitted_at":"2020-10-02T10:32:29Z","title":"Continual Learning for Natural Language Generation in Task-oriented Dialog Systems","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.00910","snapshot_observed_at":"2026-08-07T04:52:56.044301Z","title":"arXiv preprint arXiv:2010.00910","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.044301Z"},"links":{"cited_paper":"/paper/2010.00910","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:4ebbfd96d512c3906c9ce16dd6e7885c91c0296af0e78ff190a1d8934e97ffbf","observation_id":"2022c9d3-f28a-43f0-b226-0b7aa8035969","resolution":{"observed_at":"2026-08-07T04:52:56.044301Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.08173","last_updated":"2021-07-17T03:20:38Z","snapshot_observed_at":"2026-07-06T11:29:56.695924Z","submitted_at":"2021-07-17T03:20:38Z","title":"Continual Learning for Task-oriented Dialogue System with Iterative Network Pruning, Expanding and Masking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.08173","snapshot_observed_at":"2026-08-07T04:52:56.006926Z","title":"arXiv preprint arXiv:2107.08173","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.006926Z"},"links":{"cited_paper":"/paper/2107.08173","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:51daf384c79948c48f17f9d732ca03bf1fe8a5a95c85c4d201545e308bd4a075","observation_id":"fe01d1e1-f818-49c1-addc-d3bd040e3d45","resolution":{"observed_at":"2026-08-07T04:52:56.006926Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14233","last_updated":"2023-05-23T16:49:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-05-23T16:49:14Z","title":"Enhancing Chat Language Models by Scaling High-quality Instructional Conversations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14233","snapshot_observed_at":"2026-08-07T04:52:56.001442Z","title":"arXiv preprint arXiv:2305.14233","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:56.001442Z"},"links":{"cited_paper":"/paper/2305.14233","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:b046680368d49d4f85879a36d86c74ff538debd2c57d0dac2d7e52afb8b969af","observation_id":"6adee60d-47ea-46da-9f8f-d9ad82bcceb6","resolution":{"observed_at":"2026-08-07T04:52:56.001442Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10323","last_updated":"2024-06-14T17:44:08Z","snapshot_observed_at":"2026-07-06T18:31:13.039726Z","submitted_at":"2024-06-14T17:44:08Z","title":"GenQA: Generating Millions of Instructions from a Handful of Prompts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10323","snapshot_observed_at":"2026-08-07T04:52:55.994940Z","title":"arXiv preprint arXiv:2406.10323","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-07T04:52:55.994940Z"},"links":{"cited_paper":"/paper/2406.10323","citing_paper":"/paper/2506.09428"},"observation_digest":"sha256:1e8db2dea8084f7e43bcdf2c379d600ece3ce8839bc88a221c304e1704d88588","observation_id":"2fda7eb7-1f61-4008-b791-ef709c12bf72","resolution":{"observed_at":"2026-08-07T04:52:55.994940Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.09428","last_updated":"2025-06-28T02:26:03Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-09T03:32:07.542600Z","submitted_at":"2025-06-11T06:23:50Z","title":"Improved Supervised Fine-Tuning for Large Language Models to Mitigate Catastrophic Forgetting"},"reference_resolution":{"displayed":17,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":16,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":17},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 17 of 17 outbound references and 3 inbound Pith citation observations for arXiv:2506.09428."}