{"as_of":"2026-08-14T11:58:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:08ce29d50e9334ece9edeba9b08b7d6837767b90315fc6866258c9e2fad38907","coverage":[{"denominator":43,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":43,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:00:33.402147Z","state":"measured"},{"denominator":50,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":50,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":7,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":7,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T20:00:56.113141Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-06T20:00:56.113141Z","title":"arXiv preprint arXiv:2505.20416","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.04009","last_updated":"2025-07-05T11:38:59Z","snapshot_observed_at":"2026-08-12T19:19:40.791635Z","submitted_at":"2025-07-05T11:38:59Z","title":"Easy Dataset: A Unified and Extensible Framework for Synthesizing LLM Fine-Tuning Data from Unstructured Documents","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-06T20:00:56.113141Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2507.04009"},"observation_digest":"sha256:a107148db980b72f40ba243edf32f8b919026983393a75ce0dd8741df47699bd","observation_id":"dde05f2d-f7b4-4d0d-9660-c04418cba721","resolution":{"observed_at":"2026-08-06T20:00:56.113141Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2602.12705","last_updated":"2026-04-07T11:35:36Z","snapshot_observed_at":"2026-07-06T22:45:44.135338Z","submitted_at":"2026-02-13T08:19:38Z","title":"MedXIAOHE: A Comprehensive Recipe for Building Medical MLLMs","version":4},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-15T22:52:30.992054Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2602.12705"},"observation_digest":"sha256:f01386822fdf101cf73172e10098f6476fe63481c1b60d1a3797c1a6f3b3f98f","observation_id":"541e15e3-fab2-4269-9ca0-a4bc0a88c7e0","resolution":{"observed_at":"2026-05-15T22:56:50.406292Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2605.08709","last_updated":"2026-05-09T05:44:29Z","snapshot_observed_at":"2026-07-06T23:20:52.521341Z","submitted_at":"2026-05-09T05:44:29Z","title":"UniShield: Unified Face Attack Detection via KG-Informed Multimodal Reasoning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-12T01:10:21.661074Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2605.08709"},"observation_digest":"sha256:07f97976aff506dea65b8a36248cfcf0daf84f885ecf07f9a188b8c016eecc72","observation_id":"a015fdf5-7c8f-43ca-92e5-4f94a3fb8400","resolution":{"observed_at":"2026-05-12T08:26:24.434798Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2605.18025","last_updated":"2026-05-18T08:14:49Z","snapshot_observed_at":"2026-08-12T16:03:04.130920Z","submitted_at":"2026-05-18T08:14:49Z","title":"TeleCom-Bench: How Far Are Large Language Models from Industrial Telecommunication Applications?","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-20T11:23:34.256279Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2605.18025"},"observation_digest":"sha256:c1c35a8dd3b196e598ffb1a81c88240cd387911ac68e80075ce2e72d9e00e765","observation_id":"33337c6f-9e3a-452c-a134-5004acaa7301","resolution":{"observed_at":"2026-05-20T11:28:14.626268Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2605.19394","last_updated":"2026-05-19T05:40:12Z","snapshot_observed_at":"2026-08-02T23:18:14.648216Z","submitted_at":"2026-05-19T05:40:12Z","title":"EmbGen: Teaching with Reassembled Corpora","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-20T06:34:39.739666Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2605.19394"},"observation_digest":"sha256:202e7657343b83e8d0fcbf1aa0cd289b5b4cd6fe345977c068e78ff5736fd5a7","observation_id":"9337b38c-61da-4dc5-be05-4cee59c5f96c","resolution":{"observed_at":"2026-05-20T06:38:05.530130Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2606.12087","last_updated":"2026-06-10T13:49:11Z","snapshot_observed_at":"2026-08-02T15:03:50.335992Z","submitted_at":"2026-06-10T13:49:11Z","title":"FORT-Searcher: Synthesizing Shortcut-Resistant Search Tasks for Training Deep Search Agents","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-06-27T10:01:45.332920Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2606.12087"},"observation_digest":"sha256:2e4f97600706bbec45cc78d5da4664e9ee8273a110e96cabf7f8b135e841b8a8","observation_id":"88e06038-75e0-43e4-a20f-4a9faf677921","resolution":{"observed_at":"2026-07-03T10:27:56.494259Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2606.12837","last_updated":"2026-06-17T03:34:47Z","snapshot_observed_at":"2026-08-12T14:26:53.329770Z","submitted_at":"2026-06-11T03:04:32Z","title":"LoHoSearch: Benchmarking Long-Horizon Search Agents Beyond the Human Difficulty Ceiling","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-06-27T07:02:40.293652Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2606.12837"},"observation_digest":"sha256:a82e3f0e939789a5d0b12ead2e4ebf40b53a610621f6b5d36f1a2fd562793c7d","observation_id":"baccde30-33c1-41f3-b1e9-c690b5ce62f0","resolution":{"observed_at":"2026-07-03T14:28:31.706678Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.20416/citation-record","integrity":"/paper/2505.20416/integrity","json":"/paper/2505.20416/citation-record.json","paper":"/paper/2505.20416"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2407.05015","last_updated":"2024-07-06T09:10:05Z","snapshot_observed_at":"2026-08-12T23:27:42.652308Z","submitted_at":"2024-07-06T09:10:05Z","title":"How do you know that? Teaching Generative Language Models to Reference Answers to Biomedical Questions","version":1},"cited_work":{"arxiv_id":"2407.05015","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.05015","snapshot_observed_at":"2026-08-07T14:00:34.351519Z","title":"How do you know that? Teaching Generative Language Models to Reference Answers to Biomedical Questions","venue":"cs.CL","work_id":"c41d0524-5d49-4f4f-8e16-5f6a7e45925b","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:29.497217Z"},"links":{"cited_paper":"/paper/2407.05015","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:940a86a05b1a296bcbfe804283c51bb02ff79c1ca6e9731a6e40001ce0ae7297","observation_id":"5de4fea8-03f8-4057-8ed6-db863ca1c556","resolution":{"observed_at":"2026-08-07T14:00:34.410611Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.06290","last_updated":"2024-07-26T18:09:11Z","snapshot_observed_at":"2026-08-13T10:57:21.545059Z","submitted_at":"2023-07-12T16:37:31Z","title":"Instruction Mining: Instruction Data Selection for Tuning Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.06290","snapshot_observed_at":"2026-08-07T14:00:29.600015Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:29.600015Z"},"links":{"cited_paper":"/paper/2307.06290","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:9504732c9396083a4bc22ff007867dd67e0241662af261bb5ca7b3bb5d354518","observation_id":"c2e9f62e-708f-4c21-84bc-ea1c79d19ced","resolution":{"observed_at":"2026-08-07T14:00:29.600015Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:37.488225Z","title":null,"venue":null,"work_id":"d6686898-1f2d-4b9f-8874-d11324c522c2","year":2023},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:29.738159Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:d9946ff4c7f8c4cf1fbad868bc768da5468963377c3cd0a50e463158efa9643a","observation_id":"886f756d-3c45-4b6a-90e8-addf5704fe5b","resolution":{"observed_at":"2026-08-07T14:00:37.588231Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:37.343660Z","title":null,"venue":null,"work_id":"dcc814b6-3083-4e7b-9e41-321d49ecc83e","year":2017},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:29.914295Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:e22b1fa80cecedac89f5a66a26fe387947597ab2865afcdae1dcfb886ace4ca1","observation_id":"ea6c4a00-bfde-4796-8c33-c4182977a0a2","resolution":{"observed_at":"2026-08-07T14:00:37.424749Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:37.177911Z","title":null,"venue":null,"work_id":"5841b7c9-c1e4-4450-a5c5-47df84155b38","year":2022},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.089910Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:d9c79bceed762d13a7decfff7bbad8d9bb70b8774476f61099953562b84c6bd4","observation_id":"8059f304-dd4d-4f30-8fe7-cd4313097b88","resolution":{"observed_at":"2026-08-07T14:00:37.264612Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:36.991495Z","title":null,"venue":null,"work_id":"d67978ac-de11-4a4b-b9dc-c2de3af69ac1","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.238178Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:5ae6e1b76c9a16f8ca0f3c5a2554d67e20cc125145f141b36b30bcdd8edf4f57","observation_id":"43373927-271c-4ea3-a94e-82f052489631","resolution":{"observed_at":"2026-08-07T14:00:37.080345Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:36.803646Z","title":null,"venue":null,"work_id":"857815fe-e2e8-414f-9e13-adddb0a28f79","year":2017},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.288868Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:f9eda969facad99ea1cd117247a0b76409823cd35b62bb4907c695dfd59dd7c4","observation_id":"87f6604b-84cd-46fd-bb3d-71ee7d5a4266","resolution":{"observed_at":"2026-08-07T14:00:36.907393Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:36.608237Z","title":null,"venue":null,"work_id":"c7821090-84b1-4bfe-9d84-c4f07e792678","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.359664Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:23ce203cf3f2e6614a426818989f6354ef7d37f3d994a5b5f23c34cb4a2d0e19","observation_id":"11dd96d6-3b39-4e9d-aa74-11f3d5cafb96","resolution":{"observed_at":"2026-08-07T14:00:36.703652Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05779","last_updated":"2025-04-28T17:36:27Z","snapshot_observed_at":"2026-08-11T13:00:45.634760Z","submitted_at":"2024-10-08T08:00:12Z","title":"LightRAG: Simple and Fast Retrieval-Augmented Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05779","snapshot_observed_at":"2026-08-07T14:00:30.443560Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.443560Z"},"links":{"cited_paper":"/paper/2410.05779","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:3da3546732639da470e4e3ede43623af05301b5f4f3d12683475b0084f613c91","observation_id":"78d230ca-c325-4625-838c-04acee48a7dd","resolution":{"observed_at":"2026-08-07T14:00:30.443560Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:30.526434Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.526434Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:34c914ae7b0c145cbe90f55f9b6a026247d1f3391c746c1d670bbeb09f518b42","observation_id":"8eca1af2-ef67-49d1-9194-12b445dcac64","resolution":{"observed_at":"2026-08-07T14:00:30.526434Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:36.388682Z","title":null,"venue":null,"work_id":"70d331e0-c766-46cd-ae06-df6f470886d1","year":2017},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.603101Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:f0519aca3d6418e47c5df2ae7c60416d192f6ba66179ce32dcf710f37b2c43e1","observation_id":"7543da6e-9cc2-42fc-a5f8-0df8677c5d6b","resolution":{"observed_at":"2026-08-07T14:00:36.481755Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:30.678313Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.678313Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:7f344dbde182ae36b33fe61ef46e543f7aae0d5439d4d17a5ee464a7dbf4b571","observation_id":"8d1afdfb-ef07-4ec2-9b6b-296b9bfead64","resolution":{"observed_at":"2026-08-07T14:00:30.678313Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:36.137393Z","title":null,"venue":null,"work_id":"4ad4744e-d42c-4b87-a57f-463deda86144","year":2016},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.782135Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:c6b70739012f207999fa0f8bce0c99b1333cf4253155948bca1840808568717d","observation_id":"7e5587e7-c131-4354-af18-75673abbe7f0","resolution":{"observed_at":"2026-08-07T14:00:36.235786Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:30.845629Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.845629Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:e8d4e0ef05c316c68470ba64cfb3ee3bdfce851e268291ffcb00d125a4e971b3","observation_id":"7e841c88-6ea8-4a66-8005-302c34e511b1","resolution":{"observed_at":"2026-08-07T14:00:30.845629Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2001.08361","last_updated":"2020-01-23T03:59:20Z","snapshot_observed_at":"2026-08-13T17:41:53.092611Z","submitted_at":"2020-01-23T03:59:20Z","title":"Scaling Laws for Neural Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2001.08361","snapshot_observed_at":"2026-08-07T14:00:30.901388Z","title":"Brown, Benjamin Chess, Rewon Child, Scott Gray, Alec Radford, Jeffrey Wu, and Dario Amodei","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.901388Z"},"links":{"cited_paper":"/paper/2001.08361","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:69c171311b3ac7f2c6b11277b168835654a9d5b40342ce5d91edcb004e6f9b7e","observation_id":"3b045b9e-ddd9-4621-945e-658300ccf784","resolution":{"observed_at":"2026-08-07T14:00:30.901388Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:35.919708Z","title":null,"venue":null,"work_id":"c18e4f2f-e0c3-4faa-bee5-8b6c21bd1042","year":2025},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.981018Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:4a66bf5503da3340f01f66f23b02074bdc62e1f8572213400a5666c019e1af84","observation_id":"431cfa69-32a5-4461-a638-58ecba4f8a38","resolution":{"observed_at":"2026-08-07T14:00:36.014540Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08460","last_updated":"2024-10-03T15:46:13Z","snapshot_observed_at":"2026-08-13T12:00:44.648468Z","submitted_at":"2023-04-17T17:36:35Z","title":"LongForm: Effective Instruction Tuning with Reverse Instructions","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08460","snapshot_observed_at":"2026-08-07T14:00:31.089767Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.089767Z"},"links":{"cited_paper":"/paper/2304.08460","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:7ca3a29b84e16e9e31555532c4f4f4b3dfa9b509a87b4766a7a007cd169e121e","observation_id":"31c18690-2cb5-4afd-981b-b22ec3caf84b","resolution":{"observed_at":"2026-08-07T14:00:31.089767Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.16367","last_updated":"2024-06-24T07:17:59Z","snapshot_observed_at":"2026-08-12T23:36:25.369008Z","submitted_at":"2024-06-24T07:17:59Z","title":"On the Role of Long-tail Knowledge in Retrieval Augmented Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.16367","snapshot_observed_at":"2026-08-07T14:00:31.170015Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.170015Z"},"links":{"cited_paper":"/paper/2406.16367","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:fadcf5f3eef301f1b11c7f88e0a6c9963d22ed2cfe5b2ff008321971dd917ed4","observation_id":"50ed8937-6ff9-4060-bc55-45804d103389","resolution":{"observed_at":"2026-08-07T14:00:31.170015Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:35.726363Z","title":null,"venue":null,"work_id":"4db8df0b-3380-4721-8ed9-2bf9e156b5a2","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.239806Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:2fa9f9cbcfddcbe63029e09709b6aa258e0309a032c6729c2aeb53126d68cd5e","observation_id":"a1a95673-f78e-48fd-a18f-11386af5f299","resolution":{"observed_at":"2026-08-07T14:00:35.787443Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:35.530840Z","title":null,"venue":null,"work_id":"f72c6c9b-9619-4082-b59b-3dbfed2f0065","year":2023},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.324373Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:d310a4e60d7db267bb420c7b518573df780b39a3174347033b736db12a2ee782","observation_id":"9a39f985-5418-4a9f-bb1a-0dd0e521f7bc","resolution":{"observed_at":"2026-08-07T14:00:35.612378Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07503","last_updated":"2024-08-10T20:46:47Z","snapshot_observed_at":"2026-08-13T00:32:55.965962Z","submitted_at":"2024-04-11T06:34:17Z","title":"Best Practices and Lessons Learned on Synthetic Data","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07503","snapshot_observed_at":"2026-08-07T14:00:31.405117Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.405117Z"},"links":{"cited_paper":"/paper/2404.07503","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:f494cd17c2c27edab31463d9a9466ac32c58baf84d8420581821ba0a2d7986fc","observation_id":"5f8b9035-31b8-482b-abc4-c5df9f956dd4","resolution":{"observed_at":"2026-08-07T14:00:31.405117Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.15126","last_updated":"2024-06-14T07:47:09Z","snapshot_observed_at":"2026-08-12T23:42:57.310360Z","submitted_at":"2024-06-14T07:47:09Z","title":"On LLMs-Driven Synthetic Data Generation, Curation, and Evaluation: A Survey","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.15126","snapshot_observed_at":"2026-08-07T14:00:31.486151Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.486151Z"},"links":{"cited_paper":"/paper/2406.15126","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:9276c9a3ef137e5ec038cbf793bd252034a71fe58aacdea70b8cea1b2a9bf9a3","observation_id":"e988690b-12a9-44ef-9570-b074cb7ce78d","resolution":{"observed_at":"2026-08-07T14:00:31.486151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.11266","last_updated":"2025-05-19T14:51:44Z","snapshot_observed_at":"2026-08-14T09:29:57.623647Z","submitted_at":"2024-11-18T03:45:34Z","title":"VersaTune: An Efficient Data Composition Framework for Training Multi-Capability LLMs","version":5},"cited_work":{"arxiv_id":"2411.11266","doi":null,"metadata_source":"pith","pith_arxiv_id":"2411.11266","snapshot_observed_at":"2026-08-07T14:00:34.048977Z","title":"VersaTune: An Efficient Data Composition Framework for Training Multi-Capability LLMs","venue":"cs.CL","work_id":"bb296ddf-e814-4b00-bc87-8c69f59866c6","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.489767Z"},"links":{"cited_paper":"/paper/2411.11266","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:eb69242caea1566aa7031315a9829e003f9e7a7f564c082760a5f08d47af346e","observation_id":"8057caba-38b6-4025-b1b0-9f71da005ca1","resolution":{"observed_at":"2026-08-07T14:00:34.140189Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.16380","last_updated":"2024-01-29T18:19:08Z","snapshot_observed_at":"2026-08-13T04:32:26.957409Z","submitted_at":"2024-01-29T18:19:08Z","title":"Rephrasing the Web: A Recipe for Compute and Data-Efficient Language Modeling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.16380","snapshot_observed_at":"2026-08-07T14:00:31.493097Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.493097Z"},"links":{"cited_paper":"/paper/2401.16380","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:26bdd7e01f90be1b949ddf0f8c9beac4ecdb87d3957e78f17b21433a1dcb286e","observation_id":"af7264e8-e402-4814-abb7-82e5a8c83fc3","resolution":{"observed_at":"2026-08-07T14:00:31.493097Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:31.504032Z","title":null,"venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.504032Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:e070f426dbdd83a8045887eb2d09194fe5f67bbedd67a5116bcde4ebaee7c750","observation_id":"fddf5425-92bf-4879-ba64-9875b5c4c91c","resolution":{"observed_at":"2026-08-07T14:00:31.504032Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.00213","last_updated":"2024-04-02T20:09:45Z","snapshot_observed_at":"2026-08-14T09:59:52.654720Z","submitted_at":"2024-03-30T01:56:07Z","title":"Injecting New Knowledge into Large Language Models via Supervised Fine-Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.00213","snapshot_observed_at":"2026-08-07T14:00:31.608166Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.608166Z"},"links":{"cited_paper":"/paper/2404.00213","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:604193f2cff98b2885cb5744fc1c80cf11e0a17358b1a1d33c539baa0a0c120b","observation_id":"69709be5-7cdb-4e3c-8d42-31ce3339071c","resolution":{"observed_at":"2026-08-07T14:00:31.608166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.13296","last_updated":"2024-10-30T01:04:15Z","snapshot_observed_at":"2026-08-14T02:11:16.434887Z","submitted_at":"2024-08-23T14:48:02Z","title":"The Ultimate Guide to Fine-Tuning LLMs from Basics to Breakthroughs: An Exhaustive Review of Technologies, Research, Best Practices, Applied Research Challenges and Opportunities","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.13296","snapshot_observed_at":"2026-08-07T14:00:31.719031Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.719031Z"},"links":{"cited_paper":"/paper/2408.13296","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:2193cd95b7cc8b4df8c0e2d93b873920949f1580d5c992825ee749b66e135bb8","observation_id":"ff3899ac-635a-4ea3-bb34-96360efe4ed9","resolution":{"observed_at":"2026-08-07T14:00:31.719031Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:35.309263Z","title":null,"venue":null,"work_id":"229e328e-84f1-40fa-9e0e-b289c1dff73a","year":2017},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.788518Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:a61785c28cd88771b8ec70211d989b09355db0545bba7d1c2aeb08e31a13b8a4","observation_id":"b76f68ed-9d71-4b64-bc9f-80cfe12e1aad","resolution":{"observed_at":"2026-08-07T14:00:35.444372Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:35.156300Z","title":null,"venue":null,"work_id":"53679ece-3037-4538-9600-a38a709523ca","year":2016},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.906363Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:2f9dd7cdde0b79548318fbc8958fa6bd4b48ca13b1bcd308a56cd6f6788f36c9","observation_id":"9bacbbd9-88f6-45a6-82b1-0e3fa842bd0b","resolution":{"observed_at":"2026-08-07T14:00:35.213071Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:32.002622Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.002622Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:e55bcac61289fb9db6d9b3d37c66dbf92fb1c79ca95efc1fe6820bf0c147816b","observation_id":"f6af1e35-d7fd-435c-a3cf-c74b1af20c08","resolution":{"observed_at":"2026-08-07T14:00:32.002622Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:32.094121Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.094121Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:4e11f01bba8aa131e31ebd80134c6cee2ebc178180b5576c71bd3494b9b45bf5","observation_id":"84166ce3-1c08-44f4-a324-b31d1afe893e","resolution":{"observed_at":"2026-08-07T14:00:32.094121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:34.936372Z","title":null,"venue":null,"work_id":"0432fea6-39a1-4c28-a83f-38a7da1a4c65","year":2022},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.183983Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:7c9a028c383412ac90b2bb665bb4c4982a6d8e20dee4dec4856ff220f709e5a6","observation_id":"cb82ee31-057f-4611-9cca-968d857a2265","resolution":{"observed_at":"2026-08-07T14:00:35.071765Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-08-07T14:00:32.277975Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.277975Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:6c63de5947dbd599d096a36c01bb96dd5ce4151fd9c478709ce53b0660e55c92","observation_id":"7dfa25dd-518a-434d-a244-2d729dd91989","resolution":{"observed_at":"2026-08-07T14:00:32.277975Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1809.09600","last_updated":"2018-09-25T17:28:20Z","snapshot_observed_at":"2026-08-11T10:35:06.989614Z","submitted_at":"2018-09-25T17:28:20Z","title":"HotpotQA: A Dataset for Diverse, Explainable Multi-hop Question Answering","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1809.09600","snapshot_observed_at":"2026-08-07T14:00:32.393107Z","title":"Cohen, Ruslan Salakhutdinov, and Christopher D","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.393107Z"},"links":{"cited_paper":"/paper/1809.09600","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:9183e17d4c7e76f433e97be370a75a2d9ce8f3923681ca5f850bdeb479d1ff43","observation_id":"bf9b34c4-40cd-437b-9424-1955689571af","resolution":{"observed_at":"2026-08-07T14:00:32.393107Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.07431","last_updated":"2024-10-03T13:07:25Z","snapshot_observed_at":"2026-08-13T18:53:54.659653Z","submitted_at":"2024-09-11T17:21:59Z","title":"Synthetic continued pretraining","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.07431","snapshot_observed_at":"2026-08-07T14:00:32.484816Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.484816Z"},"links":{"cited_paper":"/paper/2409.07431","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:d65b98201e95697d0e67640afd48f889b66c6b59d2d8ba489b0ba32c2c685c38","observation_id":"022b3d1b-0f37-4f47-b71a-2ca76e38d9a1","resolution":{"observed_at":"2026-08-07T14:00:32.484816Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.14367","last_updated":"2024-01-25T18:14:57Z","snapshot_observed_at":"2026-08-13T04:34:35.372204Z","submitted_at":"2024-01-25T18:14:57Z","title":"Genie: Achieving Human Parity in Content-Grounded Datasets Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.14367","snapshot_observed_at":"2026-08-07T14:00:32.574415Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.574415Z"},"links":{"cited_paper":"/paper/2401.14367","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:74eb36d30686dc893a6038c535e1b8a2126b80b86ff7eeac800bd64a5599b7e5","observation_id":"b1f5342f-4abe-4484-a55d-015b20cf901a","resolution":{"observed_at":"2026-08-07T14:00:32.574415Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.13220","last_updated":"2025-05-19T15:02:59Z","snapshot_observed_at":"2026-08-07T15:43:09.026441Z","submitted_at":"2025-05-19T15:02:59Z","title":"SeedBench: A Multi-task Benchmark for Evaluating Large Language Models in Seed Science","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.13220","snapshot_observed_at":"2026-08-07T14:00:32.650100Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.650100Z"},"links":{"cited_paper":"/paper/2505.13220","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:89abf1e2f378be1b94fc03e29b262855ccb2ef4136fcda4da3038df1f9113195","observation_id":"236b59b8-3ae9-4f02-940c-250ec8b3a712","resolution":{"observed_at":"2026-08-07T14:00:32.650100Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11952","last_updated":"2023-05-19T18:26:26Z","snapshot_observed_at":"2026-08-13T11:38:29.478303Z","submitted_at":"2023-05-19T18:26:26Z","title":"Self-QA: Unsupervised Knowledge Guided Language Model Alignment","version":1},"cited_work":{"arxiv_id":"2305.11952","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11952","snapshot_observed_at":"2026-08-07T14:00:33.651259Z","title":"Self-QA: Unsupervised Knowledge Guided Language Model Alignment","venue":"cs.CL","work_id":"8e728853-820c-43d6-9348-4e7ad55220c2","year":2023},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.810716Z"},"links":{"cited_paper":"/paper/2305.11952","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:3b2bb4a56f392509cefb0d3987bcdf4ae4da46e4c43b50ffd5f1ed55419fa6ac","observation_id":"68bb9f2a-28dc-4380-98f6-bf1af4989348","resolution":{"observed_at":"2026-08-07T14:00:33.801103Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:34.718841Z","title":null,"venue":null,"work_id":"18812071-8230-4134-83f3-de82b8a46d73","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.905991Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:fa035c86c80efee53ffe20d6044a5256903bf0140f97132de38e23a4110c92f7","observation_id":"a7252192-a611-4aea-a82c-f95842ff96ab","resolution":{"observed_at":"2026-08-07T14:00:34.822257Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.14405","last_updated":"2024-11-25T17:57:55Z","snapshot_observed_at":"2026-08-12T15:10:50.113705Z","submitted_at":"2024-11-21T18:37:33Z","title":"Marco-o1: Towards Open Reasoning Models for Open-Ended Solutions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.14405","snapshot_observed_at":"2026-08-07T14:00:33.005432Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:33.005432Z"},"links":{"cited_paper":"/paper/2411.14405","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:b79a7f3f9719d88a1c61f98240cc46b3b9ebab08913081505eedf2a2607a8860","observation_id":"72fa98e2-b9cb-43a5-8a90-30b7875a2b80","resolution":{"observed_at":"2026-08-07T14:00:33.005432Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:34.521747Z","title":null,"venue":null,"work_id":"3773a2f2-dbf4-466e-b3b5-ad2f6e6c1a58","year":2022},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:33.103094Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:1920fc8ad2e69193da495414b16dbf2db9dbb6b2d284a2f388630be0421eb3da","observation_id":"69f78243-dde9-4186-a8a4-3affc971f670","resolution":{"observed_at":"2026-08-07T14:00:34.598822Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:33.241436Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:33.241436Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:9c06d3d383968bc7afad6b63e135b2f20861c4ed7321f9c8e33be1900d9f09a8","observation_id":"7b4ab52d-ed8e-4ec1-83cb-30967a0334a3","resolution":{"observed_at":"2026-08-07T14:00:33.241436Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:33.402147Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:33.402147Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:24be9425f369fe118308805e147e77f66141e093f1efc9043e2acbd855439abe","observation_id":"1d6430d5-9f3f-444a-a2a8-a6536ad0d4f7","resolution":{"observed_at":"2026-08-07T14:00:33.402147Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-13T15:26:37.993819Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation"},"reference_resolution":{"displayed":43,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":40,"verified_exact":3,"verified_fuzzy":0},"total_outbound_references":43},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 43 of 43 outbound references and 7 inbound Pith citation observations for arXiv:2505.20416."}