{"as_of":"2026-08-10T03:16:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8edbcc3271a88021d657bc3a1535be55ad02cffe86363b3b81e36d00ddd4326b","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":22,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":22,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":22,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":22,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T14:27:57.431339Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-03T17:58:46.763967Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":"2310.15393","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-07-03T17:58:46.763967Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":"8ad6c744-7172-4f32-8b84-442db774e9b6","year":2023},"citing_paper":{"arxiv_id":"2406.11794","last_updated":"2025-04-21T17:48:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-17T17:42:57Z","title":"DataComp-LM: In search of the next generation of training sets for language models","version":4},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-17T22:58:16.523267Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2406.11794"},"observation_digest":"sha256:673acaa805b68adee3ac68847690bb249c34defe53e4f8419ca3c46dfdea988c","observation_id":"6490f8cc-9cc1-4982-9a90-c9a9343b152a","resolution":{"observed_at":"2026-05-17T22:58:16.949433Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":"2310.15393","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-07-03T17:58:46.763967Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":"8ad6c744-7172-4f32-8b84-442db774e9b6","year":2023},"citing_paper":{"arxiv_id":"2502.00270","last_updated":"2026-05-13T20:43:28Z","snapshot_observed_at":"2026-08-03T15:35:37.964180Z","submitted_at":"2025-02-01T01:52:32Z","title":"DUET: Optimizing Training Data Mixtures via Feedback from Unseen Evaluation Tasks","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-23T03:58:48.967122Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2502.00270"},"observation_digest":"sha256:602955a4d979fce6032cce17800747ea99c4c18eca40bb18accb7c3f1147d54e","observation_id":"31ea5bb9-cd7b-43de-8a41-6cceb8cb0e22","resolution":{"observed_at":"2026-05-23T04:02:30.403038Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-09T14:27:57.431339Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.01804","last_updated":"2025-02-03T20:33:20Z","snapshot_observed_at":"2026-08-10T01:51:21.193055Z","submitted_at":"2025-02-03T20:33:20Z","title":"Soup-of-Experts: Pretraining Specialist Models via Parameters Averaging","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-09T14:27:57.431339Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2502.01804"},"observation_digest":"sha256:b423984ab8dc51386e1aeee106d3451ebd7e855c87a9a0b1edebe4a075efbd3d","observation_id":"253c7a5c-2696-4055-82b8-a1b0bab31ecc","resolution":{"observed_at":"2026-08-09T14:27:57.431339Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-08T04:58:32.207628Z","title":"DoGE: Domain reweighting with general- ization estimation, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.08489","last_updated":"2025-02-13T17:33:24Z","snapshot_observed_at":"2026-08-08T14:38:41.643106Z","submitted_at":"2025-02-12T15:26:08Z","title":"Salamandra Technical Report","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-08T04:58:32.207628Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2502.08489"},"observation_digest":"sha256:1f2799110af7c9bfb060df32ab9521a8af23d9fb06cbc34d010e5cb42c1b90eb","observation_id":"3f6c144d-914d-4cbb-aa39-1971afbec914","resolution":{"observed_at":"2026-08-08T04:58:32.207628Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":"2310.15393","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-07-03T17:58:46.763967Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":"8ad6c744-7172-4f32-8b84-442db774e9b6","year":2023},"citing_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-09T06:35:27.813995Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2505.09388"},"observation_digest":"sha256:ecc2ec223766f34cc14dd9292b15f0b6f0df301b89aad8fcccfad6d76985655c","observation_id":"0e69eb72-b581-491f-aef7-6907b899adb2","resolution":{"observed_at":"2026-05-09T06:35:28.492131Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-07T15:09:34.322875Z","title":"Doge: Domain reweighting with generalization estimation.arXiv preprint arXiv:2310.15393, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.16066","last_updated":"2025-05-21T22:34:13Z","snapshot_observed_at":"2026-08-09T02:51:38.941026Z","submitted_at":"2025-05-21T22:34:13Z","title":"Merge to Mix: Mixing Datasets via Model Merging","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T15:09:34.322875Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2505.16066"},"observation_digest":"sha256:19ef3f450f7cf2251adad92fec5d65fc6f8520463fb36b687f5f62d2a8a2914b","observation_id":"7fe996f9-91aa-4d5b-8d92-a85e9fb122ed","resolution":{"observed_at":"2026-08-07T15:09:34.322875Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-07T14:33:13.058223Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18458","last_updated":"2025-06-01T16:00:34Z","snapshot_observed_at":"2026-08-07T20:59:19.597449Z","submitted_at":"2025-05-24T01:57:12Z","title":"A Survey of LLM $\\times$ DATA","version":3},"reference_index":135,"source":"pdf_text","source_observed_at":"2026-08-07T14:33:13.058223Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2505.18458"},"observation_digest":"sha256:a0a2e22768cb68f32e6e7fc6153d5cca0299727c86763c522b143c21e0948db6","observation_id":"6d6104bd-1540-4835-b0ce-e9c371302825","resolution":{"observed_at":"2026-08-07T14:33:13.058223Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-07T14:09:45.950397Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19893","last_updated":"2025-05-26T12:23:26Z","snapshot_observed_at":"2026-08-09T06:26:16.829487Z","submitted_at":"2025-05-26T12:23:26Z","title":"ESLM: Risk-Averse Selective Language Modeling for Efficient Pretraining","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T14:09:45.950397Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2505.19893"},"observation_digest":"sha256:3aaf70b6dc4c7f17732a8c7ab1318c5bbd473c789977a1abdfbd8c9df08dcf74","observation_id":"eb274786-b2dd-4323-ac48-42b72e660c48","resolution":{"observed_at":"2026-08-07T14:09:45.950397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-07T14:04:22.784397Z","title":"Doge: Domain reweighting with generalization estimation, 2024 a","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20380","last_updated":"2025-05-26T17:32:14Z","snapshot_observed_at":"2026-08-09T06:26:18.435852Z","submitted_at":"2025-05-26T17:32:14Z","title":"GRAPE: Optimize Data Mixture for Group Robust Multi-target Adaptive Pretraining","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T14:04:22.784397Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2505.20380"},"observation_digest":"sha256:e0c8664da15827d9bf23cf07112bd6f79755bce919c772dd7d444c3c52151f57","observation_id":"59c3bb8f-c3bd-49c5-977d-081837ee815c","resolution":{"observed_at":"2026-08-07T14:04:22.784397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-07T13:14:25.236291Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.22308","last_updated":"2025-05-28T12:50:09Z","snapshot_observed_at":"2026-08-09T18:23:39.071500Z","submitted_at":"2025-05-28T12:50:09Z","title":"Transformers Pretrained on Procedural Data Contain Modular Structures for Algorithmic Reasoning","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-07T13:14:25.236291Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2505.22308"},"observation_digest":"sha256:471ff1ee4782ad3858fefc7437916c550932ca7bd2c2d770796c7fe630e8c239","observation_id":"a7412de9-a583-4bf4-8cec-43376413809c","resolution":{"observed_at":"2026-08-07T13:14:25.236291Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-07T00:42:57.723303Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.12966","last_updated":"2025-06-15T21:08:51Z","snapshot_observed_at":"2026-08-08T21:39:58.375574Z","submitted_at":"2025-06-15T21:08:51Z","title":"Assessing the Role of Data Quality in Training Bilingual Language Models","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T00:42:57.723303Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2506.12966"},"observation_digest":"sha256:5829f57f9edcad0623b19672db033df0e4865e0eb4a90993abe29b9265051869","observation_id":"87cb0481-1866-4cfe-b0eb-202b06052e3c","resolution":{"observed_at":"2026-08-07T00:42:57.723303Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-06T16:53:09.681117Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.12466","last_updated":"2025-07-16T17:59:45Z","snapshot_observed_at":"2026-08-09T12:12:41.025897Z","submitted_at":"2025-07-16T17:59:45Z","title":"Language Models Improve When Pretraining Data Matches Target Tasks","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-06T16:53:09.681117Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2507.12466"},"observation_digest":"sha256:17036b4485c0e22c55591d6201d69802f1a58b42941bd3f8937db61dd99feed7","observation_id":"508bcc58-c794-498b-b82c-c6a3a224f40b","resolution":{"observed_at":"2026-08-06T16:53:09.681117Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-05T17:26:28.822567Z","title":"DoGE : Domain reweighting with generalization estimation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.16279","last_updated":"2025-08-22T10:35:56Z","snapshot_observed_at":"2026-08-07T15:35:45.566152Z","submitted_at":"2025-08-22T10:35:56Z","title":"AgentScope 1.0: A Developer-Centric Framework for Building Agentic Applications","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-05T17:26:28.822567Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2508.16279"},"observation_digest":"sha256:815205973f1cbef99592ccd382ce3f7fb85f7391dd193e05977c7a34bf558de2","observation_id":"20364c03-6812-40b0-a322-477384ea3c35","resolution":{"observed_at":"2026-08-05T17:26:28.822567Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-04T11:16:09.547484Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.06048","last_updated":"2026-06-18T04:28:50Z","snapshot_observed_at":"2026-08-04T11:16:05.008855Z","submitted_at":"2025-10-07T15:42:33Z","title":"BLISS: A Lightweight Bilevel Influence Scoring Method for Data Selection in Language Model Pretraining","version":5},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-04T11:16:09.547484Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2510.06048"},"observation_digest":"sha256:c30f35b103822116d3f3aff0aee7a0b9ea797b248251a846c826d72de2608f51","observation_id":"95db6ff2-39a0-4749-87cb-42cec357cb8d","resolution":{"observed_at":"2026-08-04T11:16:09.547484Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-03T06:55:10.051129Z","title":"Doge: Domain reweighting with generalization estimation.arXiv preprint arXiv:2310.15393,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2601.21725","last_updated":"2026-05-27T19:52:53Z","snapshot_observed_at":"2026-08-04T21:23:57.701786Z","submitted_at":"2026-01-29T13:48:43Z","title":"Procedural Pretraining: Warming Up Language Models with Abstract Data","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-03T06:55:10.051129Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2601.21725"},"observation_digest":"sha256:3824ce2d0f31909d54789fd2f6838e65505bce5a2cd6f6c22e6a9e061e23a126","observation_id":"69fc3645-77f9-409b-aa8f-c51b2d80938d","resolution":{"observed_at":"2026-08-03T06:55:10.051129Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-08-03T04:17:27.173478Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.05547","last_updated":"2026-08-05T16:57:37Z","snapshot_observed_at":"2026-08-08T23:12:50.234380Z","submitted_at":"2026-02-05T11:06:37Z","title":"Multi-Task GRPO: Reliable LLM Reasoning Across Tasks","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-03T04:17:27.173478Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2602.05547"},"observation_digest":"sha256:96c2dfa97c6fdb671ec2bc0eb51f92a6ad4d25da6f99cbb1ec7adf01f066bfcc","observation_id":"3e90dcf9-4d61-4d34-8fcf-804d7f4a3014","resolution":{"observed_at":"2026-08-03T04:17:27.173478Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":"2310.15393","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-07-03T17:58:46.763967Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":"8ad6c744-7172-4f32-8b84-442db774e9b6","year":2023},"citing_paper":{"arxiv_id":"2604.08366","last_updated":"2026-04-09T15:33:00Z","snapshot_observed_at":"2026-07-06T22:57:26.678728Z","submitted_at":"2026-04-09T15:33:00Z","title":"Scaling-Aware Data Selection for End-to-End Autonomous Driving Systems","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-10T17:22:25.943097Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2604.08366"},"observation_digest":"sha256:5f48f788c8d1f2b22907b0609b714072de931d76718140f733141408261335f2","observation_id":"1ce2030e-a07b-4200-81da-e2f4711a2998","resolution":{"observed_at":"2026-05-11T06:56:02.228988Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":"2310.15393","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-07-03T17:58:46.763967Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":"8ad6c744-7172-4f32-8b84-442db774e9b6","year":2023},"citing_paper":{"arxiv_id":"2604.08519","last_updated":"2026-04-09T17:55:50Z","snapshot_observed_at":"2026-07-06T22:57:35.435713Z","submitted_at":"2026-04-09T17:55:50Z","title":"Cram Less to Fit More: Training Data Pruning Improves Memorization of Facts","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-05-10T17:42:31.465077Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2604.08519"},"observation_digest":"sha256:5ae62dbcafa2ddf72e569c081a0c45a1b9549d2d17cd313d005045f9db82f8e2","observation_id":"76ec9131-0844-4027-8ad3-3c12d15ed7b7","resolution":{"observed_at":"2026-05-11T06:15:58.891380Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":"2310.15393","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-07-03T17:58:46.763967Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":"8ad6c744-7172-4f32-8b84-442db774e9b6","year":2023},"citing_paper":{"arxiv_id":"2605.10288","last_updated":"2026-05-12T10:19:03Z","snapshot_observed_at":"2026-07-06T23:22:18.704329Z","submitted_at":"2026-05-11T09:50:10Z","title":"BROS: Bias-Corrected Randomized Subspaces for Memory-Efficient Single-Loop Bilevel Optimization","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-12T04:47:04.868735Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2605.10288"},"observation_digest":"sha256:c7d8cd316f84b80247b536353e30f14b88720b4d2efc0659878963dcd3981d54","observation_id":"2af674b5-abf9-4948-aa15-1b45113e6b0b","resolution":{"observed_at":"2026-05-12T05:56:26.614355Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":"2310.15393","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-07-03T17:58:46.763967Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":"8ad6c744-7172-4f32-8b84-442db774e9b6","year":2023},"citing_paper":{"arxiv_id":"2605.10288","last_updated":"2026-05-12T10:19:03Z","snapshot_observed_at":"2026-07-06T23:22:18.704329Z","submitted_at":"2026-05-11T09:50:10Z","title":"BROS: Bias-Corrected Randomized Subspaces for Memory-Efficient Single-Loop Bilevel Optimization","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-13T06:20:42.781613Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2605.10288"},"observation_digest":"sha256:308397277b8b016422ec49705c59b727bb410f01bd56d1edc72bb5908fe17ab8","observation_id":"cf2d659c-ce7e-46c2-866c-0edb0849dc76","resolution":{"observed_at":"2026-05-13T06:22:23.337367Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":"2310.15393","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-07-03T17:58:46.763967Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":"8ad6c744-7172-4f32-8b84-442db774e9b6","year":2023},"citing_paper":{"arxiv_id":"2606.00571","last_updated":"2026-05-30T06:47:24Z","snapshot_observed_at":"2026-08-07T13:20:18.650137Z","submitted_at":"2026-05-30T06:47:24Z","title":"On the Difficulty of Learning a Meta-network for Training Data Selection","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-06-28T19:14:24.554237Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2606.00571"},"observation_digest":"sha256:95a214c7fb35cbd8d06776d7020dac967a0e28d9f4fadd26c1134f69259b055e","observation_id":"2c7a25d6-515f-4f4b-84cf-9f872e788598","resolution":{"observed_at":"2026-06-28T19:22:34.801323Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation","version":2},"cited_work":{"arxiv_id":"2310.15393","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.15393","snapshot_observed_at":"2026-07-03T17:58:46.763967Z","title":"Doge: Domain reweighting with generalization estimation","venue":null,"work_id":"8ad6c744-7172-4f32-8b84-442db774e9b6","year":2023},"citing_paper":{"arxiv_id":"2607.01686","last_updated":"2026-07-02T04:27:45Z","snapshot_observed_at":"2026-08-09T02:25:47.839654Z","submitted_at":"2026-07-02T04:27:45Z","title":"WARP: Weight-Space Analysis for Recovering Training Data Portfolios","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-03T17:54:31.856386Z"},"links":{"cited_paper":"/paper/2310.15393","citing_paper":"/paper/2607.01686"},"observation_digest":"sha256:d5a593a056a4457e9c2aec4178ed642f141e9e229cf1e6b48abaa5e2a9966d36","observation_id":"fecef004-9f40-46c6-8c35-0b9e2f250760","resolution":{"observed_at":"2026-07-03T17:58:46.765477Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2310.15393/citation-record","integrity":"/paper/2310.15393/integrity","json":"/paper/2310.15393/citation-record.json","paper":"/paper/2310.15393"},"outbound":[],"paper":{"arxiv_id":"2310.15393","last_updated":"2024-02-05T16:33:05Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-09T06:27:27.325450Z","submitted_at":"2023-10-23T22:51:58Z","title":"DoGE: Domain Reweighting with Generalization Estimation"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 22 inbound Pith citation observations for arXiv:2310.15393."}