{"as_of":"2026-08-09T17:01:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:823e72a42ec22e5709885ab52e94c7fe60ab1696546f31330347b4828d557676","coverage":[{"denominator":228,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-25T23:26:57.147374Z","state":"measured"},{"denominator":100,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":100,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2606.24752/citation-record","integrity":"/paper/2606.24752/integrity","json":"/paper/2606.24752/citation-record.json","paper":"/paper/2606.24752"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":", author=","venue":null,"work_id":null,"year":2000},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:1fe4360027fa390731f42bbf10957044071d71b4ee40331f4ee2bc6834f4c808","observation_id":"fb217d6d-c690-4ccc-9f4c-f5cd76062ef4","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Early Word Catches the Weights , url =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:c5414c6318007f57408660c0bddaf803425c31644d3c76d15ce2fee5383062e7","observation_id":"9c6e2bcc-da03-490f-8ca1-02af4bea9bba","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"On Warm-Starting Neural Network Training , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:ab1750cd7799ffa0da7deef79e894da9acc875e6167bbe1cd0b84de24348836a","observation_id":"1b17de25-a478-4573-bb4e-e94fb3a7428c","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"39th International Conference on Machine Learning , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:aaa6d380283e898c861cf5f2e7f169dec6ff0edefb21d1a8e3ecc18a45264630","observation_id":"33a9b314-6e29-48dd-850c-19066aff42c1","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"10th International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:06a238bc57865f4e8c1116f57d0665f4a8be1fe00e7298101bab114d2f215c12","observation_id":"512b99b9-1b67-4a7b-9d64-8129fbc94203","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Proceedings of the 40th International Conference on Machine Learning , pages =","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:548ac04ea0b0bc61bca3769889919626cb9c80161ebec310f775e0f688baccba","observation_id":"e8333cfc-2773-45d2-94f1-9267569bbad3","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Nature , volume=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:acbcf5362dbe7a98c83441c44d0eb4a74966ce11b3c9b96ce433ac09e57b7d91","observation_id":"2878dc90-29e4-4ec4-ada3-17e26cfa8f77","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"41st International Conference on Machine Learning , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:bdc0b359934ab4c0bf16c631d473ed1859e0dc62f4df48387454cf34682c64fd","observation_id":"510bb74c-9d31-44a4-8383-dab1ede10def","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2024 , note=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:783d5eabf5dffb926fd76074cf7380ebb8eef3128a522136aef3b896467c9fc1","observation_id":"b708b6fb-09da-4923-af45-879e5455f2e4","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:06f3d83c2c0baa9a50492b5b38ce107fc5338249d6d3999faf3cf5c29fdd3f21","observation_id":"7161a4e8-1231-4d41-8644-010c48e0b7ca","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"12th International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:f458a4e2ec47b6d12a756383ce775c039e4f0a4a5e125e2d0ca7f1b4725ee9df","observation_id":"dc11a154-bbf5-4b89-a8bb-29291522bed2","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Rupam Mahmood , title=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:bf0f4839065804a85da1a932ea7cdd2dd5e9440c7595c94da2aa0f0b0e96994c","observation_id":"d0a4ebd4-a5a1-4194-a159-17075228b33c","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Thirty-eighth Annual Conference on Neural Information Processing Systems , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:c9c3ce24f4789d9bb5ee03765f4fd74addbc9baec5bbf5575cd63b67f85e6484","observation_id":"a246990b-f0c2-42e4-9402-010042905ae7","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"3rd Conference on Lifelong Learning Agents , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:e880169df48e822689ad1f7e85cf5f9a7c2521796f77005c31315d74248a01e8","observation_id":"97b5c321-af23-4578-b688-11249b6df1c3","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"3rd Conference on Lifelong Learning Agents , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:2ae71dc3be9cf4f16bc8c4419f1256ee070dc3603a1db62783c73071898d4f49","observation_id":"e6622406-89aa-4857-baf1-7d4537558b2d","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"4th Conference on Lifelong Learning Agents , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:01f1b48354994896f92a40041ae4aa8d1722a6f4bbc4287e21e591ec86b5e921","observation_id":"58e7eb4b-1b6f-481f-b3e0-ee546143cc4e","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Thirteenth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:9182f5372a40419778741b818647c1ca27f4ed1441620c6bdbafa360df6fcc04","observation_id":"49381c8e-8d09-4888-8705-8748231496f1","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Thirteenth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:7025323d9b66bd9d373d745fb5c0c81b434ff3df58eccddea687bea54b551ea8","observation_id":"147b9a44-228a-4809-a8f5-0f3accc7b152","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Thirty-ninth Annual Conference on Neural Information Processing Systems , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:b77d2d83f6b144497359edd007bb70aa428311038ff26d5ae196c2f5f84b9ffb","observation_id":"b9aa965b-ed80-4313-bdf3-7e5f5b62c176","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Fourteenth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:7bcb260dba91317b2d00338885d8aa02daa4e44748058210c7e389d6ecacf139","observation_id":"b659cbd9-52d0-41d1-8b73-7f7d80b52e32","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Forty-second International Conference on Machine Learning , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:0f2b0271e0ce39c92856261c1b681ea315a7dbb094cfd8ca06c22bfb22f13a90","observation_id":"ea70cbc1-8ac1-411e-996b-c284d5760ac1","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:fd04b1d45f3c6b055db860300e3a01adc4d33205c03522ecf00905953dc83a50","observation_id":"5368a2be-b6fe-4285-be13-40c6dc00c90e","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Fourteenth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:73d834a77ddf32281198afce050712982417644972a296ff683cd8248ed045f0","observation_id":"fa13215b-9a90-4d27-820f-15f6fb65a66d","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Thirty-ninth Annual Conference on Neural Information Processing Systems , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:9954dd601a1306fdeba80e3f19f07afcea3d3c0a6fabb18aba3c824f75cc8e03","observation_id":"0134e8fa-b14d-4dbd-9da1-f2c6a8595515","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Resetting the Optimizer in Deep","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:9c02fda1f14110d853b6bbaa218b5cbfd2b2660ee41d324360b67989ed2ff5b8","observation_id":"13b947b8-bfc6-4d01-a5d5-eb8d016c291b","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2026 , publisher =","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:2054b6a468632eb65ab64896289e1f95a8b5ab42bd276a810ab1ace2e9006582","observation_id":"e96024ab-74b8-488e-9f05-dcd5f6581343","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:d8cb19f5bb2b03add8a2789801d8e3a1dd20672810be9f0867015ba22b717014","observation_id":"8e9ac0e0-e7ca-4fde-8719-e41c158108e9","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:364674e677ae79eda20e198a02c2d2a2c1029e99aba98e04cd449c7901698e20","observation_id":"665e3858-196d-46aa-8df3-84c8cb34d6dc","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:f0511e2e75d351239cecf132527c44227e9adb83b924a0f80db74f790e358e4f","observation_id":"0f0a3252-5df7-4fbf-bec3-d4ca902f10ae","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2018 , publisher=","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:c8581256dca15aa304a0bb348eef06d431f9293bdcf71553b77c269023c83cd7","observation_id":"4512a284-23d8-47ba-9e99-0f90b633360a","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":", title =","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:99b3c714a02d0777d703a494f763ef6d4df014f154a5dc9d3107e6d459e24c33","observation_id":"fadff9a7-ec37-4e53-8c32-8d6e05e2b829","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:87339e7b59fe70bb145ba983ecba1d0196ec99b6a1fbc0f324485f1d7dd8df23","observation_id":"f9bdce77-3f6d-44d0-a1d4-b4476fb6d6b5","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:fefb94fe9584fea18a64d3450d8241f12f950feef4e1db0028ca957d7312d368","observation_id":"f124c51d-799b-4174-889f-8579ee2129b1","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Tuning Large Neural Networks via Zero-Shot Hyperparameter Transfer , url =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:ed1f76ef3dd7fe8f9bb96f39ad9c1392377ea740485c1bf9a54085a4c88117f8","observation_id":"d51a726f-717d-4f1e-a33c-98d8f61d85ba","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"When Attention Collapses: How Degenerate Layers in","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:ba300f024585f58557530a0fef38bc1495ebf09b975701cd68e66821f5ee403d","observation_id":"13ddb4f0-1cfc-481a-8065-2786e325aa05","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Proceedings of the 40th International Conference on Machine Learning , articleno =","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:bad6277c3893372e8aa2c8cc33d2538a2f6e3a59e022c050065de84c20823396","observation_id":"52c9ea2b-b7ac-45b7-af9a-c8c6e19355b0","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Twelfth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:dea1234e0d2f577dfd66491e5fa17db6e5b9952629a40cd7ea2ed2a1f9039e73","observation_id":"955f15a1-cdc9-4644-aee7-5b1814f2ba50","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Twelfth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:3ec244b4bad34ded8c9c7fc6f309c6f7d6fa12a7152ab32205e004b2b1f9a112","observation_id":"abbcf827-77da-4c6e-bb48-bdb0202de86e","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"A Comprehensive Survey of Continual Learning: Theory, Method and Application , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:5d50ce79e6331d5231b3a96f0475223763ca89b88e8414168681a95f2e438fd2","observation_id":"284b7879-e2f0-435b-94c2-0297d0fd6c78","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Frontiers in psychology , volume=","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:64a12f19bb3e634164a6ad7bed45c1d344289852b4662771c407892c5d6b6b84","observation_id":"6443ef38-1a0c-409b-becf-0885210ad3fe","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Psychology of Learning and Motivation , volume=","venue":null,"work_id":null,"year":1989},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:62279dbf08bdbbe8c204a2d2c972872d28b42499136f6f1e022e6f7963846ebf","observation_id":"d187b74a-5ce3-4416-a617-df67d0c95f22","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"French , keywords =","venue":null,"work_id":null,"year":1999},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:0e19805e0dd3f520d51a4d48355f1e1fc1c9949d90e5f20421af133fdc294d01","observation_id":"4469cc60-ac80-4a0c-ab4e-d278461946d8","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Proceedings of the national academy of sciences , volume=","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:1b28db4486db8f673879981dc6660ff63ef131ecb2d6374e4f907cb4eb48ef91","observation_id":"c9f241f8-5cc8-43c6-823a-bb4aba4a44a4","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"and Nguyen, Thien Huu","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:8f505e4e922f10671b39ea6a8172f9ee728ab5f95ef5e91e3fdd71fdb7251186","observation_id":"2365676e-45e9-465d-bbbf-4d479aef5c00","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:2b2cbd0217dc937d12472509dc34f10ef4af7459895271ed725b08768111c685","observation_id":"b87b3a79-aff4-406d-82a4-935a42e8fae8","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"International conference on machine learning , pages=","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:923a8fd830aeaf58faa351ba222620dafae2ced20e839ef6b8d772abf6809cf7","observation_id":"57906729-ec67-448a-a912-826edad0fb8f","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Proceedings of the European conference on computer vision (ECCV) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:85501c7f9c3767716cf7245a4ca70ae5bc4855d561b35576f2ef10db50a0eb92","observation_id":"b126caa0-3b44-4720-925b-cd275bfb9be6","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Nature Machine Intelligence , volume=","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:cf9637e31e3931564923430a1e593886e80cd55b7720655093eeb20ce4792b14","observation_id":"139fd9d3-e247-4faf-b5ed-c5d94bf840d4","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"INFORMS Journal on Applied Analytics , volume=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:cc39c5b304967fb1a9898e3d95dabf3dd1c78f478e22b693bca95e49b78a2e24","observation_id":"083f17ba-8078-4416-b659-09d7c322fe6c","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Machine Learning , pages=","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:b49f147b395aa0317e7c74b88402d1e34e901dbade85ff64708c9d9834f0467e","observation_id":"dfb987ae-c23e-4c1a-bfca-cec88608b787","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Proceedings of machine learning and systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:ca0fc999e1d1f4baea7e2d08080fd35cfdb73f43c53c13cbaec31771ab27f7c5","observation_id":"8ddb0f2c-22b0-42e1-a7d2-b307908f8975","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"and Pechenizkiy, Mykola and Mocanu, Decebal Constantin , title =","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:17e7440d88b8dde5ae9b3a7f8cdba4c330e6da036ade3f35c92e59377f6056f2","observation_id":"e89b746f-3e2d-4f61-84d2-2d397136a8cb","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"The Twelfth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:2b8abf68473e0e048b44c60f2953d7869804199a4feae8e3802ee0f6b59322ea","observation_id":"9a6bad98-dc94-44c3-9596-120849ccc1e4","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:dde2d752f119f43c7c1b3a54601aa5702901fa668df2c011bcdacdc21912c0b4","observation_id":"6a5062b2-4deb-4701-bd86-9689a322afe7","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2018 , publisher=","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:fca64945d15b7aeb789425e633b8b6cad173a169eb885b7e56c2b8a23814658d","observation_id":"4cd4321f-4733-44c9-bf9f-c0d2ebff4dea","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"International conference on machine learning , pages=","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:715fa7ad750333625baba53bef1a689be0e03c7ae1f2b17d28ac70d0e71332ae","observation_id":"9c240f88-eba0-458c-90a9-5af287388430","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:ee1e17ed108f0555d587e87127b1cc839657fa5611798aedbf7b23103a3670cc","observation_id":"9d9f5a4b-14a9-44b0-96eb-8f6493ea5b9e","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2019 , url=","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:84ee6fa70035a5c31530adb190d9168cac16c51d0687453b3bf4b41ef021fa69","observation_id":"74c145e4-c3eb-40f3-9bbe-5a53e49858ac","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Proceedings of the Thirteenth International Conference on Artificial Intelligence and Statistics , pages =","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:eaeca128ab1ee39888795d4f7a8064dab2ab14c3422e40f833b5d3adc202b347","observation_id":"7f104e83-c1ab-48b8-94da-56493efb838a","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"IEEE International Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:80ecf7393a393d38a06a20515e995c011e344c11b18569f12e8760756ba499bf","observation_id":"216aa91a-2796-4148-85e5-09953a484b46","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Proceedings of the 37th International Conference on Machine Learning , pages =","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:3e32009cd1c375c9e29ce0ca6e91716a5cc1f88de7a51c258d42563886c76cf8","observation_id":"c2dc2903-b6a2-414d-b3c2-d3a0bfc01bb1","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:0af92f50b5c8258c057f7f6d86f833a4db42c365e40df876c12ac86c659827a1","observation_id":"4da27fb8-9988-4a7b-b973-b0713cc5ae77","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Proceedings of the IEEE , volume=","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:cfa55a7f74b1f729b0229d661274ef51b30d1c74d70b181473f8be94ceea4e9b","observation_id":"129bd4cd-dddc-439f-a1f9-6d2e57c48e88","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:e177d2c205cac8d05f028ab0298b10c3c08e5246dcc196ac57210222e866ea17","observation_id":"54cdbb49-3b57-40ae-9e72-1df1961ecdfe","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Advances in neural information processing systems , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:c7cf9cf6ce01d32a1252f6c3863d2447431c1aa4abeb6d32558ca2c26ee8d780","observation_id":"443f99b7-cb30-4724-92ec-f3736acda710","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"9th International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:245513192cfbcc710c37ad0e6be46799050093e80fac5af83ca5dae5a6912e04","observation_id":"9e441efa-595c-4119-a327-5c5d85c79e7a","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2024 , eprint=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:ed93268a0c7ce94e8af0f9c1ed1759bf72aada9b8ed86eef1f96fb2a0cf68ae2","observation_id":"3982b4cb-8b19-4480-b9d9-673796161c51","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"ICML Workshop on Deep Learning for Audio, Speech and Language Processing , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:4de4be1f60429d5c1e17fbaddca819793bccddad7885d8d3129d2404c075cab0","observation_id":"793dd782-382e-4b8e-8ba0-b3f33ee7c8ce","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1606.08415","last_updated":"2023-06-06T01:53:32Z","snapshot_observed_at":"2026-07-06T05:01:27.910364Z","submitted_at":"2016-06-27T19:20:40Z","title":"Gaussian Error Linear Units (GELUs)","version":5},"cited_work":{"arxiv_id":"1606.08415","doi":"10.18653/v1/n19-1122","metadata_source":"pith","pith_arxiv_id":"1606.08415","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gaussian Error Linear Units (GELUs)","venue":"cs.LG","work_id":"0466fd22-03a1-4a61-af0a-a900e77bb023","year":2016},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"cited_paper":"/paper/1606.08415","citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:76c3b8ec15509c3ac20459fab0b6ef24e7072dbf837cf4acbaf1da346051c1e6","observation_id":"852bc18f-1942-444c-add4-07f4020b67eb","resolution":{"observed_at":"2026-07-04T17:50:00.015323Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1710.05941","last_updated":"2017-10-27T17:45:21Z","snapshot_observed_at":"2026-08-08T18:23:31.977872Z","submitted_at":"2017-10-16T18:05:45Z","title":"Searching for Activation Functions","version":2},"cited_work":{"arxiv_id":"1710.05941","doi":"10.48550/arxiv.1710.05941","metadata_source":"pith","pith_arxiv_id":"1710.05941","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Searching for Activation Functions","venue":"cs.NE","work_id":"3a43a02d-e005-47ad-8373-c166e20c9ee9","year":2017},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"cited_paper":"/paper/1710.05941","citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:86dd0d07412b05711097d1da4f7eefe982ffe0e0ba8d1b27870f1094a30108fb","observation_id":"7d2a03cc-4e78-428f-a316-58048b3a1d43","resolution":{"observed_at":"2026-07-04T17:50:00.017838Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-07-13T22:50:52.328775+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T22:50:52.328775+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"IEEE Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:b4b0f12197597b61aa13185f5c03104436c5e135a655ea539d678a36df39f091","observation_id":"9bce97f8-4bc1-49c3-a825-9dcaccb9f30b","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:664f1f3d8e28b9e52d5ff15c3b6084f746e05e0d6dc0f3fbfe29253cff99c662","observation_id":"403d898b-3f36-4a7e-a11f-721db78b2135","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"parse_uncertain"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Proceedings of the IEEE/CVF international conference on computer vision , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:8186f74eaded2ce01f728bda7b49c845060d551de84c05966a8c50d261613635","observation_id":"b317d558-75ff-4bad-974c-0791cd57f934","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"AAAI , volume=","venue":null,"work_id":null,"year":1992},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:d2d1f8958f5060de912d6fe4767c4f634b64e4aec82c3329cb8b9c1f148af7ec","observation_id":"00881b04-0875-41c8-baa1-ad9e5682e651","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.02342","last_updated":"2025-07-09T01:03:54Z","snapshot_observed_at":"2026-08-05T08:25:55.043525Z","submitted_at":"2024-02-04T04:55:54Z","title":"MetaOptimize: A Framework for Optimizing Step Sizes and Other Meta-parameters","version":6},"cited_work":{"arxiv_id":"2402.02342","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.02342","snapshot_observed_at":"2026-07-04T17:50:00.019136Z","title":"Sharifnassab, S","venue":null,"work_id":"feaa8537-640e-47db-aa40-d13a859bc581","year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"cited_paper":"/paper/2402.02342","citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:d4c505a1f000610b402d65c8ff65b47468f2423c19c08ad759cd76fe69128234","observation_id":"6838dc90-ab0e-4edd-a1ee-4bdd2bd5b3d1","resolution":{"observed_at":"2026-07-04T17:50:00.020854Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:76ab84cbe720daad46dad912d7b62139b4071e4f9911a5e80715e8e06ffd54bf","observation_id":"03bbc532-6ff3-48c6-b1e2-cbd6d8319c18","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:f9abd48a4e8905481005fdfef300d499ae4a59ba910493898da66ff1defb5059","observation_id":"f6e331d2-a0fa-4079-9c09-cab204a50403","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:af543395b8355e6d3200749ef50f50956c2fabbdadd26226353a68f4e3abf7a0","observation_id":"2db18e1c-8249-4a8e-98a7-72b8cc5c4265","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2004 , publisher=","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:b7047df66f7415e2b521b1da1e6c9fd25edd928d976e2cdbc0028cba0ebc16f6","observation_id":"fb4d5fbe-dd57-42b6-b3c9-d0856bb4fe45","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:3c14c0dfa6e12a2a3cda193decf5b2da801eab16f09c0f4b7976aa72307c6c3d","observation_id":"0609f128-bf6a-4b59-8fc5-421767e7e269","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Implicit Regularization in Deep Learning May Not Be Explainable by Norms , volume =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:917f2ab237503bc9e14cb628daeffde3b8933967cd5b4448fcae8cdbe79d4db3","observation_id":"a058e7a5-c887-4933-903f-99fc5fda6e31","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"7th International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:26adb2fb7b5848eea64e5cdee941f7649c626cb804cc16d3f629377eba169b7e","observation_id":"39c216e1-65bd-44a9-bed3-c623241b1ed6","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:ea1b7b1f0fb7f407b4befe277fe53b0af2330d9f5486672e190aa8d0586b0d6e","observation_id":"9f1d7b38-2d4d-4f97-a06e-1895ad287f97","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Understanding Batch Normalization , volume =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:7c7215d74b31dca79f3ebb71441d93eee0dc9003fd1ff70a98923cb9d8fefa37","observation_id":"31e3c03f-699c-4b28-bad4-cae7a5140c9c","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:6ffd83c69efb6f6c03c1e1f436681f8f94886148092fd619238a4528e232e273","observation_id":"8878a09a-d513-4ed9-b57a-c5de0ae3f035","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"IEEE transactions on neural networks , volume=","venue":null,"work_id":null,"year":1994},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:dfa00725ec0d275de6987ce97e7fb1c9cb16424a74356209ff43c4ddbb149497","observation_id":"0e12908e-d975-44e6-bfce-0f5eac31bb8b","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:8fcc8740e0b647095f70e38da01f0b5aea8b1206921a0f1edc4d472907ef10a5","observation_id":"994cc456-1270-4819-8c1c-6f7953b4c7a2","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:2f7e367fc56b06be7016bc29972c864fbfdd05cab87cd762a443b9821968e04d","observation_id":"0a792df8-e5cc-4543-a2f4-b076b9f2d607","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2016 , eprint=","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:981246caa82c866f9268e8f8e7b1d484c58e4b69ade740ca8e20e07da935ba9d","observation_id":"e5d6f6f0-feb4-447a-8160-cd0606ee853d","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:bebe943458b694819769528d1771aa8b5255dc8c3b747fb1f5ffce86abb5e571","observation_id":"478f441d-3ffb-447d-9ff4-58e201b95a3b","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.4208/cicp.oa-2020-0165","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Dying ReLU and initialization: Theory and numerical examples","venue":"Communications in Computational Physics","work_id":"9f58281b-4b6a-412a-99fa-701342fce63d","year":2020},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:c03194fa55ffc5baad6240a8d5151e17abcf2e57ed9438a8812a9205185c9e83","observation_id":"063eea47-0102-43ab-99b5-5c3ed1f0d4a4","resolution":{"observed_at":"2026-06-25T23:28:41.040708Z","resolver_source":"doi","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2020 , pages=","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:25c3df092eb1e6b52f32f95a23dbf0ee1d83359e9e6c842670d7de2f1e79072a","observation_id":"a8c1e1b5-bea4-4ab6-9fa9-b5a6be131168","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2017 , archivePrefix=","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:ef41f887814f5ed076ccef397a4a0040ecbac5cdea474ff2ab67c51ddf00d14c","observation_id":"d79fad2f-dc2a-4121-8bdf-b2533643bba4","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"IEEE Transactions on Neural Networks and Learning Systems , volume=","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:e825a0623294b6394d1fcc306e284947f32b00577d898b3ef391c5fd29b5f900","observation_id":"9c759af5-749e-4812-a16d-3ec8eef8d857","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Measuring Saturation in Neural Networks , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:643959a1e6ab101eb496a51532919236d783d5365f55202b2621786ba8a6553f","observation_id":"631d1c5b-e5d8-4ef0-b3d0-6e93babb4811","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"2012 , publisher=","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:e59c8517ef2dd4cfb7a7f0c74aadccee0e3ab2067cf0bf4b3e12192f51bc3dbf","observation_id":"6c8fd99b-3500-4c39-958c-0e29418cd494","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":97,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:64fdb772ff4df5ddda4c668c1967297ea161d5b83fbeb7af4dad668c5792a17f","observation_id":"34bf17ed-0a82-4a8f-b585-f6f3a471e65f","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"Mechanization of Thought Processes, Proceedings of a Symposium Held at the National Physical Laboratory , pages=","venue":null,"work_id":null,"year":1958},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":98,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:be768c978a25ed138f907332aa8e941338220b416b8f533f5eb773f7027ddf35","observation_id":"dac5cc5a-ba50-4c20-a7c7-766fe2ff16d3","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":99,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:fd371aa9c836e64a247ad212269180841718c83bd41b54f493c9fdfe21e307c4","observation_id":"977ce818-9d3c-4602-acdb-74c4ee136218","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-25T23:26:57.147374Z","title":"SIGART Bull","venue":null,"work_id":null,"year":1977},"citing_paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?","version":1},"reference_index":100,"source":"arxiv_source","source_observed_at":"2026-06-25T23:26:57.147374Z"},"links":{"citing_paper":"/paper/2606.24752"},"observation_digest":"sha256:ae646cf59e235924f25cdf6c64a5bd7f87d50859b046db83f605cefe6a9c2fc2","observation_id":"2e4a81b1-dbf5-418a-8ff7-82f241f55e82","resolution":{"observed_at":"2026-06-25T23:26:57.147374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2606.24752","last_updated":"2026-06-23T16:14:52Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-07-06T23:59:18.313042Z","submitted_at":"2026-06-23T16:14:52Z","title":"Can Scale Save Us From Plasticity Loss in Large Language Models?"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":4,"parse_uncertain":1,"unresolved":95,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":228},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 100 of 228 outbound references and 0 inbound Pith citation observations for arXiv:2606.24752."}