{"as_of":"2026-08-21T17:13:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:308f1ea850275fc9772971df4569718069fbe23a9931180d9ad69058e136c8b4","coverage":[{"denominator":91,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":91,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T18:39:44.587149Z","state":"measured"},{"denominator":93,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":93,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T10:57:00.950744Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-16T07:00:43.314495Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.07907","snapshot_observed_at":"2026-08-16T10:57:00.950744Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2504.16920","last_updated":"2025-08-11T13:29:28Z","snapshot_observed_at":"2026-08-17T18:16:06.091886Z","submitted_at":"2025-04-23T17:45:42Z","title":"Summary statistics of learning link changing neural representations to behavior","version":3},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-16T10:57:00.950744Z"},"links":{"cited_paper":"/paper/2507.07907","citing_paper":"/paper/2504.16920"},"observation_digest":"sha256:a64521e9298ffed10ef7f110ad2cda0f05e8f7480cd726d9e26bedd9198ed912","observation_id":"16d859fb-611f-4494-970c-5230899e6163","resolution":{"observed_at":"2026-08-16T10:57:00.950744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"cited_work":{"arxiv_id":"2507.07907","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.07907","snapshot_observed_at":"2026-06-23T04:13:39.099253Z","title":"and Mori, F","venue":null,"work_id":"84342f08-df28-47a4-9715-dac7c383c077","year":null},"citing_paper":{"arxiv_id":"2602.04774","last_updated":"2026-05-08T16:24:57Z","snapshot_observed_at":"2026-08-15T04:58:48.280086Z","submitted_at":"2026-02-04T17:11:36Z","title":"Theory of Optimal Learning Rate Schedules and Scaling Laws for a Random Feature Model","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-16T06:58:38.927268Z"},"links":{"cited_paper":"/paper/2507.07907","citing_paper":"/paper/2602.04774"},"observation_digest":"sha256:7fff6192ad6873c19114aaf44bfcdfd5086b8d9583e7b3b93593b257f1148d7d","observation_id":"73d71fe1-5ee9-4800-bba5-413e9bac63e0","resolution":{"observed_at":"2026-06-23T04:13:39.099253Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2507.07907/citation-record","integrity":"/paper/2507.07907/integrity","json":"/paper/2507.07907/citation-record.json","paper":"/paper/2507.07907"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:37.482331Z","title":"Botvinick and Jonathan D","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:37.482331Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:4425708382f9c56e8373bad7bad9a2ad34d0d1688a9444f5f419d16849ddd9be","observation_id":"52f1f0de-6ba2-43e4-b595-1d24dcd9fe94","resolution":{"observed_at":"2026-08-06T18:39:37.482331Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:37.543303Z","title":"The easy-to-hard effect in human (homo sapiens) and rat (rattus norvegicus) auditory identification.Journal of Comparative Psychology, 122(2):132, 2008","venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:37.543303Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:a8c8d6e48798c8b7845b35acff9810b565cd7c7f3312366d22acb862b3792abf","observation_id":"ce6fb437-88a0-4be5-8244-434513096f3b","resolution":{"observed_at":"2026-08-06T18:39:37.543303Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:37.613934Z","title":"Springer Nature, 2019","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:37.613934Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:31b94d4b95b5b2336af035a952e6f272b49cd5f1d5d785a28da53a2edf5a58b2","observation_id":"592325bb-69bf-4633-a667-c3113e620610","resolution":{"observed_at":"2026-08-06T18:39:37.613934Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:37.658215Z","title":"Random search for hyper-parameter optimization.The journal of machine learning research, 13(1):281–305, 2012","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:37.658215Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:8ab1b5ab124ee13cf32ab98eef366131f46a077af42c386f270b9b71946abe9a","observation_id":"1b2de035-50c9-4737-802c-53ddfcf13ccc","resolution":{"observed_at":"2026-08-06T18:39:37.658215Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:37.718171Z","title":"Practical bayesian optimization of machine learning algorithms.Advances in neural information processing systems, 25, 2012","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:37.718171Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:ae9f132938b7245cb13a5e3fc5d89ae5e9a7cc1796ee2234894a1b5bb597ffbe","observation_id":"2d1f71b3-9005-4812-9ca8-464b6898e919","resolution":{"observed_at":"2026-08-06T18:39:37.718171Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:37.782205Z","title":"Gradient-based hyperparameter opti- mization through reversible learning","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:37.782205Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:01fb9d11c400d384752b0e42c7cc5151fc5cfa9dd45a499e727df0b17acd640d","observation_id":"283b93ba-10f0-4ac1-bbad-ae01dc638884","resolution":{"observed_at":"2026-08-06T18:39:37.782205Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:37.844820Z","title":"Model-agnostic meta-learning for fast adaptation of deep networks","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:37.844820Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:3be1a5cbc6aefb207cd3eca2cf602be94394b40f6a145522e1a2bd185595bea9","observation_id":"a3c4d8d9-d987-406e-8bc8-12d947cea1e2","resolution":{"observed_at":"2026-08-06T18:39:37.844820Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:37.910355Z","title":"Engel and C","venue":null,"work_id":null,"year":2001},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:37.910355Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:fe1751c591d79669ac7e41fd4537df2535879ee0f08cd6b88a887858e5b2243c","observation_id":"80c20d4e-f629-41a9-92bf-5689f86bef52","resolution":{"observed_at":"2026-08-06T18:39:37.910355Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:37.999550Z","title":"Optimal errors and phase transitions in high-dimensional generalized linear models.Proceedings of the National Academy of Sciences, 116(12):5451–5460, 2019","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:37.999550Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:81a1c0256378281843fe033f8a7c60d72cd93313f7d3e1e3a4760d5d89f70306","observation_id":"8c65f617-ffdc-46f6-8b0b-a75fccb66e41","resolution":{"observed_at":"2026-08-06T18:39:37.999550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:38.052672Z","title":"Learning curves of generic features maps for realistic datasets with a teacher-student model.Advances in Neural Information Processing Systems, 34:18137–18151, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.052672Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:bea42f12d8745d56cdf28834fc3da60fe38a3d2a96021870b3b1b57537896728","observation_id":"b837e9e2-c2e6-491f-a3ed-cf25c26f8fa1","resolution":{"observed_at":"2026-08-06T18:39:38.052672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:38.133041Z","title":"The role of regularization in classification of high-dimensional noisy Gaussian mixture","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.133041Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:7d5a26d42299a24800e8e49e46c52f40646b8e45097cbb3a7fefdab507291306","observation_id":"7267e221-9cf0-4dc3-b6dc-1b59a696f661","resolution":{"observed_at":"2026-08-06T18:39:38.133041Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:55.286991Z","title":"Gener- alisation error in learning with random features and the hidden manifold model","venue":null,"work_id":"4aa7892c-8df9-4360-a033-0c28aecb6e73","year":2020},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.191066Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:2ce9552ded63095c32ce534c787bfe71109594a9beb7036f7bd8d8acb4a22ecb","observation_id":"8e2fc506-c714-4380-aa70-effed298c75d","resolution":{"observed_at":"2026-08-06T18:39:55.339313Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:55.158887Z","title":"Dy- namics of stochastic gradient descent for two-layer neural networks in the teacher-student setup","venue":null,"work_id":"1d9baf62-0917-4aa8-813e-0c5a98f1b24f","year":2019},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.263656Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:cc23e4fede282a4cd4952ac73570333088f1864028351cf968999907c15189c3","observation_id":"b0a5135e-f6ce-42e8-97d1-1a22c5dabb02","resolution":{"observed_at":"2026-08-06T18:39:55.220479Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:55.035603Z","title":"Dynamical mean-field theory for stochastic gradient descent in gaussian mixture classification.Advances in Neural Information Processing Systems, 33:9540–9550, 2020","venue":null,"work_id":"c7aecbb6-358c-4f07-844c-bbb9929743bb","year":2020},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.343415Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:f2acc76b41b0070bdd74ed73bc4e6883b3175b8ff4b0c2506d3c23744c408648","observation_id":"2b0c1db0-c7b6-4134-b087-6dea6f7834ca","resolution":{"observed_at":"2026-08-06T18:39:55.082169Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:38.407723Z","title":"Self-consistent dynamical field theory of kernel evolution in wide neural networks.Advances in Neural Information Processing Systems, 35:32240–32256, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.407723Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:5356d3d0383a036977d1bc618c2b822262950d2f99cb808dbc75edea3d5be86c","observation_id":"2c7a530f-abdc-41cd-b929-7b1eb6f1353d","resolution":{"observed_at":"2026-08-06T18:39:38.407723Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:54.905819Z","title":"An analytical theory of curriculum learning in teacher-student networks","venue":null,"work_id":"d4a29635-84ce-4609-b8ee-f6846bbdb4ff","year":2022},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.472295Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:023f21b5ed1c56e69bb7770f9ada390ab713ea6c9481f81e52f29603d621691a","observation_id":"bb764f37-37f5-4788-b0d3-498ffb1eeda1","resolution":{"observed_at":"2026-08-06T18:39:54.945914Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:54.761508Z","title":"Why do animals need shaping? a theory of task composition and curriculum learning","venue":null,"work_id":"c67e94fe-c57d-44cc-9819-85faf5002823","year":2024},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.546431Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:750825b49f8dbad99723768dbb718abe1bdb98d35679d4f96f23fcdbc0dae7e9","observation_id":"f6ee6bca-79f5-4704-a1f5-eac4089e52a4","resolution":{"observed_at":"2026-08-06T18:39:54.821832Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:54.547704Z","title":"Curriculum learning in humans and neural networks, Mar 2025","venue":null,"work_id":"3f4f7942-59e6-4da1-94b9-ab610aefdf20","year":2025},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.627073Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:c8fa9cbe633be7613112b4cd9fcb37d525a6a371dac3b4747341e11234450ff3","observation_id":"133a8815-a4bb-45cd-a5f0-909fe21117b0","resolution":{"observed_at":"2026-08-06T18:39:54.637764Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:54.374521Z","title":"High-dimensional learning of narrow neural networks.Journal of Statistical Mechanics: Theory and Experiment, 2025(2):023402, 2025","venue":null,"work_id":"ab9308f8-ce16-4e1c-bdf1-8685ef653abf","year":2025},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.700719Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:d22a0f345f27df906214a753d753954f7e31d1f882787c4a8c154eefde49437f","observation_id":"5470782c-f97d-481f-b096-84a49b1b6541","resolution":{"observed_at":"2026-08-06T18:39:54.460619Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:38.772814Z","title":"Learning by on-line gradient descent.Journal of Physics A: Mathematical and general, 28(3):643, 1995","venue":null,"work_id":null,"year":1995},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.772814Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:b1a1d723a12d1d9861c345afd48e0f73d67d2e3209bcf4523d96a78a48412d06","observation_id":"0873e939-db2b-4cff-8722-1ea390868025","resolution":{"observed_at":"2026-08-06T18:39:38.772814Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:54.109268Z","title":"Exact solution for on-line learning in multilayer neural networks","venue":null,"work_id":"cf4fe529-798d-4f57-9779-5a8b554cdefb","year":1995},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.846040Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:921fb7f1f22add9788e9481647a83f9d50fdea8526d1557f9251b64b57602c54","observation_id":"ee986a4b-e6e9-4135-ac0c-83db804774b5","resolution":{"observed_at":"2026-08-06T18:39:54.232347Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:53.955030Z","title":"Analysis of on-line training with optimal learning rates.Physical Review E, 58(5):6379, 1998","venue":null,"work_id":"1397513d-b884-4a04-9577-15a576a0cc88","year":1998},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.918212Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:4be6efa3ce102b680e234c657a26f3929fd8ba8f11b7a38ba9cc982601c89cb4","observation_id":"ddce2ffb-9f31-4ac6-9842-d09b1f2d0117","resolution":{"observed_at":"2026-08-06T18:39:54.033205Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.19919","last_updated":"2024-07-15T12:07:03Z","snapshot_observed_at":"2026-08-21T05:56:43.954455Z","submitted_at":"2023-10-30T18:29:26Z","title":"Meta-Learning Strategies through Value Maximization in Neural Networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.19919","snapshot_observed_at":"2026-08-06T18:39:38.993799Z","title":"Meta-learning strategies through value maximization in neural networks.arXiv preprint arXiv:2310.19919, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:38.993799Z"},"links":{"cited_paper":"/paper/2310.19919","citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:2c3ac2ac6d83cb2f876cf74b6c183db697ab85abe5d39eb5e444e99ccc2cc60f","observation_id":"3ad444c6-b75d-4f91-a3f6-00d765b92091","resolution":{"observed_at":"2026-08-06T18:39:38.993799Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:39.054136Z","title":"Courier Corporation, 2004","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.054136Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:8272093a4a7eaba5274c512eeb3d5af3c6a7089dac93751705737bb652860232","observation_id":"c9872b6a-b8ab-4ab2-845c-04d8e6ae3952","resolution":{"observed_at":"2026-08-06T18:39:39.054136Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:39.120325Z","title":"SIAM, 2010","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.120325Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:9d57bb98996cfdf26105c52d32720221b2ccb261b9eda4e4b5739736931a919d","observation_id":"31a490cd-0281-4cc4-a5b0-f3c92416853a","resolution":{"observed_at":"2026-08-06T18:39:39.120325Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:53.633953Z","title":"Practical recommendations for gradient-based training of deep architectures","venue":null,"work_id":"5d01ebdc-021a-4abc-b845-172a310a0f8f","year":2012},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.186637Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:45b5c2e9100515a1a1760e5ff26a6e0bb0a88e6f47382123b13695055a3d7fcf","observation_id":"f9c912ea-a254-485c-8d0f-aae36345d3a3","resolution":{"observed_at":"2026-08-06T18:39:53.788367Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:53.456397Z","title":"Why warmup the learning rate? underlying mecha- nisms and improvements.Advances in Neural Information Processing Systems, 37:111760–111801, 2024","venue":null,"work_id":"2c1f18c7-d13b-4eda-b59b-7afa5a4f0372","year":2024},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.222636Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:316fe615da1de7b9ba3e6abb88a593a686fc6cfebf0ce1c0fc834713970758c8","observation_id":"a5c71fdd-f0fc-4144-b5b3-3aa5f6a141d0","resolution":{"observed_at":"2026-08-06T18:39:53.532450Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:39.272763Z","title":"Sgdr: Stochastic gradient descent with warm restarts","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.272763Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:341740a6d20403017485fa09f43d6f3228795779832851265798f0d9898b2812","observation_id":"30a7c17c-b881-4505-bb64-0e5b9baf9795","resolution":{"observed_at":"2026-08-06T18:39:39.272763Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:53.240573Z","title":"Online learning rate adaptation with hypergradient descent","venue":null,"work_id":"48255eac-9ea1-4ffc-9706-dd9b5d1fb1bb","year":2018},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.346949Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:ea98df9aeda1067fc7e3bfdc21f110cbd485e65992294b865bcdea93ee30ae94","observation_id":"1f596d8d-8405-488a-ad0d-f0e0f9dbaa97","resolution":{"observed_at":"2026-08-06T18:39:53.334964Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:53.029480Z","title":"Globally optimal parameters for on-line learning in multilayer neural networks.Physical review letters, 79(13):2578, 1997","venue":null,"work_id":"060a7af1-2d84-44f5-b88a-0f7c1558b0d1","year":1997},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.399890Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:fe069aa92d82947a090f8053ab6eb0a1a810bd6695cbbebdfc862a56c758c4a4","observation_id":"c3d0100f-2999-44ca-a94b-e7440979f63e","resolution":{"observed_at":"2026-08-06T18:39:53.116098Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:52.827119Z","title":"Optimization of on-line principal component analysis.Journal of Physics A: Mathematical and General, 32(22):4061, 1999","venue":null,"work_id":"a3bb18c7-b67c-4072-8b3f-f36710a90df7","year":1999},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.459449Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:3e3bd18e12289e611c342306f08a40d43456ba86270d065c627a407ba3aa29e5","observation_id":"087bf65e-18c5-4e4c-9e32-d696057851df","resolution":{"observed_at":"2026-08-06T18:39:52.908411Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2202.04509","last_updated":"2022-02-09T15:15:39Z","snapshot_observed_at":"2026-08-19T08:34:16.554353Z","submitted_at":"2022-02-09T15:15:39Z","title":"Optimal learning rate schedules in high-dimensional non-convex optimization problems","version":1},"cited_work":{"arxiv_id":"2202.04509","doi":null,"metadata_source":"pith","pith_arxiv_id":"2202.04509","snapshot_observed_at":"2026-08-06T18:39:45.044616Z","title":"Optimal learning rate schedules in high-dimensional non-convex optimization problems","venue":"cs.LG","work_id":"fe014e07-cd2b-4509-8ad6-15ce7a9a78c6","year":2022},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.539635Z"},"links":{"cited_paper":"/paper/2202.04509","citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:cbc7fdbba36f72c2260e4ce62560b0d9327daebbdc4270a9e419fd82f258a56f","observation_id":"f25371e4-785c-4b18-b5ac-b47287052cdb","resolution":{"observed_at":"2026-08-06T18:39:45.097555Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:52.619568Z","title":"Optimal protocols for contin- ual learning via statistical physics and control theory","venue":null,"work_id":"e7cad6de-957c-44c0-9d48-2a18fb98c0ff","year":2025},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.612502Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:4db99c1c4814c6364fa7f1a3d711037292a99db8c9c589b05d7d95d3a7fc1d1e","observation_id":"ec5a4f11-f90d-43bd-ba8c-42ac9317ecb0","resolution":{"observed_at":"2026-08-06T18:39:52.715039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:52.405309Z","title":"Don’t decay the learning rate, increase the batch size","venue":null,"work_id":"c9c96119-bb7e-4541-a87b-b01fda114a5d","year":2018},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.683558Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:4b4ed9859d5fd5fb170f9909f5e0852f5fe9812ca66b5837fe706f598805a154","observation_id":"3a18c318-9fcc-4246-8672-8a54379891f8","resolution":{"observed_at":"2026-08-06T18:39:52.507181Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:52.202870Z","title":"Continual learning in the teacher-student setup: Impact of task similarity","venue":null,"work_id":"5dc899b6-2e14-4f22-b244-6ded12e78f9c","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.747483Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:ecb05a3b8b77d44d21ac8fdfdabc17ec6cc223e3a7b0c70bfe3f0116654569d7","observation_id":"2aedaa07-b91c-413e-9a29-baa686cf8fa1","resolution":{"observed_at":"2026-08-06T18:39:52.284503Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:39.809701Z","title":"How catas- trophic can catastrophic forgetting be in linear regression? InConference on Learning Theory, pages 4028–4079","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.809701Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:1f547e77a8ed3f30a2f0c1368799941fa2e5fd5ec9e9f705db295c1da10ae5dc","observation_id":"c14331a8-7684-4f47-afef-fd2293af8acf","resolution":{"observed_at":"2026-08-06T18:39:39.809701Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10315","last_updated":"2025-01-26T04:27:17Z","snapshot_observed_at":"2026-08-16T13:34:20.244967Z","submitted_at":"2024-07-14T20:22:36Z","title":"Order parameters and phase transitions of continual learning in deep neural networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10315","snapshot_observed_at":"2026-08-06T18:39:39.875669Z","title":"Order parameters and phase transitions of continual learning in deep neural networks.arXiv preprint arXiv:2407.10315, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.875669Z"},"links":{"cited_paper":"/paper/2407.10315","citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:7fb177085afe4455cf65c3e86a987c334526a847a3c151d74bf534d41ef5696c","observation_id":"e6c319e7-50f3-4aa4-a436-b51d324a83ea","resolution":{"observed_at":"2026-08-06T18:39:39.875669Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:52.015485Z","title":"Provable advantage of curriculum learn- ing on parity targets with mixed inputs","venue":null,"work_id":"b323c836-4eb5-4f59-9963-1f6e5bb68d1c","year":2023},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:39.941072Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:fe83e865013ab7d7eaf3b06ff8992dd1870e5ffca77d44e3fb3ed0cbf87d695d","observation_id":"11d5977d-4ec7-4293-9992-cb062df1f4ed","resolution":{"observed_at":"2026-08-06T18:39:52.104731Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:51.752336Z","title":"Restoring data balance via generative models of t-cell receptors for antigen-binding prediction.bioRxiv, pages 2024–07, 2024","venue":null,"work_id":"84b6fcaf-850d-4752-9eb0-12ab2c4b3d45","year":2024},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.026989Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:f70e564601119cf3ce36606eeb5dffbd730dd52a63244302fd03493f75d8802b","observation_id":"2a5483a1-e3b9-4011-bb3f-db4f63e3d885","resolution":{"observed_at":"2026-08-06T18:39:51.907605Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:51.561570Z","title":"Bias-inducing geometries: exactly solvable data model with fairness implications","venue":null,"work_id":"550e537c-20e5-4fb7-8229-d80cf8a50469","year":2024},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.106104Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:491c1aededd7dca3638fefee5220d37700ac2a101b34f79866fe4020149cda83","observation_id":"320d9c14-234c-4bd5-977a-eb5bcde07cc6","resolution":{"observed_at":"2026-08-06T18:39:51.649899Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:51.373251Z","title":"Bias in motion: Theoretical insights into the dynamics of bias in sgd training","venue":null,"work_id":"d9cc8af7-c80d-4d13-b3cd-f818d8a77a6b","year":2024},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.179025Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:59e12e749c86414b924729d18062397ff056a406afd7e800cd5b435855da1ed9","observation_id":"e6d109c5-9762-4929-b86f-0294b1ec16af","resolution":{"observed_at":"2026-08-06T18:39:51.476035Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:40.240849Z","title":"Dropout: a simple way to prevent neural networks from overfitting.The journal of machine learning research, 15(1):1929–1958, 2014","venue":null,"work_id":null,"year":1929},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.240849Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:527ff714b4ec7fa0e183c8d89a363fb763be98108ab1aeea21395bdda91db0dc","observation_id":"dcd81077-e14c-4a6a-89dc-0e939e936305","resolution":{"observed_at":"2026-08-06T18:39:40.240849Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:51.144163Z","title":"Curriculum dropout","venue":null,"work_id":"635fd5bb-9c9c-4078-bea9-865714f7d147","year":2017},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.312494Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:3159e1fc8606ea1724857078929c74bf57227eff3fbd3c11a4f43dbe869287b2","observation_id":"538afb0f-065c-400e-a146-1bc58b0226f7","resolution":{"observed_at":"2026-08-06T18:39:51.222926Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:50.967932Z","title":"Dropout reduces under- fitting","venue":null,"work_id":"255e75c5-f293-4020-9a81-6ba2a1d6dacf","year":2023},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.370839Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:57f3ce12a4916e857a6c020a8e81fe4256d0d4eb7da305c2eb3a61eaea6646ff","observation_id":"607a7d12-7da2-4552-9f2a-1c16f332b9f8","resolution":{"observed_at":"2026-08-06T18:39:51.060653Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:50.796212Z","title":"Analytic theory of dropout regularization.Phys","venue":null,"work_id":"e7fa13b0-3522-4628-bbe7-fa9e8136aa62","year":2025},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.434947Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:976a8290494927d116df53f46deaadbe1b9416566d7038e9601dc5fd24e0cce4","observation_id":"51ceb9d1-f9ee-4d53-8e8a-b3cb791b24b2","resolution":{"observed_at":"2026-08-06T18:39:50.861781Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:40.508200Z","title":"Outrageously large neural networks: The sparsely-gated mixture-of-experts layer","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.508200Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:a297ccba13fe91d874204c0afd75dba4f93431b6bdbf251d603e3f3385c496f7","observation_id":"6140ae34-ea7b-496d-92db-0a352f367168","resolution":{"observed_at":"2026-08-06T18:39:40.508200Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:50.626842Z","title":"Learning phrase representations using RNN encoder–decoder for statistical machine translation","venue":null,"work_id":"6a71a2a8-84be-4700-8261-c8f8994f7511","year":2014},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.577747Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:ada869c46e9a6a97e4668cb40da0f75a8eda51094542bc4bc26878a4491447af","observation_id":"5de31c2d-e066-4e91-b718-6ac9da08e044","resolution":{"observed_at":"2026-08-06T18:39:50.687588Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:50.355721Z","title":"Gated linear networks","venue":null,"work_id":"df765f4e-5143-477a-a474-7ceff027fdc4","year":2021},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.626790Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:6f6f1e436284ccae4ab0d639b19c2ba151bf164f65993b154dfb964105bac8f2","observation_id":"17d4c5db-05c7-4a3b-a5f7-62d8ff77bdb1","resolution":{"observed_at":"2026-08-06T18:39:50.502074Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:50.217442Z","title":"Globally gated deep linear networks.Advances in Neural Information Processing Systems, 35:34789–34801, 2022","venue":null,"work_id":"bb6bc083-4221-4a25-9ae5-6d8e3c795e96","year":2022},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.676027Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:8115c1ba424186d077b8396dbdaf12f1c710d0e5054224f94ccfc5522c8b3290","observation_id":"23613871-7513-4824-a0f1-2086228c8d96","resolution":{"observed_at":"2026-08-06T18:39:50.274973Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:50.056400Z","title":"The neural race reduction: Dynamics of abstraction in gated networks","venue":null,"work_id":"11c2dbf3-f021-4228-b25e-25ddc1035fa5","year":2022},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.739092Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:16d79e5227521fb4a3826666d3b7ff94b6bca32989d889b0afe3496e3ee4bd16","observation_id":"8316a1bd-b9c5-4840-9f6b-e319d5022bff","resolution":{"observed_at":"2026-08-06T18:39:50.127031Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:49.898414Z","title":"Nonlinear classification of neural manifolds with contextual information.Physical Review E, 111(3):035302, 2025","venue":null,"work_id":"d81b2253-ef26-4e52-9e75-227bb4fd465e","year":2025},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.798765Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:bb555f6c0474f2459de3b5831423189f5bf62e3cd3152cadba964e6323bb6167","observation_id":"c8888dbb-f764-4364-bfe2-a1df5aa6b25c","resolution":{"observed_at":"2026-08-06T18:39:49.957601Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:40.868401Z","title":"Attention is all you need.Advances in neural information processing systems, 30, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.868401Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:a0ef502c02d5d43468595fe174578a2708d797c8490f07011f979b5768444498","observation_id":"6fc760d4-40aa-445b-8776-cd3025c1ef6e","resolution":{"observed_at":"2026-08-06T18:39:40.868401Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:40.943435Z","title":"Efficient content-based sparse attention with routing transformers.Transactions of the Association for Computational Linguistics, 9:53–68, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:40.943435Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:b453a6e8ae67cbde1e47346ed65fded2422a5ed26f99e6a3fce405452fcb57ec","observation_id":"a6ecd214-c1df-447a-9777-20e2154c23ee","resolution":{"observed_at":"2026-08-06T18:39:40.943435Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:49.725651Z","title":"Adaptive atten- tion span in transformers","venue":null,"work_id":"ca64d697-5139-401e-9968-63de581a3ca6","year":2019},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.034785Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:744dd1e9b1080e98726f4a3ea6cdce1381bfc1b0d2e155ac6d604b9571243e0c","observation_id":"6dfaa829-cd7f-442b-92cd-a45c03418a9f","resolution":{"observed_at":"2026-08-06T18:39:49.779690Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:41.097955Z","title":"Are sixteen heads really better than one?Advances in neural information processing systems, 32, 2019","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.097955Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:c43fe2286cd30fbd2f10c313f6218db0de08a6f79381d33b9956a1819109ede7","observation_id":"dbe94e78-2aaa-424b-a747-179edcd0cdb8","resolution":{"observed_at":"2026-08-06T18:39:41.097955Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:49.492053Z","title":"The effects of information order and learning mode on schema abstraction.Memory & cognition, 12(1):20–30, 1984","venue":null,"work_id":"5530d6f3-8817-4fc9-a576-598b3f18806c","year":1984},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.173408Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:65a83a52f28c8d205234d2654314882df99537635ef6822ceaaed60e001e0af3","observation_id":"1c6716b1-6ac6-4d12-9d27-3abf771821fa","resolution":{"observed_at":"2026-08-06T18:39:49.615918Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:49.292705Z","title":"When does fading enhance perceptual category learning? Journal of Experimental Psychology: Learning, Memory, and Cognition, 39(4):1162, 2013","venue":null,"work_id":"48fe9154-1156-4ec0-a1a4-37b16f0abbf6","year":2013},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.252558Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:c6b10ad9c687328d5dbc5104455bc86f065eca1bc9b6f45e129886bd871ca4ba","observation_id":"baa6e266-40be-47a4-85d0-b0e451067fcc","resolution":{"observed_at":"2026-08-06T18:39:49.362898Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:41.339550Z","title":"Curriculum learning","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.339550Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:f051ac1465ef2504fdf1f6a7ea891aad2aff69c85c8dd37faafbc1b3d7a30aa4","observation_id":"4b2db901-1f1d-4894-bf0f-0e9857e09656","resolution":{"observed_at":"2026-08-06T18:39:41.339550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:41.395215Z","title":"A survey on curriculum learning.IEEE transactions on pattern analysis and machine intelligence, 44(9):4555–4576, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.395215Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:43b67ec38c5b9e0260a5d02c94d3a7f685f18838f3a9bf30527a23eafb6b9c68","observation_id":"c1ba915b-b44d-4ab4-971e-d262f5fbcba4","resolution":{"observed_at":"2026-08-06T18:39:41.395215Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:49.071482Z","title":"On the power of curriculum learning in training deep networks","venue":null,"work_id":"0fa7c1ef-e30a-4396-98b4-60b8b7cdc679","year":2019},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.517224Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:62e18ba50afde376395a436d4f39f3e1acda2518546a39aab22a54441b0b59f7","observation_id":"9498c7bb-2b60-4853-a3c3-ff384c817cc5","resolution":{"observed_at":"2026-08-06T18:39:49.159581Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:48.957625Z","title":"When do curricula work? InInternational Conference on Learning Representations, 2021","venue":null,"work_id":"874d66e0-60a7-4a9a-a74d-cc536b6b108d","year":2021},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.616993Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:707fc09e1b275167b6d02bdd896dbcba6bcc5f6f7ffab61d53ea844ae5a3c845","observation_id":"bf678ef4-9b42-46d8-9c67-cc00b275df85","resolution":{"observed_at":"2026-08-06T18:39:49.023007Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:48.834625Z","title":"Theory of curriculum learning, with convex loss functions","venue":null,"work_id":"3956d647-2c5b-4ba0-b566-75072e351243","year":2020},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.704001Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:172918a4b40a1e7c52e631c7a6052c7c888513a17fc1bf76d59e34a353bcecc5","observation_id":"18e144f6-6b0b-47fe-81f5-d9722b119d36","resolution":{"observed_at":"2026-08-06T18:39:48.887455Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:48.701850Z","title":null,"venue":null,"work_id":"d3bae9b9-4e28-47af-9df3-3543b73cf6c7","year":2024},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.824329Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:b7786806f20b817d7cadea29c47442fb7e7bec8f7636af680e84c31c65edcdf3","observation_id":"2a18e7ce-4785-4060-b52e-e7469709107b","resolution":{"observed_at":"2026-08-06T18:39:48.745524Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:48.558949Z","title":"Curriculum learning by optimizing learning dynam- ics","venue":null,"work_id":"5f0d1b5f-5d04-4a46-9016-3ac1f55e3c89","year":2021},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:41.918014Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:da32ab3803ffcc48b9295ff9bb99ffb2d4cbca4c4a2f51323294cb0a52a2ca7c","observation_id":"294fba1d-ef4d-476f-88c3-937ad5604d3a","resolution":{"observed_at":"2026-08-06T18:39:48.620820Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:48.398184Z","title":"Rennie, Vaibhava Goel, and Samuel Thomas","venue":null,"work_id":"b5a8b895-56cc-4f75-9cea-46c31dd623b4","year":2014},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:42.038549Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:140fe6dfa6fcce105f76f49b749ff835eb9abdd5e3e3616f3c7394def82219f6","observation_id":"15407a7a-1f06-4e1d-804d-c1fb58140c43","resolution":{"observed_at":"2026-08-06T18:39:48.433448Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:48.240584Z","title":"Extracting and composing robust features with denoising autoencoders","venue":null,"work_id":"84d11c66-fed0-434e-b54b-ef2befb6f472","year":2008},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:42.154847Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:0f400c5ecf000b74a50bf0bb8a57326161f0218892af768dba1091c236a37b2f","observation_id":"0e0b0eec-91bb-4204-b93b-00053f386750","resolution":{"observed_at":"2026-08-06T18:39:48.283751Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:42.280382Z","title":"Denoising diffusion probabilistic models","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:42.280382Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:9c40c269020d286e7d0036bdb861b4646b63a7d0a5996ecbd4c6f569f24fdd3a","observation_id":"0edaf180-9986-44f7-b8c5-894d63ab8c11","resolution":{"observed_at":"2026-08-06T18:39:42.280382Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:48.054372Z","title":"Learning dynamics of linear denoising au- toencoders","venue":null,"work_id":"66e7c98e-0db0-423a-ae32-c02484421c14","year":2018},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:42.393848Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:e01da61500d109f190e7b83dc611163925a339df428b418168ac15cfc9dec8b8","observation_id":"110e0016-9c47-4476-a2fd-81d6369c485c","resolution":{"observed_at":"2026-08-06T18:39:48.147690Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:47.889897Z","title":"High-dimensional asymptotics of denoising autoencoders.Ad- vances in Neural Information Processing Systems, 36:11850–11890, 2023","venue":null,"work_id":"9c2fb7d3-8029-4e57-aad0-3da4d9c49902","year":2023},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:42.518648Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:f9dedf1bc817e0def81457d5623683fb8b718f175943e7c027f316625a3a740b","observation_id":"61128815-a93e-4f82-b686-8b7c7bd73b59","resolution":{"observed_at":"2026-08-06T18:39:47.991457Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:47.777006Z","title":"A solvable model of learning generative diffusion: theory and insights.Advances in Neural Information Processing Systems, 38:5253–5296, 2026","venue":null,"work_id":"d221308a-2831-4611-abd1-10084718afa4","year":2026},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:42.633370Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:1a25a31fcdbfcd4bb4ad6beb27422df64bc635bf0189b54f2ae02f126d29eef9","observation_id":"f4bf8ce0-a759-43c8-861f-4defec3a0c28","resolution":{"observed_at":"2026-08-06T18:39:47.820017Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:47.640352Z","title":"Geras and Charles Sutton","venue":null,"work_id":"54ac36df-6a2c-43fe-9ac7-1eea2db49ecf","year":2015},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:42.726970Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:1806f01a25e8407ffbc72f1fc8a95fb2851c64a845948ce629042a0980855f6a","observation_id":"200773d0-0722-41d5-ab0a-596c529269c1","resolution":{"observed_at":"2026-08-06T18:39:47.688160Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:47.538014Z","title":"Non-uniform timestep sampling: Towards faster diffusion model training","venue":null,"work_id":"335d758c-e988-444a-b6ee-c0b100c8f0e7","year":2024},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:42.811578Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:3f59cb4ba02a835412d6f356e97f013a40f664ee7e34bdd366e7c5e81a7e6f92","observation_id":"c01921b0-d2c3-4c48-9496-f01b4fd9dacb","resolution":{"observed_at":"2026-08-06T18:39:47.594273Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:47.419823Z","title":"Marginalized denoising auto- encoders for nonlinear representations","venue":null,"work_id":"504b33a3-7fa5-479f-a739-a9911db5759e","year":2014},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:42.922180Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:58f407d25b2295e1b9379359074dbada6d69412897114a96ecc567f474873d67","observation_id":"ce968619-6731-4813-b25b-bcf84f996514","resolution":{"observed_at":"2026-08-06T18:39:47.440290Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:47.345806Z","title":"Modeling the influence of data structure on learning in neural networks: The hidden manifold model.Physical Review X, 10(4):041044, 2020","venue":null,"work_id":"d969f486-56a7-4b7e-b565-d306c89be8b7","year":2020},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:43.018635Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:a2f8e0278387d13f6d58a60e7aaedede54d9537bce8e26c1d78bb5bf821b233f","observation_id":"2910953f-2e6c-45a2-89f5-a5361f93c10b","resolution":{"observed_at":"2026-08-06T18:39:47.386593Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:47.215202Z","title":"Classification of heavy-tailed features in high dimensions: a superstatistical approach","venue":null,"work_id":"b0887323-145f-4da1-9829-93c790937694","year":2023},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:43.107114Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:6bccf52ed7a00c19ce7cf6589bc9d7cdcc423b25cd399a4b64f02b89190e35c2","observation_id":"aefd29e1-dc19-43bf-8fdd-62a6a03d64a1","resolution":{"observed_at":"2026-08-06T18:39:47.258091Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:47.091351Z","title":"Wakhloo, Tamara J","venue":null,"work_id":"b4062e04-8d3c-4de1-9b83-e2a0be1ba649","year":2023},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:43.220206Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:0e754003a1238092198a295cde42f86005a2ccfb0ef4e7ff3af712bf93e20f90","observation_id":"7b010d21-1895-4ea0-8d82-3cdccc35e717","resolution":{"observed_at":"2026-08-06T18:39:47.126682Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19644","last_updated":"2025-07-25T19:34:35Z","snapshot_observed_at":"2026-08-17T18:02:32.233346Z","submitted_at":"2025-07-25T19:34:35Z","title":"Hierarchical clustering and dimensional reduction for optimal control of large-scale agent-based models","version":1},"cited_work":{"arxiv_id":"2507.19644","doi":null,"metadata_source":"pith","pith_arxiv_id":"2507.19644","snapshot_observed_at":"2026-08-06T18:39:44.668560Z","title":"Hierarchical clustering and dimensional reduction for optimal control of large-scale agent-based models","venue":"math.OC","work_id":"b5baf479-fde8-45bc-aa3e-8dd266a0e76e","year":2025},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:43.341374Z"},"links":{"cited_paper":"/paper/2507.19644","citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:b676124a4c3f6a7fa0211fcba72eb4d48a8dff050bfded7cf6073bdae877cc4f","observation_id":"30545258-c068-4a1b-8c4b-1a8b0f234001","resolution":{"observed_at":"2026-08-06T18:39:44.876307Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:46.961899Z","title":"Some mathematical problems arising in connection with the theory of optimal au- tomatic control systems","venue":null,"work_id":"53e3bb63-070e-42b7-9910-fc64428fd161","year":1957},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:43.416490Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:ba8f0b1f391dac61552d40c9633f113798ab3c6526ff4829567a15f8f8d5de26","observation_id":"3e1593d8-677b-4137-b4c1-aade41f72657","resolution":{"observed_at":"2026-08-06T18:39:47.019611Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:46.855674Z","title":"Casadi: a software framework for nonlinear optimization and optimal control.Mathematical Programming Computation, 11:1–36, 2019","venue":null,"work_id":"6c21a375-f067-45e8-8b1d-9207e8cd3092","year":2019},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:43.507966Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:5baa48e60807d93650c9ecfba258d035d8a8970d31449ae66b5ccff4125d383a","observation_id":"adbb0dba-1874-44a2-86b3-9135c9391f32","resolution":{"observed_at":"2026-08-06T18:39:46.913846Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:46.720220Z","title":"training time","venue":null,"work_id":"05dc7082-b62f-426f-8c43-f33b3ef6c101","year":2025},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:43.608880Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:948dc919d3e6d7daf96fe60942b3ee22093156c7585ffed9d2ac968f05f63366","observation_id":"b199d466-69e9-4813-9fc9-b1c5b9e5cde3","resolution":{"observed_at":"2026-08-06T18:39:46.784038Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:46.624176Z","title":null,"venue":null,"work_id":"8c0f4394-0cb7-4398-ab44-0615ab3e617f","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:43.743647Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:cbe00d4eb408a90240ac15fc6c835650d9d5e8758eed61755eb834793f070a73","observation_id":"af963b0e-fd76-4152-be97-1b3a6528f2f0","resolution":{"observed_at":"2026-08-06T18:39:46.681522Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:46.491063Z","title":null,"venue":null,"work_id":"39a55409-a1b4-4893-b468-b65219a7641c","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:43.867005Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:906df62455fdf16533179ee7a5d5c508c363722d6f55eabbaa09eb74e6356c7b","observation_id":"e44eb0cb-b26f-4db0-ba51-4d2df4eb1536","resolution":{"observed_at":"2026-08-06T18:39:46.548044Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:46.371202Z","title":null,"venue":null,"work_id":"fa10e873-6ec7-4ea3-9b7e-a42aeabdc170","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:43.988866Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:4ae7e669db0d95b07805f4ac7df12094997ed4899710195a2946670e59f4f69d","observation_id":"2be27758-0ade-4d3a-a750-4890e34a2ea0","resolution":{"observed_at":"2026-08-06T18:39:46.417450Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:46.248999Z","title":null,"venue":null,"work_id":"892503f7-54ef-484b-8f91-a102ee69f9ad","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:44.106634Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:29d1d33585f3dc9475c62e7deed6b3372bec59c875f38472cc1bb998e2ecc067","observation_id":"fe6c4fc1-5a08-4caf-b2ca-6e1bdb5700fc","resolution":{"observed_at":"2026-08-06T18:39:46.294812Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:46.083413Z","title":null,"venue":null,"work_id":"ec83c61b-8326-45b9-92be-b11ba14f9351","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:44.230407Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:2c9e5498a20b888079e357252bdbfb88062001a8e19c88537d52fb2e1a12cfd7","observation_id":"6a8acd41-21fd-42d6-bbe3-672d85516348","resolution":{"observed_at":"2026-08-06T18:39:46.186570Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:45.983253Z","title":null,"venue":null,"work_id":"2d7abf28-1434-4278-9eda-f749011de41e","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:44.327050Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:9e1bb30f8bf216dd862329e44c686d7322feaa530425003ce28f6fc092f2009e","observation_id":"4c114b19-0231-43e0-b9f1-2ef556e3654f","resolution":{"observed_at":"2026-08-06T18:39:46.043042Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:45.853823Z","title":null,"venue":null,"work_id":"c187fcfb-f7ed-4a1f-9725-08e1f76dd737","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:44.408632Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:ba07fc1ca32cfc58227bd47fb8ed473f21247ef1dbaa32dcfc24e2a103ff3e82","observation_id":"1c6f30cc-62c7-4340-a09a-88c8ce9801d2","resolution":{"observed_at":"2026-08-06T18:39:45.894187Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:45.736233Z","title":null,"venue":null,"work_id":"07434975-908e-490d-877e-2f017bc229e2","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:44.455229Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:6b15c83ae6b5104ed5ca227d02a3c466e2f4f92e1798b06fd838856b09031a23","observation_id":"8784a780-b126-4670-9859-70f4b9028ab2","resolution":{"observed_at":"2026-08-06T18:39:45.779867Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:45.576458Z","title":null,"venue":null,"work_id":"0c148ec2-8124-4470-a95f-e4e6424803e5","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:44.493638Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:a97ca54de35668184d761e6209b745d18ee2aae8525c681628dd5f903f0505dc","observation_id":"ff6fc80c-f053-4ba9-a721-afb5321064e0","resolution":{"observed_at":"2026-08-06T18:39:45.664158Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:45.357388Z","title":null,"venue":null,"work_id":"3a44bf40-f460-4b8f-b946-5491c881b919","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:44.537064Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:1db5df204fbbdc2b56afa94c1f9e137aa242b1ffbc162117cf5162a2260c8c6f","observation_id":"0d7b4d1e-0331-4f24-9fcb-83f3ecd6f128","resolution":{"observed_at":"2026-08-06T18:39:45.508078Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:39:45.242561Z","title":"We typically choose the damping parameterγ damp >0.9","venue":null,"work_id":"4205f4c8-4ace-4315-8323-0bd7c9be64db","year":null},"citing_paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning","version":2},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:44.587149Z"},"links":{"citing_paper":"/paper/2507.07907"},"observation_digest":"sha256:fc1c3d8d1b1e8a5f772d8151b024f7bde68520957e48d61a20b9bdc2f95e60f7","observation_id":"8b114ec7-8a9c-4d89-8fa3-4a01e983e5e8","resolution":{"observed_at":"2026-08-06T18:39:45.283472Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2507.07907","last_updated":"2026-06-22T15:49:10Z","latest_version":2,"primary_category":"cond-mat.dis-nn","snapshot_observed_at":"2026-08-18T16:51:51.683910Z","submitted_at":"2025-07-10T16:39:46Z","title":"A statistical physics framework for optimal learning"},"reference_resolution":{"displayed":91,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":38,"verified_exact":2,"verified_fuzzy":51},"total_outbound_references":91},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 21 August 2026, this Paper Citation Record lists 91 of 91 outbound references and 2 inbound Pith citation observations for arXiv:2507.07907."}