{"as_of":"2026-08-12T20:20:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2fa3dd4f5dddc22af74b5ec6ddaa242b261b720158ad59212d9f25246fddbacb","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":8,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":8,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":8,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":8,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T12:32:47.947181Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T07:06:44.997380Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2305.18270","last_updated":"2025-06-03T19:12:56Z","snapshot_observed_at":"2026-08-10T00:04:17.648426Z","submitted_at":"2023-05-29T17:43:44Z","title":"How Two-Layer Neural Networks Learn, One (Giant) Step at a Time","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.18270","snapshot_observed_at":"2026-08-12T12:32:47.947181Z","title":"How two-layer neural networks learn, one (giant) step at a time, 2023 b","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.17201","last_updated":"2024-11-26T08:14:48Z","snapshot_observed_at":"2026-08-12T12:21:30.416946Z","submitted_at":"2024-11-26T08:14:48Z","title":"Learning Hierarchical Polynomials of Multiple Nonlinear Features with Three-Layer Networks","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-12T12:32:47.947181Z"},"links":{"cited_paper":"/paper/2305.18270","citing_paper":"/paper/2411.17201"},"observation_digest":"sha256:e3a2a6fe41c7736ef46ddc86d61ae108b8b70ebfe021156bf1fcefd0bb907bad","observation_id":"30473510-18fe-4bff-b315-a1aa1c0189ab","resolution":{"observed_at":"2026-08-12T12:32:47.947181Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18270","last_updated":"2025-06-03T19:12:56Z","snapshot_observed_at":"2026-08-10T00:04:17.648426Z","submitted_at":"2023-05-29T17:43:44Z","title":"How Two-Layer Neural Networks Learn, One (Giant) Step at a Time","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.18270","snapshot_observed_at":"2026-08-11T04:36:32.044680Z","title":"Learning two-layer neural networks, one (giant) step at a time","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.18756","last_updated":"2024-12-25T03:03:58Z","snapshot_observed_at":"2026-08-12T19:52:15.938751Z","submitted_at":"2024-12-25T03:03:58Z","title":"Towards a Statistical Understanding of Neural Networks: Beyond the Neural Tangent Kernel Theories","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-11T04:36:32.044680Z"},"links":{"cited_paper":"/paper/2305.18270","citing_paper":"/paper/2412.18756"},"observation_digest":"sha256:a3081f6e5954b3b6ad4880e3eaf3eed2dace55fb2535c4e67b03fe49d3b55909","observation_id":"fa86431a-df41-4314-9fb5-a8535fcabfaf","resolution":{"observed_at":"2026-08-11T04:36:32.044680Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18270","last_updated":"2025-06-03T19:12:56Z","snapshot_observed_at":"2026-08-10T00:04:17.648426Z","submitted_at":"2023-05-29T17:43:44Z","title":"How Two-Layer Neural Networks Learn, One (Giant) Step at a Time","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.18270","snapshot_observed_at":"2026-08-08T15:34:16.642638Z","title":"How two-layer neural networks learn, one (giant) step at a time","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06443","last_updated":"2025-07-21T10:48:50Z","snapshot_observed_at":"2026-08-10T00:04:40.232039Z","submitted_at":"2025-02-10T13:19:30Z","title":"Low-dimensional Functions are Efficiently Learnable under Randomly Biased Distributions","version":2},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-08T15:34:16.642638Z"},"links":{"cited_paper":"/paper/2305.18270","citing_paper":"/paper/2502.06443"},"observation_digest":"sha256:083de74230d6ffdefed786ae66a83259ac8e6d65db3b7e40c5ea61db56027f1c","observation_id":"a261ff52-bf6c-4edb-a572-9f73bdb300c6","resolution":{"observed_at":"2026-08-08T15:34:16.642638Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18270","last_updated":"2025-06-03T19:12:56Z","snapshot_observed_at":"2026-08-10T00:04:17.648426Z","submitted_at":"2023-05-29T17:43:44Z","title":"How Two-Layer Neural Networks Learn, One (Giant) Step at a Time","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.18270","snapshot_observed_at":"2026-08-07T14:40:54.621375Z","title":"Learning two-layer neural networks, one (giant) step at a time","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.18346","last_updated":"2025-05-23T20:09:09Z","snapshot_observed_at":"2026-08-09T18:01:10.152830Z","submitted_at":"2025-05-23T20:09:09Z","title":"On the Mechanisms of Weak-to-Strong Generalization: A Theoretical Perspective","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T14:40:54.621375Z"},"links":{"cited_paper":"/paper/2305.18270","citing_paper":"/paper/2505.18346"},"observation_digest":"sha256:e5bbe825ad8580ca11ec17e2eb482991b351ad7554cc0bbcf264fa01174daa74","observation_id":"538f54cd-b6d1-4db0-a809-544ffa8652d4","resolution":{"observed_at":"2026-08-07T14:40:54.621375Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18270","last_updated":"2025-06-03T19:12:56Z","snapshot_observed_at":"2026-08-10T00:04:17.648426Z","submitted_at":"2023-05-29T17:43:44Z","title":"How Two-Layer Neural Networks Learn, One (Giant) Step at a Time","version":4},"cited_work":{"arxiv_id":"2305.18270","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2305.18270","snapshot_observed_at":"2026-07-02T07:06:44.997380Z","title":"How two-layer neural networks learn, one (giant) step at a time.arXiv preprint arXiv:2305.18270","venue":null,"work_id":"f7b71d5d-aead-488e-b0cb-38f05965355d","year":2025},"citing_paper":{"arxiv_id":"2601.19208","last_updated":"2026-05-12T21:37:09Z","snapshot_observed_at":"2026-08-10T21:14:18.300968Z","submitted_at":"2026-01-27T05:22:34Z","title":"How Do Transformers Learn to Associate Tokens: Gradient Leading Terms Bring Mechanistic Interpretability","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-16T11:20:33.400885Z"},"links":{"cited_paper":"/paper/2305.18270","citing_paper":"/paper/2601.19208"},"observation_digest":"sha256:78a20ac3b85778b0d046eba186d775dc2e7157d794edfc0373884d4a348c8e5e","observation_id":"b92bcd50-dfe8-41fa-8cef-b6b523db09c1","resolution":{"observed_at":"2026-05-16T11:20:52.659015Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18270","last_updated":"2025-06-03T19:12:56Z","snapshot_observed_at":"2026-08-10T00:04:17.648426Z","submitted_at":"2023-05-29T17:43:44Z","title":"How Two-Layer Neural Networks Learn, One (Giant) Step at a Time","version":4},"cited_work":{"arxiv_id":"2305.18270","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2305.18270","snapshot_observed_at":"2026-07-02T07:06:44.997380Z","title":"How two-layer neural networks learn, one (giant) step at a time.arXiv preprint arXiv:2305.18270","venue":null,"work_id":"f7b71d5d-aead-488e-b0cb-38f05965355d","year":2025},"citing_paper":{"arxiv_id":"2606.04662","last_updated":"2026-06-03T09:40:30Z","snapshot_observed_at":"2026-08-12T19:27:04.962406Z","submitted_at":"2026-06-03T09:40:30Z","title":"Why Muon Outperforms Adam: A Curvature Perspective","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-06-28T07:04:21.012269Z"},"links":{"cited_paper":"/paper/2305.18270","citing_paper":"/paper/2606.04662"},"observation_digest":"sha256:82a6ae1e8a508d3bc6872135995666e8b8993bbbccf0c4dda876c9a740e3aa84","observation_id":"eec2132c-6579-4d21-98d0-7bf1a82f197e","resolution":{"observed_at":"2026-07-02T07:06:44.999214Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18270","last_updated":"2025-06-03T19:12:56Z","snapshot_observed_at":"2026-08-10T00:04:17.648426Z","submitted_at":"2023-05-29T17:43:44Z","title":"How Two-Layer Neural Networks Learn, One (Giant) Step at a Time","version":4},"cited_work":{"arxiv_id":"2305.18270","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2305.18270","snapshot_observed_at":"2026-07-02T07:06:44.997380Z","title":"How two-layer neural networks learn, one (giant) step at a time.arXiv preprint arXiv:2305.18270","venue":null,"work_id":"f7b71d5d-aead-488e-b0cb-38f05965355d","year":2025},"citing_paper":{"arxiv_id":"2606.28486","last_updated":"2026-06-26T18:00:01Z","snapshot_observed_at":"2026-08-04T02:14:38.861171Z","submitted_at":"2026-06-26T18:00:01Z","title":"Spectral phase transitions and trainability in neural network learning dynamics","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-30T01:22:17.359656Z"},"links":{"cited_paper":"/paper/2305.18270","citing_paper":"/paper/2606.28486"},"observation_digest":"sha256:64cc9b8c67aa30c79bb5e42409a480de24d65d256aa739ea268e919770783310","observation_id":"faf00f46-23d0-4e6b-a86c-1129ebdfc471","resolution":{"observed_at":"2026-07-01T15:35:47.552849Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18270","last_updated":"2025-06-03T19:12:56Z","snapshot_observed_at":"2026-08-10T00:04:17.648426Z","submitted_at":"2023-05-29T17:43:44Z","title":"How Two-Layer Neural Networks Learn, One (Giant) Step at a Time","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.18270","snapshot_observed_at":"2026-07-12T03:07:54.001891Z","title":"Learning two-layer neural networks, one (giant) step at a time","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.03347","last_updated":"2026-07-03T14:02:42Z","snapshot_observed_at":"2026-08-04T04:21:48.715795Z","submitted_at":"2026-07-03T14:02:42Z","title":"The Multiscale Single-Index Model: A Stylized Model for Hierarchical Feature Learning","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-07-12T03:07:54.001891Z"},"links":{"cited_paper":"/paper/2305.18270","citing_paper":"/paper/2607.03347"},"observation_digest":"sha256:baed941013d756b9e7328c62eb5bd19a095fb1b3f60c0514529956d2ef61410d","observation_id":"35c1a2cd-d361-4f9c-bc7f-b499cafbcb1c","resolution":{"observed_at":"2026-07-12T03:07:54.001891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2305.18270/citation-record","integrity":"/paper/2305.18270/integrity","json":"/paper/2305.18270/citation-record.json","paper":"/paper/2305.18270"},"outbound":[],"paper":{"arxiv_id":"2305.18270","last_updated":"2025-06-03T19:12:56Z","latest_version":4,"primary_category":"stat.ML","snapshot_observed_at":"2026-08-10T00:04:17.648426Z","submitted_at":"2023-05-29T17:43:44Z","title":"How Two-Layer Neural Networks Learn, One (Giant) Step at a Time"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 8 inbound Pith citation observations for arXiv:2305.18270."}