{"as_of":"2026-08-14T06:59:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b5791dd89d2249ca8e9c0c85a9ac63a159509fadfb863cc33fb6531bdd901508","coverage":[{"denominator":52,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":52,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T21:05:06.139774Z","state":"measured"},{"denominator":52,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":52,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2509.08241/citation-record","integrity":"/paper/2509.08241/integrity","json":"/paper/2509.08241/citation-record.json","paper":"/paper/2509.08241"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.171808Z","title":"Kleff, A","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.171808Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:55f76b86e97f29e2397763aac0745deb27363c91fb1ecc394e6e3f29bd1f2525","observation_id":"8a347263-995f-40f0-911d-650a812434c3","resolution":{"observed_at":"2026-08-04T21:05:04.171808Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.189937Z","title":"Grandia, F","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.189937Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:c5b236eba72d02ba5d63432f503ccd72e92ae4a04a7985f240c093ed5a4379ea","observation_id":"515fbfed-d9c0-45bd-81bb-1c437a1f33c5","resolution":{"observed_at":"2026-08-04T21:05:04.189937Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2022.32283","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:08.144057Z","title":"Meduri, P","venue":null,"work_id":"681cc30a-ce95-4cc8-a2a0-d7348ecc16af","year":2023},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.212528Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:f5fc5adf54c27880dbef57bfb561b923676eb227a50159bbe17311bbb6a0ddb1","observation_id":"de8cd6ae-a19b-403a-9d2b-bb4a70c4b296","resolution":{"observed_at":"2026-08-04T21:05:08.190979Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2024.33515","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:07.926948Z","title":"Le Cleac’h, T","venue":null,"work_id":"53b5cf9c-0554-406a-b65f-2c08ac4f8216","year":2024},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.226196Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:830488b2eb2a96834d2ee5185a453f795f37b53c314523831ae618aa995598ed","observation_id":"2755e7bb-25ca-4709-9c54-583e84663b17","resolution":{"observed_at":"2026-08-04T21:05:07.992606Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-04T21:05:04.237137Z","title":"Schulman, F","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.237137Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:b3b3c9ab78ab3ebfab6fce8e4951e6944150c8c365ec88f452520e222e8bb990","observation_id":"86d5dab0-ef1a-4a30-8671-f61e2acf4b13","resolution":{"observed_at":"2026-08-04T21:05:04.237137Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:11.008854Z","title":"Fujimoto, H","venue":null,"work_id":"23cb8364-8fb7-42a5-a54a-fa940394ee6e","year":2018},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.247783Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:f80a05738af6cfa755ce2eb1f9eb254d0fb6c09cf6269832b7cb8ece5d24bee8","observation_id":"55459a89-72c0-4e4f-81af-ef0d3e336363","resolution":{"observed_at":"2026-08-04T21:05:11.023639Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:10.885790Z","title":"Haarnoja, A","venue":null,"work_id":"c6b75f4f-02e6-4898-b36e-75b4cce65de4","year":2018},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.256052Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:665bc6f2c16602c2d1bbb9e6e9977c9e99691a434c9d21143f7ed9183f59464f","observation_id":"e95d90e8-ad8b-4266-8ef8-39defe5eb9f5","resolution":{"observed_at":"2026-08-04T21:05:10.938778Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:10.746521Z","title":"Jitosho, T","venue":null,"work_id":"a4ea56fd-1597-4707-87b5-e7a9e3386a91","year":2023},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.264381Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:bc224b52d767f3e9e77e415c602e35a00bf188b1abfee1ba16bcec9742a77650","observation_id":"6f5ee181-6d62-4090-b559-5d284e156cc6","resolution":{"observed_at":"2026-08-04T21:05:10.799647Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.273283Z","title":"Rajeswaran, V","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.273283Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:209aec47c1c937bdda62b93a49fbbb2edb8fce8be3e2c9cc1ae25633c4795e69","observation_id":"504a485c-1bb9-4146-bbfc-2f063b09fa70","resolution":{"observed_at":"2026-08-04T21:05:04.273283Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:10.664441Z","title":null,"venue":null,"work_id":"a31fb0ff-382e-41d7-9583-ca31b2114f8b","year":2022},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.284199Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:873ccb832812e2ec2dd1181c08c0765d2c81a669114b01060c14f0e5a77423c9","observation_id":"85591910-9621-449c-b7fe-a50fe0f04d9b","resolution":{"observed_at":"2026-08-04T21:05:10.718072Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.294781Z","title":"Williams, N","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.294781Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:345336e7c216754d286a49acde0a6af83c8b08f27f01a87191ca4da587c0821b","observation_id":"c139b2a4-f876-4f0b-8938-f9df54b40469","resolution":{"observed_at":"2026-08-04T21:05:04.294781Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.305318Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.305318Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:4a23bce75dc2237bb9a2469207ab3aee0c4de11262638c06beb94e41362498b1","observation_id":"9be20d2d-ec62-4c61-8098-dd4bdaf8b909","resolution":{"observed_at":"2026-08-04T21:05:04.305318Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2023.15100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:07.780936Z","title":null,"venue":null,"work_id":"66fb6e11-fd61-493f-94e2-6f35241ff156","year":2023},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.315346Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:ddfb320935852ba8db97d75fd598443f68ef7a4bfc87ac0272b4b715790435c6","observation_id":"818ccfd4-b742-4473-b724-360ce215188b","resolution":{"observed_at":"2026-08-04T21:05:07.835678Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:10.512035Z","title":null,"venue":null,"work_id":"7158fd69-0675-40bc-a7f3-0e14517b22ee","year":2013},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.332320Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:efc3d7bd3639c3e6d152551af33138435678f4d6d6f88a2a0252db0e67e37e79","observation_id":"e2ef6318-09f5-4354-8d9f-38c3080485ac","resolution":{"observed_at":"2026-08-04T21:05:10.580579Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.342337Z","title":null,"venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.342337Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:c329384fa1143b376ec10d9f5021c3aef476fae5d8c2277a0cef8fdcf4b56615","observation_id":"e02fd41f-0fc1-4ac7-ab3d-9edcb8b877b9","resolution":{"observed_at":"2026-08-04T21:05:04.342337Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:10.433817Z","title":null,"venue":null,"work_id":"a85b1659-ef9f-4d18-afa0-ce2c2242a50d","year":2017},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.349730Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:96c07abbdf2f9ad079db7fb86795db7d9147d6ef5bd0b26fdce0ffff7cede247","observation_id":"f350b2b8-26ad-41b4-9b6b-8a30d796319c","resolution":{"observed_at":"2026-08-04T21:05:10.471810Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.355894Z","title":"Geneva and N","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.355894Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:3d0b0681d2627a946d8ba05da3753e301f5b3499ee98055a87fccb20129695f9","observation_id":"10c81c8a-efb4-412d-a3b2-3c19bc36f66e","resolution":{"observed_at":"2026-08-04T21:05:04.355894Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2010.55893","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:07.657789Z","title":"Susuki and I","venue":null,"work_id":"eb741371-1ec7-48df-8a86-c3c2e49adfc3","year":2010},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.366040Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:3749ff7cde5b68e1a1630160c6af809ec92b3830b82d75f6c3c4d3f39aa724a8","observation_id":"61909c8f-4c13-41cd-97e6-477511d63ed2","resolution":{"observed_at":"2026-08-04T21:05:07.715756Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.376748Z","title":"Bruder, B","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.376748Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:a82d5d713a2d8561708548f1af25945ada97992df2c96d2c045bce010ff66ee1","observation_id":"3230d915-bd3f-4d31-84e4-d739acfaaa53","resolution":{"observed_at":"2026-08-04T21:05:04.376748Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1177/02783649241272114","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Bruder, D","venue":"The International Journal of Robotics Research","work_id":"e7dd3d11-f346-4a33-bc69-8930ad445cdd","year":2024},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.418409Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:1cd23098e4658a0301944b01a9382fc96430ca2e0c6acd4a352243d8fe701c98","observation_id":"987a8179-630d-4fb7-9bd1-fbc0d878378a","resolution":{"observed_at":"2026-08-04T21:05:06.653561Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2019.29238","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:07.537895Z","title":"Abraham and T","venue":null,"work_id":"cc077bfd-0a67-4b2b-a592-5b7e3530855e","year":2019},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.489289Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:ebe8c5bb86eb1c767c83ce73bd3dfdf06ef5e14a0e870d40548e9327a7f4d041","observation_id":"c7a36f35-bac9-470f-a78c-0b4554c3574e","resolution":{"observed_at":"2026-08-04T21:05:07.600632Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2021.30765","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:07.402922Z","title":"Mamakoukas, M","venue":null,"work_id":"11faf57c-4050-413c-8a6a-82e219bfd27b","year":2021},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.507898Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:5895e0211d72cd32bce63e12e5917e69e8486650e22d68ae9d3174404057ddb0","observation_id":"0eebd01a-ce9b-44a2-9571-2f871d049f41","resolution":{"observed_at":"2026-08-04T21:05:07.501678Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:10.329389Z","title":null,"venue":null,"work_id":"32d2aede-5500-40cf-8b12-e617d59db7b7","year":2025},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.540346Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:d592161c1b6367b3511fd1d51736b350d05277ca56336a1e275a8c75739f7696","observation_id":"91d5acf1-0272-4034-9c9a-7dbfb3164f1e","resolution":{"observed_at":"2026-08-04T21:05:10.369792Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:10.194531Z","title":null,"venue":null,"work_id":"6849b6e6-1d8f-4413-8f9a-9960cd1951ae","year":null},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.577408Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:4d6440f60b1f934ef7db6562a0bce83822f9485bd62ad888afbccfe2cef207e6","observation_id":"044a9c0a-19ca-42c2-ad0d-0576e35a2c54","resolution":{"observed_at":"2026-08-04T21:05:10.282622Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.727583Z","title":null,"venue":null,"work_id":null,"year":1996},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.727583Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:97b8d62580016e82864571e3e3302b81b1f57f638221755fc3addf46ff089098","observation_id":"dcd148f4-df8c-4035-a27c-8d30a2210724","resolution":{"observed_at":"2026-08-04T21:05:04.727583Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:09.951508Z","title":"Nagabandi, K","venue":null,"work_id":"1c6293bb-800f-4a18-9a6a-a2ce03b21aa5","year":null},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.798756Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:3b881cb0b2c4bba436737fa51b30418f79c81aba9319183a3f6c0378fab8b158","observation_id":"68e15270-7760-4b29-afa8-f54f92b4e8c6","resolution":{"observed_at":"2026-08-04T21:05:10.008961Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1137/18m1192329","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Zhang, C","venue":"SIAM Journal on Applied Dynamical Systems","work_id":"4d878824-25be-4171-9363-78b05fa3888e","year":2019},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.873250Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:d096cae0819bbccc73804db99d867b1aa379eaefa5af8036408c72a7bdeed070","observation_id":"494f5936-34ab-44cb-974a-e29b3035ecf3","resolution":{"observed_at":"2026-08-04T21:05:06.529917Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.910368Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.910368Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:7268d9f10b705b393090c693c36bf78b451f9803636805c99d0815be0e0b9230","observation_id":"5b401602-3256-459d-9761-b03b8f8876cd","resolution":{"observed_at":"2026-08-04T21:05:04.910368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:04.996420Z","title":"N ¨uske, S","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.996420Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:53508e410f8620a6f6891a541fd27d5b26b69cf223d71f0737d7031711410cbb","observation_id":"b10bf17b-7379-452e-bd9e-b3c9ff4ce308","resolution":{"observed_at":"2026-08-04T21:05:04.996420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:09.758286Z","title":"Zhang and E","venue":null,"work_id":"7327324f-589d-4390-94b1-5549d8b465a3","year":2023},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.052964Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:b9cb37db48e215bde3a35dd64d6f34a36183bb99da5df7b7415e3f40200f8a8b","observation_id":"1fe6c94c-4ce9-4956-8bc8-0be37e09d567","resolution":{"observed_at":"2026-08-04T21:05:09.783965Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.02494","last_updated":"2024-05-23T09:04:32Z","snapshot_observed_at":"2026-08-14T06:35:16.004725Z","submitted_at":"2024-02-04T13:58:48Z","title":"Variance representations and convergence rates for data-driven approximations of Koopman operators","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.02494","snapshot_observed_at":"2026-08-04T21:05:05.127093Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.127093Z"},"links":{"cited_paper":"/paper/2402.02494","citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:3848685e697cf8cf8ccd3508f3656a86d806b1ce38cd83c073533d2817bbf1b3","observation_id":"a9d4ef65-5aed-4eaf-a5bb-3c13ba5f6860","resolution":{"observed_at":"2026-08-04T21:05:05.127093Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:09.660695Z","title":null,"venue":null,"work_id":"ca91b40e-5301-4a44-b61e-fcb0faf3a05e","year":1960},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.179258Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:1c37a5f3da5f3af1b6a2550fd3adbfeeef53626e785164e2e74df8a34de70ea2","observation_id":"34920a5c-2947-4122-a5f4-c45d050657a0","resolution":{"observed_at":"2026-08-04T21:05:09.698041Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:05.234896Z","title":null,"venue":null,"work_id":null,"year":1931},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.234896Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:1bb2c5136437bb226dddd48129403f0d477dfb0695c9230b7c9bf9d4f2664910","observation_id":"b39feebb-cc6e-4030-98ca-23dd956771d5","resolution":{"observed_at":"2026-08-04T21:05:05.234896Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2023.32535","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:07.272207Z","title":null,"venue":null,"work_id":"d86ad04f-3221-4f9b-953f-f6205af7ddd5","year":2023},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.314436Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:7a8687d9e3460be224274f8a301e2a2f9cda2bb53ffbc3a394bd245c0a4ef981","observation_id":"70df344f-486e-4f66-8f3d-7c070c7095e5","resolution":{"observed_at":"2026-08-04T21:05:07.335566Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:05.354839Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.354839Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:d353d6abc58bf2face8dfed64fd3b95745bef1275c17fb00119ba425301beda4","observation_id":"8c90d553-d0c3-4c2f-9678-50f69d81ec38","resolution":{"observed_at":"2026-08-04T21:05:05.354839Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"stable/2236561","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:07.119773Z","title":"Sherman and W","venue":null,"work_id":"6dfcdec5-705a-43ac-a276-79318244df1e","year":1950},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.392479Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:befd8fc6be0de6d81faacd40cde69915693a790a8cf2fbf75acfe27615a20f02","observation_id":"749bd4fe-34d2-4c31-be52-0f442118a21b","resolution":{"observed_at":"2026-08-04T21:05:07.186634Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2016.25967","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:06.987301Z","title":null,"venue":null,"work_id":"4bea7636-bf27-468a-8b86-7942156305bf","year":2016},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.459326Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:67c7f3b98713a3e98f4d62cf1b04f0535d55caba04bbb61c2b61cab825f43fc5","observation_id":"58d7a302-1905-4911-b63c-2257fbf863e5","resolution":{"observed_at":"2026-08-04T21:05:07.052656Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1177/02783649211037697","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Nishimura and M","venue":"The International Journal of Robotics Research","work_id":"150624e3-e919-4dfe-b9ff-fc310304e494","year":2021},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.518653Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:6d8aaf6018ffabffe73d2b1303d1a5694bffe097e161b616655899d25bb2f3ce","observation_id":"deea914f-54cb-407c-9cda-2128190ffe28","resolution":{"observed_at":"2026-08-04T21:05:06.415320Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:09.545682Z","title":"Ketchum, J","venue":null,"work_id":"dfeec9b7-68a4-43e7-86d7-af63cd1e878b","year":2025},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.595151Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:70d64281a626621b2d02420ac042e5d5fbfe88d4ce2d88632692b29fd4a7333a","observation_id":"21f1f06e-c1f0-4f62-ab22-845352fd0f57","resolution":{"observed_at":"2026-08-04T21:05:09.614995Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:05.607354Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.607354Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:5fd99e17a7ba5a2b1cf49ee2275f6f509cfd24691d5eaae50f06df70cf7c2b44","observation_id":"f0644007-7678-4864-8845-b5a5e58711e2","resolution":{"observed_at":"2026-08-04T21:05:05.607354Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.15607/rss.2017.xiii.052","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Abraham, G","venue":null,"work_id":"51960487-2f25-4fb5-91ee-d2a0b37f5d4d","year":2017},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.660725Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:b0b5604b2241021113f62254bf326f53b3decdf59067bb430f576e0ec2b7d6ee","observation_id":"e405f2ff-c653-49c6-a77b-67c9978c7444","resolution":{"observed_at":"2026-08-04T21:05:06.262459Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:09.445018Z","title":null,"venue":null,"work_id":"099307ff-46b2-4872-9c5a-c9112b34b703","year":null},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.699189Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:ff3cac653031452f3db6cecf2bba544d81b9f324b3e424b7798e234ed3fadeae","observation_id":"c1416384-f9a9-452d-9964-1da9a1c2389f","resolution":{"observed_at":"2026-08-04T21:05:09.486751Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:09.213330Z","title":"Avtges, J","venue":null,"work_id":"7cb8b335-c16c-4ccc-aace-d58d4007baae","year":2025},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.838867Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:be101006a444ec91d8b7b4ac82766cddac4ef48fa0cd5685424695fd135e268d","observation_id":"a9533242-ddc1-4a2a-bbb0-5d717f534351","resolution":{"observed_at":"2026-08-04T21:05:09.252378Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:09.070730Z","title":"Boyd and L","venue":null,"work_id":"f2883199-32e5-4677-97ea-2afc0f0e58d5","year":2018},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.869737Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:8d60ad61dc89cc18cba90245b87dbeba1d45135314b20429f95a8ff9c3b26631","observation_id":"12202301-4789-4e78-b7a3-ed4f3a753e7b","resolution":{"observed_at":"2026-08-04T21:05:09.129772Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:08.924507Z","title":null,"venue":null,"work_id":"fca85b83-da8a-44c1-bc4c-9605564c9430","year":1969},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.895336Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:839f4e8842e653db40089bd7483585fcd7e54191258485651dbe527bd6be85f1","observation_id":"e952bee5-1261-4ba6-90e6-02fdfc3f514f","resolution":{"observed_at":"2026-08-04T21:05:08.980629Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:08.770106Z","title":"Foundation","venue":null,"work_id":"30875bab-cf83-4d0d-9f07-8f6ded6e8e0c","year":null},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.961038Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:cc9586d9794c9075e6fe9a29a7cd6a2acfbc59fa244aadeb4e0110f8bfb5bbf2","observation_id":"d4614651-6b73-471c-83c4-65f76b1cf320","resolution":{"observed_at":"2026-08-04T21:05:08.826416Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:08.518577Z","title":null,"venue":null,"work_id":"9cc11fa8-2200-4d4f-9451-db272cacd652","year":2024},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:06.021893Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:33b522179bc200ec1d660d7e17d1c713c66ff678029d47c9d070ef142b77c88e","observation_id":"b4a13b58-fd35-42d4-9519-c946dd7c050b","resolution":{"observed_at":"2026-08-04T21:05:08.667585Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:08.313706Z","title":"Raffin, A","venue":null,"work_id":"960fda06-a46d-403c-ae8b-31f28820e77c","year":2021},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:06.076252Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:5fc99aeeaa02b721ca5f05ad8dafff1c3b33f702eb0ab29a3435699f8880ab7d","observation_id":"e035135c-4bb4-4578-ab81-ca376ffa129f","resolution":{"observed_at":"2026-08-04T21:05:08.412644Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2016.74872","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:06.776486Z","title":"Williams, P","venue":null,"work_id":"2bed1223-24ab-4e1b-be91-3ee73b01e48d","year":2016},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:06.139774Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:9ec24151175f1cab60022621571542d18f29b11c312021df5e8e043d5b9e9416","observation_id":"a877439f-12f0-4487-82b0-a6c3bf563907","resolution":{"observed_at":"2026-08-04T21:05:06.873831Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:09.834410Z","title":"URLhttps://proceedings.mlr.press/v100/ nagabandi20a.html","venue":null,"work_id":"a6c619a5-9820-43b6-b608-fae98fb05f29","year":2020},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":1112,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.835881Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:337e2642c15a7ed85cbe29f6f3088fe4e5e32264a4825cfd94c9b02af63e4c5c","observation_id":"0d277d62-12bc-47e8-b3c2-4fb071431a10","resolution":{"observed_at":"2026-08-04T21:05:09.872022Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:10.082660Z","title":null,"venue":null,"work_id":"85bb9181-c938-4382-9fc4-25d6b1b2426d","year":1988},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":1988,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:04.647365Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:ac559dcd04483de6423d0fb2e05c3d6a82fa6d4c01ab02741ee7b1e8772abb71","observation_id":"6e4e0e16-2209-49c6-9784-d4de61ff3a45","resolution":{"observed_at":"2026-08-04T21:05:10.113984Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:05:09.341475Z","title":null,"venue":null,"work_id":"8fd1b6cf-4b8e-4205-b266-c7585e67cb3d","year":null},"citing_paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-04T21:05:05.770512Z"},"links":{"citing_paper":"/paper/2509.08241"},"observation_digest":"sha256:340af5400efcb712728a328d6f050b51adeb68fa3266adc8986abc85d5815561","observation_id":"fd2484ff-d542-4301-bbf9-83f18ba2ab3a","resolution":{"observed_at":"2026-08-04T21:05:09.382596Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2509.08241","last_updated":"2025-09-10T02:47:42Z","latest_version":1,"primary_category":"cs.RO","snapshot_observed_at":"2026-08-13T23:23:01.000675Z","submitted_at":"2025-09-10T02:47:42Z","title":"Sample-Efficient Online Control Policy Learning with Real-Time Recursive Model Updates"},"reference_resolution":{"displayed":52,"state_counts":{"malformed_identifier":5,"metadata_mismatch":7,"parse_uncertain":0,"unresolved":24,"verified_exact":5,"verified_fuzzy":11},"total_outbound_references":52},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 52 of 52 outbound references and 0 inbound Pith citation observations for arXiv:2509.08241."}