{"as_of":"2026-08-10T09:32:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c16c6de60196671e4e8a5f79a720f74ca850a55f5852034c336b1e0bae12f918","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":15,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":15,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":15,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":15,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T14:30:08.919449Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T09:59:44.626163Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-08-09T14:30:08.919449Z","title":"The ingredients of real-world robotic reinforcement learning","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2502.01800","last_updated":"2025-05-05T19:40:42Z","snapshot_observed_at":"2026-08-10T07:01:34.329509Z","submitted_at":"2025-02-03T20:25:50Z","title":"Flow-based Domain Randomization for Learning and Sequencing Robotic Skills","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-09T14:30:08.919449Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2502.01800"},"observation_digest":"sha256:d0a7ca6da13f37d5bd47a52ffd66b859a3499adaf5f800bb1def6a25fd020dec","observation_id":"1fc3d64b-cddb-42dd-8f93-1b7d4114c1fb","resolution":{"observed_at":"2026-08-09T14:30:08.919449Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":"2004.12570","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-07-04T09:59:44.626163Z","title":"arXiv preprint arXiv:2004.12570 (2020) 2","venue":null,"work_id":"58233088-569f-42a2-bad2-8195f07616cc","year":2004},"citing_paper":{"arxiv_id":"2505.18719","last_updated":"2025-05-24T14:42:51Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-24T14:42:51Z","title":"VLA-RL: Towards Masterful and General Robotic Manipulation with Scalable Reinforcement Learning","version":1},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-05-16T12:55:40.245908Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2505.18719"},"observation_digest":"sha256:caf248deca62ed8ffb32c860f4ec63b988964ee945aa580f1d697ce3b8940feb","observation_id":"7b83022a-a003-4080-a99f-2f33b4a44360","resolution":{"observed_at":"2026-05-16T12:55:40.405729Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-08-07T05:49:54.928309Z","title":"The ingredients of real-world robotic re- inforcement learning.arXiv preprint arXiv:2004.12570, 2020","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2506.07006","last_updated":"2025-06-08T06:05:32Z","snapshot_observed_at":"2026-08-09T11:00:17.510668Z","submitted_at":"2025-06-08T06:05:32Z","title":"CARoL: Context-aware Adaptation for Robot Learning","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T05:49:54.928309Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2506.07006"},"observation_digest":"sha256:e79deecd2a623dfaf4fb948488c1436e711b2dca43084f45c98c9017f3043fa7","observation_id":"e679ad12-e1af-41e9-a591-c6726325a71e","resolution":{"observed_at":"2026-08-07T05:49:54.928309Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-08-06T23:55:12.511151Z","title":"The ingredients of real-world robotic re- inforcement learning.arXiv preprint arXiv:2004.12570, 2020","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2506.15847","last_updated":"2025-06-18T19:55:10Z","snapshot_observed_at":"2026-08-08T04:12:43.706263Z","submitted_at":"2025-06-18T19:55:10Z","title":"SafeMimic: Towards Safe and Autonomous Human-to-Robot Imitation for Mobile Manipulation","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T23:55:12.511151Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2506.15847"},"observation_digest":"sha256:459a88930a201bf5c3072921a0187f088f6dac92255bfff2a945a84f73b65cfa","observation_id":"ce0ac35c-94a5-48b9-b3fa-d6b9e18f8096","resolution":{"observed_at":"2026-08-06T23:55:12.511151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-08-06T21:17:36.076495Z","title":"The ingredients of real-world robotic reinforcement learning,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.00611","last_updated":"2025-07-01T09:43:57Z","snapshot_observed_at":"2026-08-06T21:08:39.566215Z","submitted_at":"2025-07-01T09:43:57Z","title":"Residual Reward Models for Preference-based Reinforcement Learning","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T21:17:36.076495Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2507.00611"},"observation_digest":"sha256:743d642932c59dc9d6022f4d9870feef7f25501eb0a32296c6967cc216535d4b","observation_id":"d0d5d731-8015-4ce6-a8e5-88b4ac8a5704","resolution":{"observed_at":"2026-08-06T21:17:36.076495Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-08-06T19:53:06.317839Z","title":"The ingredients of real-world robotic reinforcement learning,","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2507.04452","last_updated":"2025-07-06T16:18:40Z","snapshot_observed_at":"2026-08-08T10:00:23.689092Z","submitted_at":"2025-07-06T16:18:40Z","title":"SimLauncher: Launching Sample-Efficient Real-world Robotic Reinforcement Learning via Simulation Pre-training","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T19:53:06.317839Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2507.04452"},"observation_digest":"sha256:3539659f753ae044a98f139a2d4056250a8c5daa8f0cecf51ea6d9e3294cedec","observation_id":"33ee60cd-a81e-48d3-8c09-ebdd61e8a0d1","resolution":{"observed_at":"2026-08-06T19:53:06.317839Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-08-06T13:21:59.449740Z","title":"The ingredients of real-world robotic reinforcement learning","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2507.20853","last_updated":"2025-07-28T14:06:44Z","snapshot_observed_at":"2026-08-10T04:44:01.568944Z","submitted_at":"2025-07-28T14:06:44Z","title":"Geometry of Neural Reinforcement Learning in Continuous State and Action Spaces","version":1},"reference_index":139,"source":"arxiv_source","source_observed_at":"2026-08-06T13:21:59.449740Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2507.20853"},"observation_digest":"sha256:a72a3379fb8a78d8c45cecf3923b626efa69c7d3b425c55dbe686d35bbfcd354","observation_id":"e31bd30f-f515-41b1-90a0-0c01d78464e3","resolution":{"observed_at":"2026-08-06T13:21:59.449740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":"2004.12570","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-07-04T09:59:44.626163Z","title":"arXiv preprint arXiv:2004.12570 (2020) 2","venue":null,"work_id":"58233088-569f-42a2-bad2-8195f07616cc","year":2004},"citing_paper":{"arxiv_id":"2511.14759","last_updated":"2025-11-19T04:34:49Z","snapshot_observed_at":"2026-07-06T22:36:13.287872Z","submitted_at":"2025-11-18T18:58:55Z","title":"$\\pi^{*}_{0.6}$: a VLA That Learns From Experience","version":2},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-05-12T10:34:59.134604Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2511.14759"},"observation_digest":"sha256:d1abb637b6f0ef4f90b34ce5c1db18ef50114b72f10f3bb7f98ee83bc2a9a144","observation_id":"adb13b4b-71bd-4495-b751-9bdfe8dba9c1","resolution":{"observed_at":"2026-05-12T10:34:59.439726Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-08-03T18:19:01.453963Z","title":"The ingredients of real-world robotic reinforcement learning","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2512.06244","last_updated":"2026-07-31T22:43:03Z","snapshot_observed_at":"2026-08-09T03:45:26.166178Z","submitted_at":"2025-12-06T02:04:50Z","title":"Auto-exploration for online reinforcement learning","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-03T18:19:01.453963Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2512.06244"},"observation_digest":"sha256:11a157e1ba63b0d4bca11dd71fc541dab8da111f26a8549c9a99ba5922684c83","observation_id":"71344456-bc8b-40c8-9b23-eaf0b07ed8ab","resolution":{"observed_at":"2026-08-03T18:19:01.453963Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-08-04T06:51:04.150269Z","title":"The ingredients of real-world robotic reinforcement learning","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2512.06244","last_updated":"2026-07-31T22:43:03Z","snapshot_observed_at":"2026-08-09T03:45:26.166178Z","submitted_at":"2025-12-06T02:04:50Z","title":"Auto-exploration for online reinforcement learning","version":3},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-04T06:51:04.150269Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2512.06244"},"observation_digest":"sha256:5429351bdafb485b3f8b0a0fe31850f7bd550b8134e7400416b15c2d7bed6ec8","observation_id":"4796b188-2383-420a-abcf-946678d6cf8f","resolution":{"observed_at":"2026-08-04T06:51:04.150269Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-07-14T23:46:32.301737Z","title":"The ingredients of real- world robotic reinforcement learning.arXiv preprint arXiv:2004.12570,","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2603.10263","last_updated":"2026-07-01T06:52:08Z","snapshot_observed_at":"2026-08-07T08:16:48.418429Z","submitted_at":"2026-03-10T22:49:46Z","title":"From Prior to Pro: Efficient Skill Mastery via Distribution Contractive RL Finetuning","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-07-14T23:46:32.301737Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2603.10263"},"observation_digest":"sha256:6e11af56d3e13f87aa99b670820003491f0b99f43ebff07a2fc259dc0e7230d2","observation_id":"3ef674cb-e39e-49f4-ade1-ec2e4348e9e2","resolution":{"observed_at":"2026-07-14T23:46:32.301737Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":"2004.12570","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-07-04T09:59:44.626163Z","title":"arXiv preprint arXiv:2004.12570 (2020) 2","venue":null,"work_id":"58233088-569f-42a2-bad2-8195f07616cc","year":2004},"citing_paper":{"arxiv_id":"2604.23073","last_updated":"2026-04-30T20:04:38Z","snapshot_observed_at":"2026-07-31T07:02:53.376869Z","submitted_at":"2026-04-24T23:57:45Z","title":"RL Token: Bootstrapping Online RL with Vision-Language-Action Models","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-08T11:56:34.978806Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2604.23073"},"observation_digest":"sha256:73b828016c18ec1e7f154def0e03d67fdb7d6b704b81985871c65f0da44a8912","observation_id":"5c86aa38-7478-473e-a601-fdf79fc51a6d","resolution":{"observed_at":"2026-05-11T19:26:09.269122Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":"2004.12570","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-07-04T09:59:44.626163Z","title":"arXiv preprint arXiv:2004.12570 (2020) 2","venue":null,"work_id":"58233088-569f-42a2-bad2-8195f07616cc","year":2004},"citing_paper":{"arxiv_id":"2606.10927","last_updated":"2026-06-09T14:35:53Z","snapshot_observed_at":"2026-08-01T15:29:53.973770Z","submitted_at":"2026-06-09T14:35:53Z","title":"AllDayNav: Lifelong Navigation via Real-World Reinforcement Learning","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-06-27T13:25:59.194721Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2606.10927"},"observation_digest":"sha256:ca3f3f93f234e43aaf08f351152ae7721e1dc5066dfb5f52eadf29ff5cd1dec2","observation_id":"534b0b5b-e3b5-4285-8ce5-91dbb54b8b8b","resolution":{"observed_at":"2026-07-03T05:07:38.984009Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":"2004.12570","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-07-04T09:59:44.626163Z","title":"arXiv preprint arXiv:2004.12570 (2020) 2","venue":null,"work_id":"58233088-569f-42a2-bad2-8195f07616cc","year":2004},"citing_paper":{"arxiv_id":"2606.12372","last_updated":"2026-06-10T17:38:24Z","snapshot_observed_at":"2026-08-06T11:07:11.202179Z","submitted_at":"2026-06-10T17:38:24Z","title":"UniIntervene: Agentic Intervention for Efficient Real-World Reinforcement Learning","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-06-27T09:46:59.746745Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2606.12372"},"observation_digest":"sha256:fc8c516ccac4bed88316cd6b683f9e82deeac78c705e57ba6804c1f4025eb261","observation_id":"0cbc1fb2-9f07-4fa9-963b-18f5808bb34a","resolution":{"observed_at":"2026-07-03T10:58:02.679820Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning","version":1},"cited_work":{"arxiv_id":"2004.12570","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2004.12570","snapshot_observed_at":"2026-07-04T09:59:44.626163Z","title":"arXiv preprint arXiv:2004.12570 (2020) 2","venue":null,"work_id":"58233088-569f-42a2-bad2-8195f07616cc","year":2004},"citing_paper":{"arxiv_id":"2606.23640","last_updated":"2026-06-22T17:30:24Z","snapshot_observed_at":"2026-08-02T09:22:48.099895Z","submitted_at":"2026-06-22T17:30:24Z","title":"Learning Process Rewards via Success Visitation Matching for Efficient RL","version":1},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-06-26T09:20:35.062060Z"},"links":{"cited_paper":"/paper/2004.12570","citing_paper":"/paper/2606.23640"},"observation_digest":"sha256:d65a066dda94b496e406a112cfe104c5667d1ca0706ea85e76c9017e382274ac","observation_id":"a877d542-c656-4cad-b672-20f8ee46bd91","resolution":{"observed_at":"2026-07-04T09:59:44.627806Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2004.12570/citation-record","integrity":"/paper/2004.12570/integrity","json":"/paper/2004.12570/citation-record.json","paper":"/paper/2004.12570"},"outbound":[],"paper":{"arxiv_id":"2004.12570","last_updated":"2020-04-27T03:36:10Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-09T09:30:54.167911Z","submitted_at":"2020-04-27T03:36:10Z","title":"The Ingredients of Real-World Robotic Reinforcement Learning"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 15 inbound Pith citation observations for arXiv:2004.12570."}