{"as_of":"2026-08-09T13:12:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:cc89a70cf91567c165a78ec013999715e95688eabcc1b27ff6ca89ecd0d10052","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":23,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":23,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":23,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":23,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T10:46:06.424312Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":50,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2205.06175","last_updated":"2022-11-11T10:04:29Z","snapshot_observed_at":"2026-08-08T03:18:33.595658Z","submitted_at":"2022-05-12T16:03:26Z","title":"A Generalist Agent","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-13T06:24:49.833638Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2205.06175"},"observation_digest":"sha256:c83f85914d66b26641e7bf87334859dcbc446ed76c873a83165cd58b504a2b7a","observation_id":"986c05f3-628c-4994-b984-5400d767e7ae","resolution":{"observed_at":"2026-05-13T06:24:49.880515Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2301.04104","last_updated":"2024-04-17T17:41:20Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-01-10T18:12:16Z","title":"Mastering Diverse Domains through World Models","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-11T09:08:21.677362Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2301.04104"},"observation_digest":"sha256:5235bc4518812c8e405a9887828eb88a86a7e2a06124a0c259c0c6b046d02938","observation_id":"277ff47c-34e9-4d15-8d37-cc93fd2ce4df","resolution":{"observed_at":"2026-05-11T09:08:22.082585Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2302.01560","last_updated":"2024-07-08T05:56:47Z","snapshot_observed_at":"2026-08-02T14:50:04.461434Z","submitted_at":"2023-02-03T06:06:27Z","title":"Describe, Explain, Plan and Select: Interactive Planning with Large Language Models Enables Open-World Multi-Task Agents","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-16T03:27:40.524895Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2302.01560"},"observation_digest":"sha256:d9e1d0f9c216cbecfdaa825610586e4baf9c7e738da363a374f604efb754e701","observation_id":"ef1029a3-d927-436a-bacf-a94fcdbd3cdf","resolution":{"observed_at":"2026-05-16T03:27:40.743136Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2305.16291","last_updated":"2023-10-19T16:27:03Z","snapshot_observed_at":"2026-08-07T08:29:46.650400Z","submitted_at":"2023-05-25T17:46:38Z","title":"Voyager: An Open-Ended Embodied Agent with Large Language Models","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T13:11:40.995345Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2305.16291"},"observation_digest":"sha256:8a338f26299118935a408745275ed8ebe3f04126c799fbda460822315b80fa1a","observation_id":"f8c883c3-b099-4fcc-95d4-879a812a740b","resolution":{"observed_at":"2026-05-10T13:11:41.059888Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2309.16797","last_updated":"2023-09-28T19:01:07Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-28T19:01:07Z","title":"Promptbreeder: Self-Referential Self-Improvement Via Prompt Evolution","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-05-16T08:12:30.984870Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2309.16797"},"observation_digest":"sha256:e9ae93ee829ac521b87b69da143ebcdb9fcbf6f8fcf54a8cbb9b9c09ffb9a352","observation_id":"fce01e15-727f-481a-a06c-901bd4b06853","resolution":{"observed_at":"2026-05-16T08:12:31.274505Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2310.06114","last_updated":"2024-09-26T17:14:09Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-09T19:42:22Z","title":"Learning Interactive Real-World Simulators","version":3},"reference_index":219,"source":"arxiv_source","source_observed_at":"2026-05-16T02:15:18.265190Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2310.06114"},"observation_digest":"sha256:6ce623b05c902ebec32e9ff61df62cc8e44e16ccd8cafc8ffe4b3b616bc3fbac","observation_id":"617edcce-b5d0-4431-ab19-6ed9ae3abe8b","resolution":{"observed_at":"2026-05-16T02:15:18.445102Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2412.02125","last_updated":"2026-05-01T15:42:42Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-03T03:27:48Z","title":"Preference Goal Tuning: Post-Training as Latent Control for Frozen Policies","version":2},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-05-23T08:20:05.898025Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2412.02125"},"observation_digest":"sha256:d8bf61f9c5d27c83172e08273fd2a17bef8326b2fcea2f3dc7fb586d5c4836c0","observation_id":"676efef6-edf1-49f2-ac33-e253b0d5d32a","resolution":{"observed_at":"2026-05-23T08:22:44.385372Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2505.12705","last_updated":"2025-06-17T22:33:35Z","snapshot_observed_at":"2026-07-06T21:26:01.534136Z","submitted_at":"2025-05-19T04:55:39Z","title":"DreamGen: Unlocking Generalization in Robot Learning through Video World Models","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-15T23:50:45.332466Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2505.12705"},"observation_digest":"sha256:bf67cfdeb472cecade85b512949ed49ea9f1ab5f35fbd48f7cf312264708438c","observation_id":"45edeb9c-7083-458f-8200-889add293bd6","resolution":{"observed_at":"2026-05-15T23:50:45.650952Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-06T10:46:06.424312Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.23698","last_updated":"2025-07-31T16:20:02Z","snapshot_observed_at":"2026-08-06T12:52:22.195774Z","submitted_at":"2025-07-31T16:20:02Z","title":"Scalable Multi-Task Reinforcement Learning for Generalizable Spatial Intelligence in Visuomotor Agents","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-06T10:46:06.424312Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2507.23698"},"observation_digest":"sha256:10531003cd44cf9893bb9c611a866c8fad8eac35bd1637f697042f9c2e5a3d4e","observation_id":"dcf482dd-1dd7-426a-ae45-b605024bf623","resolution":{"observed_at":"2026-08-06T10:46:06.424312Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T13:46:45.747217Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.00361","last_updated":"2025-08-30T04:53:32Z","snapshot_observed_at":"2026-08-09T09:27:18.513708Z","submitted_at":"2025-08-30T04:53:32Z","title":"Generative Visual Foresight Meets Task-Agnostic Pose Estimation in Robotic Table-Top Manipulation","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-05T13:46:45.747217Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2509.00361"},"observation_digest":"sha256:d87ee85bf0f5d901aeac24416fe49972d725dd37e08cc1744b691e94191b99e1","observation_id":"6737d5a0-3899-4a16-84f7-71eba7eb3c2a","resolution":{"observed_at":"2026-08-05T13:46:45.747217Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-04T23:58:44.340988Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.06235","last_updated":"2025-09-07T22:51:12Z","snapshot_observed_at":"2026-08-09T10:09:57.218826Z","submitted_at":"2025-09-07T22:51:12Z","title":"PillagerBench: Benchmarking LLM-Based Agents in Competitive Minecraft Team Environments","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-04T23:58:44.340988Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2509.06235"},"observation_digest":"sha256:07ac29e41254aed2f46412938f2e85f577b014a969cb87af52e2671085e3c500","observation_id":"49b175a3-22db-4fad-a340-e10665f6aa08","resolution":{"observed_at":"2026-08-04T23:58:44.340988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2605.13918","last_updated":"2026-05-13T12:52:35Z","snapshot_observed_at":"2026-08-02T07:48:11.711356Z","submitted_at":"2026-05-13T12:52:35Z","title":"CA2: Code-Aware Agent for Automated Game Testing","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-15T05:53:26.558042Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2605.13918"},"observation_digest":"sha256:be07a943bf802e2d6acea584a6d12aff7286361ad711057b6d0de26b54168f54","observation_id":"a9059ad9-2301-4408-90f9-24b24924b0ef","resolution":{"observed_at":"2026-05-15T05:55:04.989946Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2605.14211","last_updated":"2026-06-05T16:41:58Z","snapshot_observed_at":"2026-08-06T05:54:33.547619Z","submitted_at":"2026-05-14T00:10:12Z","title":"ASH: Agents that Self-Hone via Embodied Learning","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-15T02:51:53.328977Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2605.14211"},"observation_digest":"sha256:5e110822a238c1ebbae9218e90e29c736fda64cfe3e3f0f68123b3c737ee70e6","observation_id":"9ce58f77-7b49-406b-897e-8d4da5292e23","resolution":{"observed_at":"2026-05-15T02:53:33.482966Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2605.14211","last_updated":"2026-06-05T16:41:58Z","snapshot_observed_at":"2026-08-06T05:54:33.547619Z","submitted_at":"2026-05-14T00:10:12Z","title":"ASH: Agents that Self-Hone via Embodied Learning","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-30T21:12:54.453886Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2605.14211"},"observation_digest":"sha256:5b183a3be9d2f756486df5155d7ed3cfc931bf9599b8ac4e65828e16c2dd9eae","observation_id":"cdb9c999-d720-4bd8-8ec8-da56f66bc0d6","resolution":{"observed_at":"2026-06-30T21:15:04.116021Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2605.19503","last_updated":"2026-05-20T07:37:16Z","snapshot_observed_at":"2026-08-01T16:57:53.244755Z","submitted_at":"2026-05-19T07:54:40Z","title":"ARC-RL: A Reinforcement Learning Playground Inspired by ARC Raiders","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-20T05:28:50.354662Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2605.19503"},"observation_digest":"sha256:c05025d8adab1d87980392fc790f4877b6421f3959e022ce89d917f414e9addc","observation_id":"bc03e82d-e8dc-42d2-879a-8df64c52aee6","resolution":{"observed_at":"2026-05-20T05:33:04.182122Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2605.19503","last_updated":"2026-05-20T07:37:16Z","snapshot_observed_at":"2026-08-01T16:57:53.244755Z","submitted_at":"2026-05-19T07:54:40Z","title":"ARC-RL: A Reinforcement Learning Playground Inspired by ARC Raiders","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-21T07:36:12.214949Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2605.19503"},"observation_digest":"sha256:fa1a02c42c5e404faeb9c0760d2bf53bf0512311a4322c592bbc8c8350296dee","observation_id":"302b840e-9cec-4864-b8a9-e6955826f61d","resolution":{"observed_at":"2026-05-21T07:39:49.272258Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2606.26025","last_updated":"2026-07-03T15:34:07Z","snapshot_observed_at":"2026-07-12T12:05:57.041170Z","submitted_at":"2026-06-24T16:53:36Z","title":"In-Context World Modeling for Robotic Control","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-06-25T19:12:22.513577Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2606.26025"},"observation_digest":"sha256:8d51afa8b50f89a25cf08073d77b5e0c90e4be0bba35d1b70d1d699c77bda533","observation_id":"d357f8c1-c2f0-4f40-8fb9-2991c3e585e8","resolution":{"observed_at":"2026-07-04T21:00:09.589878Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2606.26025","last_updated":"2026-07-03T15:34:07Z","snapshot_observed_at":"2026-07-12T12:05:57.041170Z","submitted_at":"2026-06-24T16:53:36Z","title":"In-Context World Modeling for Robotic Control","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-06-26T05:11:07.089829Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2606.26025"},"observation_digest":"sha256:94a6ae9a449f1b6cb95575b69b285766206bab780ffc09ac3c10c593bbf6147b","observation_id":"dda212b4-94fa-47a6-8991-c5c1ed9fd39d","resolution":{"observed_at":"2026-07-04T13:29:51.595338Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-07-12T12:05:57.682386Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2606.26025","last_updated":"2026-07-03T15:34:07Z","snapshot_observed_at":"2026-07-12T12:05:57.041170Z","submitted_at":"2026-06-24T16:53:36Z","title":"In-Context World Modeling for Robotic Control","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-07-12T12:05:57.682386Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2606.26025"},"observation_digest":"sha256:3d0cefb2e533bc293aea9c5bd687b6b566a5384660744d4912838ba0b275e4c8","observation_id":"de3919ae-247d-4613-8294-24c54513fe49","resolution":{"observed_at":"2026-07-12T12:05:57.682386Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2606.26694","last_updated":"2026-06-28T04:16:09Z","snapshot_observed_at":"2026-07-07T00:01:01.437499Z","submitted_at":"2026-06-25T07:27:09Z","title":"PhysEditWorld: A Large-Scale Dataset Toward Physics-Editable World Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-26T05:46:21.198781Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2606.26694"},"observation_digest":"sha256:7e60fa197ffcf512a03ee2c76430f8557c508afd33be849b07f3d8390fa57aeb","observation_id":"1b6a113c-79dd-44de-884c-bd4feec47771","resolution":{"observed_at":"2026-07-04T12:49:53.233336Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":"2206.11795","doi":"10.48550/arxiv.2206.11795","metadata_source":"arxiv_reference","pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Video pretraining (vpt): Learning to act by watching unlabeled online videos","venue":"arXiv (Cornell University)","work_id":"2bf44de9-4b8f-4ab3-816e-2ffb5278c578","year":2022},"citing_paper":{"arxiv_id":"2606.26694","last_updated":"2026-06-28T04:16:09Z","snapshot_observed_at":"2026-07-07T00:01:01.437499Z","submitted_at":"2026-06-25T07:27:09Z","title":"PhysEditWorld: A Large-Scale Dataset Toward Physics-Editable World Models","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-30T10:19:06.268547Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2606.26694"},"observation_digest":"sha256:89c34c6a59e9c0760c8ecf5e1674f454d73b27b4bc85e9930b72117326c7548e","observation_id":"1801af48-6d70-4ac6-95db-e513da3ce0c3","resolution":{"observed_at":"2026-06-30T12:04:39.292257Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-07-13T04:50:59.096166Z","title":"Video PreTraining (VPT): Learning to act by watching unlabeled online videos,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.09185","last_updated":"2026-07-10T08:20:27Z","snapshot_observed_at":"2026-08-02T15:57:39.943812Z","submitted_at":"2026-07-10T08:20:27Z","title":"Causally Debiased Latent Action Model for Embodied Action Conditioned World Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-13T04:50:59.096166Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2607.09185"},"observation_digest":"sha256:b194096c21ed0c3d04be4f8a47764d085fee8414a8cbd892041c45bb6a803078","observation_id":"f7f8dd81-8c56-470d-a611-dedaaa8bd49e","resolution":{"observed_at":"2026-07-13T04:50:59.096166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.11795","snapshot_observed_at":"2026-08-01T17:45:08.507221Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos , publisher =","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.17560","last_updated":"2026-07-20T05:09:36Z","snapshot_observed_at":"2026-08-05T13:38:45.921055Z","submitted_at":"2026-07-20T05:09:36Z","title":"Reinforcement Learning: From Algorithms To Foundation Models","version":1},"reference_index":122,"source":"arxiv_source","source_observed_at":"2026-08-01T17:45:08.507221Z"},"links":{"cited_paper":"/paper/2206.11795","citing_paper":"/paper/2607.17560"},"observation_digest":"sha256:3607ad84761fa17f98f862f79169349733b46436ec9d0afef552cbea428f5d23","observation_id":"c22d914d-ad88-437d-8db1-25ec4149b7b7","resolution":{"observed_at":"2026-08-01T17:45:08.507221Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2206.11795/citation-record","integrity":"/paper/2206.11795/integrity","json":"/paper/2206.11795/citation-record.json","paper":"/paper/2206.11795"},"outbound":[],"paper":{"arxiv_id":"2206.11795","last_updated":"2022-06-23T16:01:11Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-07-06T13:24:02.173160Z","submitted_at":"2022-06-23T16:01:11Z","title":"Video PreTraining (VPT): Learning to Act by Watching Unlabeled Online Videos"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 23 inbound Pith citation observations for arXiv:2206.11795."}