{"as_of":"2026-08-12T19:28:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:1eae092feca62c6d9eacc9885220fa73e6a9126ab0c56a2556ad8ada252146ba","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":33,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":33,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":33,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":33,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T00:48:43.700829Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T03:19:31.235420Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-11T14:57:26.605711Z","title":"Data-efficient reinforcement learning with self-predictive representa- tions","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2412.11484","last_updated":"2024-12-16T06:53:00Z","snapshot_observed_at":"2026-08-11T14:50:40.452739Z","submitted_at":"2024-12-16T06:53:00Z","title":"Efficient Policy Adaptation with Contrastive Prompt Ensemble for Embodied Agents","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T14:57:26.605711Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2412.11484"},"observation_digest":"sha256:3058bb57450eef12779f8f2ccdb471cbe30253d32a781f8254f929a2ddecf757","observation_id":"221b2ec4-cd73-4e73-b74b-6bf2952ae033","resolution":{"observed_at":"2026-08-11T14:57:26.605711Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-10T18:56:44.630670Z","title":"Schwarzer, A","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2501.10893","last_updated":"2025-01-18T22:34:41Z","snapshot_observed_at":"2026-08-10T22:27:08.878421Z","submitted_at":"2025-01-18T22:34:41Z","title":"Learn-by-interact: A Data-Centric Framework for Self-Adaptive Agents in Realistic Environments","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-10T18:56:44.630670Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2501.10893"},"observation_digest":"sha256:9940b0dbb35db5642d2fb6ae0afb8e37f10fb358862fdb05b0fb763cd3144b36","observation_id":"388d9945-f561-48f6-b01b-1f82ff7a5414","resolution":{"observed_at":"2026-08-10T18:56:44.630670Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-09T19:18:45.093521Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2502.00379","last_updated":"2025-06-12T16:28:52Z","snapshot_observed_at":"2026-08-12T11:06:40.262139Z","submitted_at":"2025-02-01T09:35:51Z","title":"Latent Action Learning Requires Supervision in the Presence of Distractors","version":5},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-09T19:18:45.093521Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2502.00379"},"observation_digest":"sha256:be24c30bcd5ee1f1ab8b32ae462a3d3cb7e9db2e1913a92d0852940d829ba4a0","observation_id":"b64c9e3c-9d54-4576-aa5f-63536dc43a18","resolution":{"observed_at":"2026-08-09T19:18:45.093521Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-07T15:23:19.335384Z","title":"Schwarzer, A","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2505.15345","last_updated":"2025-05-23T11:18:10Z","snapshot_observed_at":"2026-08-12T03:31:36.344688Z","submitted_at":"2025-05-21T10:19:49Z","title":"Hadamax Encoding: Elevating Performance in Model-Free Atari","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T15:23:19.335384Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2505.15345"},"observation_digest":"sha256:cc46c854a15365abbcbc0612f332ccb750ef6e16303a8a6fcffa63f56a483b55","observation_id":"f4549351-4bbe-424a-b645-6fdda44065d0","resolution":{"observed_at":"2026-08-07T15:23:19.335384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-07T14:36:12.851155Z","title":"Data-efficient reinforcement learning with self-predictive representations.arXiv preprint arXiv:2007.05929, 2020","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2505.18671","last_updated":"2025-05-24T12:18:19Z","snapshot_observed_at":"2026-08-09T05:57:00.511284Z","submitted_at":"2025-05-24T12:18:19Z","title":"Self-Supervised Evolution Operator Learning for High-Dimensional Dynamical Systems","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T14:36:12.851155Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2505.18671"},"observation_digest":"sha256:b2898773c7e546702a2e51360f9e4be9d8aa7610c6698dcb8c8d39d1559213a5","observation_id":"12076b5b-a367-4d47-a1dc-78d24df1c461","resolution":{"observed_at":"2026-08-07T14:36:12.851155Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-07T13:49:59.950868Z","title":"Schwarzer, A","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2505.21119","last_updated":"2025-06-02T16:01:02Z","snapshot_observed_at":"2026-08-07T21:47:08.434516Z","submitted_at":"2025-05-27T12:38:19Z","title":"Universal Value-Function Uncertainties","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-07T13:49:59.950868Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2505.21119"},"observation_digest":"sha256:a44af7f8110f85c5d6e3041cd02ef44aeda1e5bc6ae94f7af05fcba834aea079","observation_id":"20cb155f-a281-40ae-af06-71b95863f5fe","resolution":{"observed_at":"2026-08-07T13:49:59.950868Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-07T12:10:07.046569Z","title":"Data-efficient reinforcement learning with self-predictive representations","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2506.00563","last_updated":"2025-09-08T18:56:14Z","snapshot_observed_at":"2026-08-11T18:14:55.762917Z","submitted_at":"2025-05-31T13:43:41Z","title":"Understanding Behavioral Metric Learning: A Large-Scale Study on Distracting Reinforcement Learning Environments","version":2},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-07T12:10:07.046569Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2506.00563"},"observation_digest":"sha256:6a72dd0845c34ceee235d1f4910b474e1c3181e48881acdcdf896849d05ac0ab","observation_id":"aacff5eb-c4b3-4ba5-ac6d-3b9a6edfe064","resolution":{"observed_at":"2026-08-07T12:10:07.046569Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-07T10:47:32.908294Z","title":"Data-efficient reinforcement learn- ing with self-predictive representations.arXiv preprint arXiv:2007.05929,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.05418","last_updated":"2025-06-05T00:36:54Z","snapshot_observed_at":"2026-08-11T22:55:37.940262Z","submitted_at":"2025-06-05T00:36:54Z","title":"Self-Predictive Dynamics for Generalization of Vision-based Reinforcement Learning","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T10:47:32.908294Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2506.05418"},"observation_digest":"sha256:ac7ec05ea3f8f53c604f307572d7c7b6b14a154006bbe92e2c6a260bedcf6690","observation_id":"d2fbd067-0290-40da-a443-9be89c70d7c2","resolution":{"observed_at":"2026-08-07T10:47:32.908294Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2506.08902","last_updated":"2026-05-11T20:21:12Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-10T15:27:46Z","title":"Intention-Conditioned Flow Occupancy Models","version":4},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-05-19T10:24:52.160209Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2506.08902"},"observation_digest":"sha256:3ca9e74e31c79063fcd774c6d1fd5be4245e47f8ee1f4ab6ec4985fba0b87ac6","observation_id":"db04718e-1378-4546-b164-cd9b2dd2f3b0","resolution":{"observed_at":"2026-05-19T10:27:14.543460Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2506.10137","last_updated":"2026-04-19T14:49:08Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-11T19:32:41Z","title":"Self-Predictive Representations for Combinatorial Generalization in Behavioral Cloning","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-19T09:15:45.511104Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2506.10137"},"observation_digest":"sha256:71388b8da77e7a793e335d4438e0a433de99b7d5b7d8e140fcb52c506ac602d8","observation_id":"1bc51aeb-a694-41a9-9a81-2578af3f264e","resolution":{"observed_at":"2026-05-19T09:17:14.006364Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-06T23:34:41.901260Z","title":"Data-efficient reinforcement learning with self-predictive representations","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2506.17518","last_updated":"2025-06-20T23:47:04Z","snapshot_observed_at":"2026-08-10T10:35:48.301748Z","submitted_at":"2025-06-20T23:47:04Z","title":"A Survey of State Representation Learning for Deep Reinforcement Learning","version":1},"reference_index":111,"source":"arxiv_source","source_observed_at":"2026-08-06T23:34:41.901260Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2506.17518"},"observation_digest":"sha256:3c5d2511b1af8b1d27a033e812d7be02b6391499dde1a410c42508a300d0fbd8","observation_id":"63fedd0e-50f5-4233-b480-487939272151","resolution":{"observed_at":"2026-08-06T23:34:41.901260Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-06T22:30:08.476025Z","title":"Pierre Sermanet, Corey Lynch, Yevgen Chebotar, Jasmine Hsu, Eric Jang, Stefan Schaal, Sergey Levine, and Google Brain","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2506.21367","last_updated":"2025-06-26T15:16:35Z","snapshot_observed_at":"2026-08-09T00:13:58.355869Z","submitted_at":"2025-06-26T15:16:35Z","title":"rQdia: Regularizing Q-Value Distributions With Image Augmentation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T22:30:08.476025Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2506.21367"},"observation_digest":"sha256:7cbbd83a5859534fd129a15e91c6dc366113f16d0f75d963d4ce94df06a20b80","observation_id":"93c144cf-2dc6-48fa-8e0a-d91d7ec9ec76","resolution":{"observed_at":"2026-08-06T22:30:08.476025Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-06T19:12:58.088970Z","title":"Data-efficient reinforcement learning with self-predictive representations,","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2507.06326","last_updated":"2025-07-08T18:30:26Z","snapshot_observed_at":"2026-08-10T16:12:17.184621Z","submitted_at":"2025-07-08T18:30:26Z","title":"Sample-Efficient Reinforcement Learning Controller for Deep Brain Stimulation in Parkinson's Disease","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T19:12:58.088970Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2507.06326"},"observation_digest":"sha256:b3c46c6a62a3ab20c6cb018b8fbc6e2bdc0f2d368b22ed469124bae02197e395","observation_id":"d81c0600-606e-4fb7-a4e2-bbbdbaeac6d6","resolution":{"observed_at":"2026-08-06T19:12:58.088970Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-06T15:55:36.073397Z","title":"Data-efficient reinforcement learning with self-predictive representations, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.14736","last_updated":"2025-07-19T19:53:08Z","snapshot_observed_at":"2026-08-10T04:33:14.186488Z","submitted_at":"2025-07-19T19:53:08Z","title":"Balancing Expressivity and Robustness: Constrained Rational Activations for Reinforcement Learning","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-06T15:55:36.073397Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2507.14736"},"observation_digest":"sha256:15b0a873af4fbe8650fb6d5963f7391c8751610112fbeefb2835751662b8fdb5","observation_id":"3c24405b-13cf-45ee-9636-e5c7338eda5b","resolution":{"observed_at":"2026-08-06T15:55:36.073397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-04T08:03:20.231886Z","title":"Data-efficient reinforcement learning with self-predictive representations","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2510.23216","last_updated":"2026-06-02T10:40:21Z","snapshot_observed_at":"2026-08-12T01:52:38.126122Z","submitted_at":"2025-10-27T11:06:00Z","title":"Human-Like Goalkeeping in a Realistic Football Simulation: a Sample-Efficient Reinforcement Learning Approach","version":4},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-04T08:03:20.231886Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2510.23216"},"observation_digest":"sha256:f040ceb0bb60c9c5039b98cc8261eadd58fc02010cec78f0b77e01d731d714a9","observation_id":"44b43f42-3d22-423e-804b-418aea13d67b","resolution":{"observed_at":"2026-08-04T08:03:20.231886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2603.07083","last_updated":"2026-04-14T13:41:32Z","snapshot_observed_at":"2026-08-10T22:11:41.520272Z","submitted_at":"2026-03-07T07:41:28Z","title":"Dreamer-CDP: Improving Reconstruction-free World Models Via Continuous Deterministic Representation Prediction","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-15T14:46:18.645009Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2603.07083"},"observation_digest":"sha256:5233ecf5ae39d495ee41ce485620467e1cfd6565cf0e6e20188daf556a953184","observation_id":"1b91d14b-6869-4daa-ae7e-6b4283beb883","resolution":{"observed_at":"2026-05-15T14:50:05.076527Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2604.03023","last_updated":"2026-04-03T13:13:47Z","snapshot_observed_at":"2026-08-11T04:39:15.871543Z","submitted_at":"2026-04-03T13:13:47Z","title":"Behavior-Constrained Reinforcement Learning with Receding-Horizon Credit Assignment for High-Performance Control","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-13T19:30:23.447901Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2604.03023"},"observation_digest":"sha256:6688acfb2771f001ebec2da46de6765e2c6ad95e2ccb4bb37e99986c15c8b55c","observation_id":"308ea5c0-99ae-4d98-96ca-95b6ac0ed558","resolution":{"observed_at":"2026-05-13T19:33:10.238247Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2604.03208","last_updated":"2026-06-16T20:46:31Z","snapshot_observed_at":"2026-07-13T13:29:04.173053Z","submitted_at":"2026-04-03T17:32:36Z","title":"Hierarchical Planning with Latent World Models","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-05-13T20:13:58.991298Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2604.03208"},"observation_digest":"sha256:c0c9d614831d7ea25d9a98a9bc0ee5d654276edf164eba2f233d0bb2d006845e","observation_id":"16cf2cd0-cf7f-4dda-bd87-35236b489e52","resolution":{"observed_at":"2026-05-13T20:18:13.756848Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2604.15289","last_updated":"2026-04-16T17:53:16Z","snapshot_observed_at":"2026-08-12T14:33:35.219687Z","submitted_at":"2026-04-16T17:53:16Z","title":"Abstract Sim2Real through Approximate Information States","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-10T10:39:57.845600Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2604.15289"},"observation_digest":"sha256:877c8bd05e7a33f4027b1609817d249058e36de4b68a8fe3721bbbf7eef2024b","observation_id":"864bbf0c-aaa1-4dd7-b1cb-cb89dfaea35b","resolution":{"observed_at":"2026-05-10T10:44:38.083839Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2604.21130","last_updated":"2026-04-22T22:40:11Z","snapshot_observed_at":"2026-08-11T19:04:11.122840Z","submitted_at":"2026-04-22T22:40:11Z","title":"Self-Predictive Representation for Autonomous UAV Object-Goal Navigation","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-09T23:31:49.766698Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2604.21130"},"observation_digest":"sha256:a2e0bbc854d09ca6e0b09499d220ad1a877d54bb79e5bc9f4157f50fa42eab19","observation_id":"79ac0ef4-a397-45fa-bbc5-90a7ef811000","resolution":{"observed_at":"2026-05-11T14:06:03.869664Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2605.00347","last_updated":"2026-05-01T02:05:56Z","snapshot_observed_at":"2026-08-11T00:38:42.624588Z","submitted_at":"2026-05-01T02:05:56Z","title":"Odysseus: Scaling VLMs to 100+ Turn Decision-Making in Games via Reinforcement Learning","version":1},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-05-09T20:22:58.061772Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2605.00347"},"observation_digest":"sha256:2434d9a68cd22be4efe772205beaab868e2a34f223f0e7b298bc1d1e99a70e4e","observation_id":"2fd12924-e127-4769-bd31-e37d3ec5b271","resolution":{"observed_at":"2026-05-11T15:16:08.814352Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2605.07278","last_updated":"2026-05-08T05:43:33Z","snapshot_observed_at":"2026-08-01T00:55:28.214005Z","submitted_at":"2026-05-08T05:43:33Z","title":"Predictive but Not Plannable: RC-aux for Latent World Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-11T02:11:26.526418Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2605.07278"},"observation_digest":"sha256:c0b39aca079e7b5ee194583543b4fdad7bc67a4b1b1e86cd2d0157d584e70fe4","observation_id":"acb10449-413e-4bd2-a7be-4c65ccc9e078","resolution":{"observed_at":"2026-05-11T03:50:57.219964Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2605.09364","last_updated":"2026-05-10T06:27:20Z","snapshot_observed_at":"2026-08-11T06:51:53.423293Z","submitted_at":"2026-05-10T06:27:20Z","title":"Multi-scale Predictive Representations for Goal-conditioned Reinforcement Learning","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-12T03:35:37.739085Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2605.09364"},"observation_digest":"sha256:b5824bdd38cb076320f47eaedb803b57823461f83c79d61f8fd1d6bc4830090b","observation_id":"ef9b80c9-0dc1-4021-925d-761bb39ee295","resolution":{"observed_at":"2026-05-12T07:16:26.310030Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2605.24517","last_updated":"2026-05-23T11:08:46Z","snapshot_observed_at":"2026-08-12T19:11:11.805917Z","submitted_at":"2026-05-23T11:08:46Z","title":"ECHO: Terminal Agents Learn World Models for Free","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-30T14:57:03.095107Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2605.24517"},"observation_digest":"sha256:be423b57efe7f21445b6ca5458c58e51eecedb538e7145aac322f8e15b71b346","observation_id":"e22b60f7-f737-45f1-bdf4-862ca67203dc","resolution":{"observed_at":"2026-06-30T15:04:46.501267Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2605.25782","last_updated":"2026-05-26T02:20:50Z","snapshot_observed_at":"2026-08-03T21:15:11.785642Z","submitted_at":"2026-05-25T12:29:47Z","title":"ParkourFormer: Integrating Predictive Supervision and Sequence Modeling into Parkour Locomotion","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-29T21:37:13.833256Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2605.25782"},"observation_digest":"sha256:275ee0712b749505b54ae05779b3da25a97444b5496b6483c15883fb21dabbae","observation_id":"10d7ab89-9dcf-4f5d-990d-8842368fc212","resolution":{"observed_at":"2026-06-29T21:43:59.611385Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2605.27532","last_updated":"2026-05-26T18:06:19Z","snapshot_observed_at":"2026-08-12T08:56:19.012106Z","submitted_at":"2026-05-26T18:06:19Z","title":"SCALE-COMM: Shared, Contrastively-Aligned Latent Embeddings for MARL Communication","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-29T17:12:21.761535Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2605.27532"},"observation_digest":"sha256:02e2fa9000ac38b90f2c30a9af60cabc2831307469cae0880dd1ba5a493ec115","observation_id":"5cd1a620-781d-48d1-84a2-0666f1610554","resolution":{"observed_at":"2026-06-29T17:13:44.511218Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2606.18469","last_updated":"2026-06-16T20:28:39Z","snapshot_observed_at":"2026-08-07T17:34:01.303298Z","submitted_at":"2026-06-16T20:28:39Z","title":"Structured Representation Learning with Locally Linear Embeddings and Adaptive Feature Fusion","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-06-27T01:24:04.887225Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2606.18469"},"observation_digest":"sha256:a614d2c85f436031f1ed6236d82a4ced80fb92256c25297e5ce2324363a8b563","observation_id":"223507fe-f61c-4e66-b364-8134ecbf89a3","resolution":{"observed_at":"2026-07-03T20:18:57.470577Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":"2007.05929","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-04T03:19:31.235420Z","title":"D., Courville, A., and Bachman, P","venue":null,"work_id":"dd401acf-5b85-4e74-ad8c-49274e40b5b7","year":2007},"citing_paper":{"arxiv_id":"2606.20411","last_updated":"2026-06-18T15:58:48Z","snapshot_observed_at":"2026-08-08T10:53:09.850580Z","submitted_at":"2026-06-18T15:58:48Z","title":"Direct Advantage Estimation for Scalable and Sample-efficient Deep Reinforcement Learning","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-06-26T18:12:00.111067Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2606.20411"},"observation_digest":"sha256:5a1baaf184e475671d9234d761bc4f94d2446b8692ccced6045be61ecf0c00bb","observation_id":"015906af-65cb-4a51-b26d-6ec2cfeed68a","resolution":{"observed_at":"2026-07-04T03:19:31.238331Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-01T16:26:17.964188Z","title":"arXiv preprint arXiv:2007.05929 , year=","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2607.18004","last_updated":"2026-07-20T14:36:11Z","snapshot_observed_at":"2026-08-10T04:32:40.105685Z","submitted_at":"2026-07-20T14:36:11Z","title":"PAMD: Structured Adaptive Distances for Bisimulation Representations in Visual Reinforcement Learning","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-01T16:26:17.964188Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2607.18004"},"observation_digest":"sha256:4d023a8f07f9b08e49dbae6cc09f7028c09dd2520175ef10f7241f4dc4f3e39d","observation_id":"b3a18fbf-78e6-49f4-a2f1-ebcb1b4f1ffd","resolution":{"observed_at":"2026-08-01T16:26:17.964188Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-07-31T21:44:39.632400Z","title":"Data-efficient reinforcement learning with self-predictive representations.arXiv preprint arXiv:2007.05929, 2020","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2607.27973","last_updated":"2026-07-30T10:17:55Z","snapshot_observed_at":"2026-08-08T00:19:31.081356Z","submitted_at":"2026-07-30T10:17:55Z","title":"TAPO: Transition-Aware Policy Optimization for LLM Agents","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-31T21:44:39.632400Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2607.27973"},"observation_digest":"sha256:ae68d03c07e304eba996a13b9970fb4173b537c25234033447021bb4794723e4","observation_id":"5328c020-c40e-4471-a32f-1bef29fd3f08","resolution":{"observed_at":"2026-07-31T21:44:39.632400Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-10T04:22:03.644762Z","title":"Data-efficient reinforcement learning with self-predictive representations.arXiv preprint arXiv:2007.05929,","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2608.06595","last_updated":"2026-08-06T21:11:16Z","snapshot_observed_at":"2026-08-12T19:10:33.873806Z","submitted_at":"2026-08-06T21:11:16Z","title":"Flowing Through States: Neural ODE Regularization for Reinforcement Learning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T04:22:03.644762Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2608.06595"},"observation_digest":"sha256:b5e422adaf675d6809322ed01e8088fc93b0c146c46739c3f6574e38e6b4efe0","observation_id":"e991aa44-d435-4ed5-805b-87bdaf3f1930","resolution":{"observed_at":"2026-08-10T04:22:03.644762Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-10T10:11:39.526442Z","title":"arXiv preprint arXiv:2007.05929 , year =","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2608.07335","last_updated":"2026-08-07T15:29:15Z","snapshot_observed_at":"2026-08-12T19:12:30.908222Z","submitted_at":"2026-08-07T15:29:15Z","title":"Aftab: A Comprehensive Benchmark of CNN Encoders and Advanced Value Functions in Parallelized Q-Networks","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-10T10:11:39.526442Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2608.07335"},"observation_digest":"sha256:0d5fd8d3221b7bdbdcd74e4306b553323519fa1176808789d45565d69beb19e2","observation_id":"be2169b7-5fc5-4226-a1f6-2341a3d1462e","resolution":{"observed_at":"2026-08-10T10:11:39.526442Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.05929","snapshot_observed_at":"2026-08-12T00:48:43.700829Z","title":"arXiv preprint arXiv:2007.05929 , year=","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2608.07870","last_updated":"2026-08-08T02:44:43Z","snapshot_observed_at":"2026-08-12T19:15:08.604220Z","submitted_at":"2026-08-08T02:44:43Z","title":"V-Simba: Unleashing the Architectural Potential of RL in Visual Continuous Control","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-12T00:48:43.700829Z"},"links":{"cited_paper":"/paper/2007.05929","citing_paper":"/paper/2608.07870"},"observation_digest":"sha256:3491b9fbe7dbce7b5087dd5195e6533f0a634a07c2a80e1d7d35f61923e0993a","observation_id":"e823bf09-6b87-401b-a16f-93f22e1f97e0","resolution":{"observed_at":"2026-08-12T00:48:43.700829Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2007.05929/citation-record","integrity":"/paper/2007.05929/integrity","json":"/paper/2007.05929/citation-record.json","paper":"/paper/2007.05929"},"outbound":[],"paper":{"arxiv_id":"2007.05929","last_updated":"2021-05-20T09:15:57Z","latest_version":4,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-10T04:32:16.941126Z","submitted_at":"2020-07-12T07:38:15Z","title":"Data-Efficient Reinforcement Learning with Self-Predictive Representations"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 33 inbound Pith citation observations for arXiv:2007.05929."}