{"as_of":"2026-08-09T15:44:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:81196c041d70bc5467e774a483a197c0be454ef86f59670cd094974ec01ebbae","coverage":[{"denominator":69,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":69,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T14:27:41.185161Z","state":"measured"},{"denominator":84,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":84,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":15,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":15,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T13:52:12.773873Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T13:09:50.104673Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-08-07T13:52:12.773873Z","title":"From foresight to forethought: Vlm-in-the-loop policy steering via latent alignment.arXiv preprint arXiv:2502.01828, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.20897","last_updated":"2025-06-22T13:53:33Z","snapshot_observed_at":"2026-08-08T02:45:56.557199Z","submitted_at":"2025-05-27T08:40:20Z","title":"Cross from Left to Right Brain: Adaptive Text Dreamer for Vision-and-Language Navigation","version":2},"reference_index":95,"source":"pdf_text","source_observed_at":"2026-08-07T13:52:12.773873Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2505.20897"},"observation_digest":"sha256:735bbf79ab856ca65249bdfa692a50f82c016758cde1fedf2adc1ff8d40eba35","observation_id":"1d242b21-9e3d-4f4d-8f34-46a78bc0a80c","resolution":{"observed_at":"2026-08-07T13:52:12.773873Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-08-07T00:48:49.489677Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12678","last_updated":"2025-06-15T01:17:30Z","snapshot_observed_at":"2026-08-09T11:54:17.542443Z","submitted_at":"2025-06-15T01:17:30Z","title":"Adapting by Analogy: OOD Generalization of Visuomotor Policies via Functional Correspondence","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T00:48:49.489677Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2506.12678"},"observation_digest":"sha256:b400dfa89d37f853a9f899d5406f1d48222af3dafd2285b805498d64188f3c0e","observation_id":"2bde91d5-ee5e-4bfc-b3bb-ef3d5ec8dc41","resolution":{"observed_at":"2026-08-07T00:48:49.489677Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-08-05T23:06:33.877275Z","title":"From foresight to forethought: Vlm-in- the-loop policy steering via latent alignment","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.05941","last_updated":"2026-07-01T21:54:22Z","snapshot_observed_at":"2026-08-08T16:23:41.867796Z","submitted_at":"2025-08-08T02:07:08Z","title":"Latent Policy Barrier: Learning Robust Visuomotor Policies by Staying In-Distribution","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-05T23:06:33.877275Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2508.05941"},"observation_digest":"sha256:517ce84f115b5339a588efb1bd5c20ef5a4a97cb48305371f69f8debbc6430e8","observation_id":"216ce25f-48eb-4a3c-8dee-35ce4f13e195","resolution":{"observed_at":"2026-08-05T23:06:33.877275Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-08-05T13:52:52.539441Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.00271","last_updated":"2026-06-17T20:45:44Z","snapshot_observed_at":"2026-08-08T09:36:37.528608Z","submitted_at":"2025-08-29T22:56:32Z","title":"Learn from What We HAVE: History-Aware VErifier that Reasons about Past Interactions Online","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-05T13:52:52.539441Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2509.00271"},"observation_digest":"sha256:6ab96a27b2ef8b145851892e83bf1dac4bc428573a9243832ed49c3d71fdc125","observation_id":"86e1d926-f588-4dd3-a39f-fbc0046a07b9","resolution":{"observed_at":"2026-08-05T13:52:52.539441Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-08-03T14:10:36.253138Z","title":"From foresight to forethought: Vlm-in-the-loop policy steer- ing via latent alignment.arXiv preprint arXiv:2502.01828,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2512.21430","last_updated":"2026-06-03T18:06:11Z","snapshot_observed_at":"2026-08-03T15:13:16.061241Z","submitted_at":"2025-12-24T21:36:34Z","title":"EVE: A Generator-Verifier System for Generative Policies","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-03T14:10:36.253138Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2512.21430"},"observation_digest":"sha256:9ec5d31256fed5886c866e2db9c935a0625b1ed44fc695a3857d19854d98fd33","observation_id":"726db4cb-2725-42bb-baaf-08956206569f","resolution":{"observed_at":"2026-08-03T14:10:36.253138Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":"2502.01828","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-07-04T13:09:50.104673Z","title":"From foresight to forethought: Vlm-in-the-loop policy steering via latent alignment","venue":null,"work_id":"a2cfc3a7-33f9-488f-b5cc-047547da4cc5","year":2025},"citing_paper":{"arxiv_id":"2602.13193","last_updated":"2026-04-06T03:04:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-02-13T18:57:56Z","title":"Steerable Vision-Language-Action Policies for Embodied Reasoning and Hierarchical Control","version":3},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-05-15T22:05:39.797848Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2602.13193"},"observation_digest":"sha256:d8b9e045ecee57ca663ad2d5bacfedc1bfc47faf9a3c5c72ef1ee38b20e807ae","observation_id":"3b2a703c-f0ff-41bf-89aa-ade26ebe71fc","resolution":{"observed_at":"2026-05-15T22:06:42.942117Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":"2502.01828","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-07-04T13:09:50.104673Z","title":"From foresight to forethought: Vlm-in-the-loop policy steering via latent alignment","venue":null,"work_id":"a2cfc3a7-33f9-488f-b5cc-047547da4cc5","year":2025},"citing_paper":{"arxiv_id":"2606.01036","last_updated":"2026-05-31T05:56:28Z","snapshot_observed_at":"2026-07-06T23:41:39.172310Z","submitted_at":"2026-05-31T05:56:28Z","title":"Position: Good Embodied Reward Models Need Bad Behavior Data","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-28T17:18:17.337336Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2606.01036"},"observation_digest":"sha256:c9482b741be99d73320a27803196f29909f40eab58f24d2e2af41881718c7cfd","observation_id":"4eca8077-1d50-42af-bf21-787a67b6195b","resolution":{"observed_at":"2026-06-28T17:22:24.960996Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":"2502.01828","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-07-04T13:09:50.104673Z","title":"From foresight to forethought: Vlm-in-the-loop policy steering via latent alignment","venue":null,"work_id":"a2cfc3a7-33f9-488f-b5cc-047547da4cc5","year":2025},"citing_paper":{"arxiv_id":"2606.03954","last_updated":"2026-06-02T17:42:17Z","snapshot_observed_at":"2026-08-04T08:58:29.065242Z","submitted_at":"2026-06-02T17:42:17Z","title":"VLESA: Vision-Language Embodied Safety Agent for Human Activity Monitoring","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-28T10:30:47.513529Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2606.03954"},"observation_digest":"sha256:cf127bc6894c0996feaba2ae40a80de37e4e3e636b98fc6e6301f4613bfccaa7","observation_id":"35950b30-a425-4cde-9359-48840d7643b1","resolution":{"observed_at":"2026-07-02T02:56:29.043003Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":"2502.01828","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-07-04T13:09:50.104673Z","title":"From foresight to forethought: Vlm-in-the-loop policy steering via latent alignment","venue":null,"work_id":"a2cfc3a7-33f9-488f-b5cc-047547da4cc5","year":2025},"citing_paper":{"arxiv_id":"2606.12366","last_updated":"2026-06-10T17:34:25Z","snapshot_observed_at":"2026-07-06T23:51:17.719874Z","submitted_at":"2026-06-10T17:34:25Z","title":"APT: Action Expert Pretraining Improves Instruction Generalization of Vision-Language-Action Policies","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-06-27T09:49:56.894300Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2606.12366"},"observation_digest":"sha256:79f36cfc3c50d8630bcfe2c01f4d2318b3db1e3f9d32887cc60f8e46f0804b52","observation_id":"26b762c2-5d3a-4258-bc31-2d7175a5612b","resolution":{"observed_at":"2026-07-03T10:48:02.727730Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":"2502.01828","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-07-04T13:09:50.104673Z","title":"From foresight to forethought: Vlm-in-the-loop policy steering via latent alignment","venue":null,"work_id":"a2cfc3a7-33f9-488f-b5cc-047547da4cc5","year":2025},"citing_paper":{"arxiv_id":"2606.18589","last_updated":"2026-06-17T01:28:07Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-06-17T01:28:07Z","title":"DREAM-Chunk: Reactive Action Chunking with Latent World Model","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-26T21:25:03.953677Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2606.18589"},"observation_digest":"sha256:a4c0edef1aa405d8e72297d8a485499a064b31fb77401c474448dfddc66c2ee0","observation_id":"41b0bc13-cd81-4e9c-8519-82d281cb7887","resolution":{"observed_at":"2026-07-04T00:09:15.270179Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":"2502.01828","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-07-04T13:09:50.104673Z","title":"From foresight to forethought: Vlm-in-the-loop policy steering via latent alignment","venue":null,"work_id":"a2cfc3a7-33f9-488f-b5cc-047547da4cc5","year":2025},"citing_paper":{"arxiv_id":"2606.20458","last_updated":"2026-06-18T16:40:07Z","snapshot_observed_at":"2026-08-07T03:20:01.050993Z","submitted_at":"2026-06-18T16:40:07Z","title":"Slow Brain, Fast Planner: Latency-Resilient VLM-Augmented Urban Navigation","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-26T17:16:52.173943Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2606.20458"},"observation_digest":"sha256:3a27051153eb90256f2a0a0c20aa2c4ea5cf26ce67e5793d4fb6426a7a6d250c","observation_id":"ec05cca0-bdac-44cb-bd0e-41436194c5f3","resolution":{"observed_at":"2026-07-04T04:09:33.934494Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":"2502.01828","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-07-04T13:09:50.104673Z","title":"From foresight to forethought: Vlm-in-the-loop policy steering via latent alignment","venue":null,"work_id":"a2cfc3a7-33f9-488f-b5cc-047547da4cc5","year":2025},"citing_paper":{"arxiv_id":"2606.21572","last_updated":"2026-06-19T16:14:12Z","snapshot_observed_at":"2026-07-06T23:56:31.454439Z","submitted_at":"2026-06-19T16:14:12Z","title":"Robot Critics that Sweat the Small Stuff","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-26T14:20:37.905355Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2606.21572"},"observation_digest":"sha256:eba6bcf58fefa146b278bf5ec926d7731e66bb9ea5f09025181407aedbfa0243","observation_id":"bb1f8008-4fff-4629-9b57-ac41cd6a5cee","resolution":{"observed_at":"2026-07-04T06:39:37.258087Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":"2502.01828","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-07-04T13:09:50.104673Z","title":"From foresight to forethought: Vlm-in-the-loop policy steering via latent alignment","venue":null,"work_id":"a2cfc3a7-33f9-488f-b5cc-047547da4cc5","year":2025},"citing_paper":{"arxiv_id":"2606.22998","last_updated":"2026-06-23T07:19:46Z","snapshot_observed_at":"2026-07-06T23:57:47.047975Z","submitted_at":"2026-06-22T08:14:35Z","title":"TEXEDO : Test Time Scaling for Controller-aware Language-conditioned Humanoid Motion Generation","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-26T08:50:47.113217Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2606.22998"},"observation_digest":"sha256:59de39410f3f925126338dfd9dcfc468fe90e6b49de5a8824e1963d7bd588d7e","observation_id":"f92489a7-9155-46f6-9020-878fce3872e4","resolution":{"observed_at":"2026-07-04T10:29:44.868915Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":"2502.01828","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-07-04T13:09:50.104673Z","title":"From foresight to forethought: Vlm-in-the-loop policy steering via latent alignment","venue":null,"work_id":"a2cfc3a7-33f9-488f-b5cc-047547da4cc5","year":2025},"citing_paper":{"arxiv_id":"2606.26588","last_updated":"2026-06-25T04:20:08Z","snapshot_observed_at":"2026-08-07T09:51:22.326073Z","submitted_at":"2026-06-25T04:20:08Z","title":"Inference-Time Robot Behavior Steering through Physically-Aware Reconfiguration of Task-Structure","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-06-26T05:32:52.472573Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2606.26588"},"observation_digest":"sha256:c31bbb0beaf2e996d4a0f3be101bab0384cb43625436d119b068c6493d97b22f","observation_id":"6617e871-9fb1-4b45-b2fe-ab32a93ca062","resolution":{"observed_at":"2026-07-04T13:09:50.106851Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.01828","snapshot_observed_at":"2026-07-12T06:30:52.991927Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.02865","last_updated":"2026-07-03T02:08:50Z","snapshot_observed_at":"2026-08-08T05:47:50.300415Z","submitted_at":"2026-07-03T02:08:50Z","title":"DREAMSTEER: Latent World Models Can Steer VLA Policies During Deployment Without Any Finetuning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-12T06:30:52.991927Z"},"links":{"cited_paper":"/paper/2502.01828","citing_paper":"/paper/2607.02865"},"observation_digest":"sha256:a7a6554a4b0a0d9377fc7621707c22f16c0872b0beb6dd6e0b5b72dad652619d","observation_id":"6d0e9dd3-8709-47ac-879a-b08065b458ee","resolution":{"observed_at":"2026-07-12T06:30:52.991927Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2502.01828/citation-record","integrity":"/paper/2502.01828/integrity","json":"/paper/2502.01828/citation-record.json","paper":"/paper/2502.01828"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.506234Z","title":"Unpacking failure modes of generative policies: Runtime monitoring of consistency and progress","venue":null,"work_id":"f0767e43-d4cf-41f8-b813-3299c29d37a7","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.900935Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:a10913e6db6a341cf9c2ba43229b154c31575c4f97918b657a7ca3f457deb948","observation_id":"8b589a07-d766-4520-a45d-36ff9a022fca","resolution":{"observed_at":"2026-08-09T14:27:42.510508Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.492420Z","title":"Anonymous title","venue":null,"work_id":"fa47f51f-1ea4-4d08-80e5-b34968f49e3e","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.906218Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:ce1aa87ad2a3b828cc122c553ec657161540fc7bd88def1abe6b8269356edd48","observation_id":"4939cea9-0537-4b07-b364-3a64607abeff","resolution":{"observed_at":"2026-08-09T14:27:42.496820Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.478766Z","title":"Policy search by dynamic programming","venue":null,"work_id":"f71bd879-a751-4966-ba02-4a40bc598dd6","year":2003},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.910683Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:7b14cb82c4ebe70baa808d5821d8faaa85d4bc842e01585ba711b574e6aa5ae9","observation_id":"5610a58e-881f-404d-bab5-66c38e0b8f41","resolution":{"observed_at":"2026-08-09T14:27:42.482960Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2212.06817","last_updated":"2023-08-11T17:45:27Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-12-13T18:55:15Z","title":"RT-1: Robotics Transformer for Real-World Control at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.06817","snapshot_observed_at":"2026-08-09T14:27:40.915216Z","title":"Rt-1: Robotics transformer for real-world control at scale","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.915216Z"},"links":{"cited_paper":"/paper/2212.06817","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:00451ed5c5925ce56b435934c49ea348e733358636fc4fd006cb845e7d9b2fb7","observation_id":"a3f063cd-8682-4b32-81a6-bd9afebf7ac7","resolution":{"observed_at":"2026-08-09T14:27:40.915216Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15818","last_updated":"2023-07-28T21:18:02Z","snapshot_observed_at":"2026-08-02T16:17:50.621617Z","submitted_at":"2023-07-28T21:18:02Z","title":"RT-2: Vision-Language-Action Models Transfer Web Knowledge to Robotic Control","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15818","snapshot_observed_at":"2026-08-09T14:27:40.920085Z","title":"Rt-2: Vision-language-action models transfer web knowledge to robotic control","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.920085Z"},"links":{"cited_paper":"/paper/2307.15818","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:2f4376bf981e81b6d9559bbc7a11e8c8bd79f5b4b13f848238152a8cd27bb865","observation_id":"84edf15a-0331-4ac3-8a42-b324757204a8","resolution":{"observed_at":"2026-08-09T14:27:40.920085Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:40.924859Z","title":"Diffusion policy: Visuomotor policy learning via action diffusion","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.924859Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:d44fb825beac6a884f875f7add367e008bb67e16fbda30111b65d22074c34786","observation_id":"163aa7f5-81cd-4f82-93a6-5c6530c37eb7","resolution":{"observed_at":"2026-08-09T14:27:40.924859Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:40.929885Z","title":"Universal manipulation interface: In- the-wild robot teaching without in-the-wild robots","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.929885Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:4fe53af53280b5469bc0c28dd5ce0ce76eea9e0405daed3dd4a0f70891deedb3","observation_id":"756d77b4-ce36-4aee-93fe-55aac4b37b75","resolution":{"observed_at":"2026-08-09T14:27:40.929885Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:40.934232Z","title":"Agibot world colosseum","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.934232Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:14fbd7755868169b6a808c6845736022bccb0979461659e5bc43200285c2bf7a","observation_id":"814dc27f-aa7d-47a2-850e-720e9bd79542","resolution":{"observed_at":"2026-08-09T14:27:40.934232Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.438182Z","title":"The complexity of theorem-proving procedures","venue":null,"work_id":"d857bb68-13f5-41cc-b892-a38d60bd1208","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.938592Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:d5058bd46484baf1744ba90264fcf21e7f6e433d63116486462b724c2629bf46","observation_id":"8419f360-de04-48a5-98cd-233a762884b1","resolution":{"observed_at":"2026-08-09T14:27:42.442408Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.425407Z","title":"Aha: A vision- language-model for detecting and reasoning over failures in robotic manipulation","venue":null,"work_id":"1df5c206-4693-47bd-9921-bd5bcbe4c62a","year":2025},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.943037Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:32d2ff52169d6092242ae3f6014d7e5c84b3065131abe9e86d00404517675bb0","observation_id":"1a827c34-c1b2-45c5-98f3-2000790e90cc","resolution":{"observed_at":"2026-08-09T14:27:42.429294Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-09T14:27:40.947290Z","title":"The llama 3 herd of models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.947290Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:7c5cc2a885a8905c02da52f0305caf362bcdd8aea1c6beb122891d53f14a9782","observation_id":"2c8ae09d-3ebe-4103-a2dc-dec892ea20ac","resolution":{"observed_at":"2026-08-09T14:27:40.947290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.13162","last_updated":"2025-04-02T16:32:52Z","snapshot_observed_at":"2026-08-07T16:57:58.030473Z","submitted_at":"2025-03-17T13:35:55Z","title":"Efficient Imitation under Misspecification","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.13162","snapshot_observed_at":"2026-08-09T14:27:40.951522Z","title":"Efficient imitation under misspecifi- cation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.951522Z"},"links":{"cited_paper":"/paper/2503.13162","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:6890cfa533cc3639148dec9e5d44b4bc74269fd85da3ba6ba2ebcab1207a7204","observation_id":"9ceea0d6-912b-4ddc-97c6-e599a2b5e76d","resolution":{"observed_at":"2026-08-09T14:27:40.951522Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.413091Z","title":"Rh20t: A robotic dataset for learning diverse skills in one-shot","venue":null,"work_id":"8ec8a5d7-d1e1-4abb-9c6a-bcc22e48f7e5","year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.955734Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:20bf3969c5067e83098b49c1f2a024ba626751f63538e66a4f0f3a711754c15b","observation_id":"4057891c-f703-44c0-ad0e-4adfbdf713ca","resolution":{"observed_at":"2026-08-09T14:27:42.417197Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.399341Z","title":"Zhao, and Chelsea Finn","venue":null,"work_id":"57b628e8-e177-4da3-b241-fd80c317b5e9","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.959466Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:c6664dd2858fb9a23e2cc07bdd704e55f69eb634527d9314392ba51923a3e4bc","observation_id":"02b8ebf6-7c3c-409f-8292-b3ce412bdc6b","resolution":{"observed_at":"2026-08-09T14:27:42.404070Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.386623Z","title":"Letter to john von neumann, 1956","venue":null,"work_id":"578f1252-36fb-4a9a-b306-052b02173cdb","year":1956},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.963284Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:599c1c3e34057158cca17b60862fc25a55aadb02538a4ddc2d66fa6596f04141","observation_id":"2264474e-7fb7-4ee9-bba1-bf8f95cc88e0","resolution":{"observed_at":"2026-08-09T14:27:42.390958Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.372881Z","title":"Task success is not enough: Investigating the use of video-language models as behavior critics for catching undesirable agent behav- iors","venue":null,"work_id":"f3c771d3-eea9-4091-963d-39b76d81dafd","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.967090Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:1c5b4f2c54ff52673d24c0ccfa152875ea97eb90d92e5dab7efe7ed5479e421e","observation_id":"88eab4e3-1eac-4a74-ae05-a78cf9932e6f","resolution":{"observed_at":"2026-08-09T14:27:42.377554Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.359058Z","title":"Inverse reward design","venue":null,"work_id":"3e6e853d-81d8-47e2-9fda-cba437cdcc8b","year":2017},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.970700Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:06b88c9fb30c994b094353f1fe5cc977ef0c34d5cca1840e7af25750cc7e145c","observation_id":"b574df31-06e6-4515-8b46-a66648f4099a","resolution":{"observed_at":"2026-08-09T14:27:42.363397Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.04104","last_updated":"2024-04-17T17:41:20Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-01-10T18:12:16Z","title":"Mastering Diverse Domains through World Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.04104","snapshot_observed_at":"2026-08-09T14:27:40.974991Z","title":"Mastering diverse domains through world models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.974991Z"},"links":{"cited_paper":"/paper/2301.04104","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:411d326b79544749ec79a93e031a66daf6144e15d58575b1b79332c3239a19a5","observation_id":"e8641933-51fc-4a0d-a8f5-37a9b77a8c88","resolution":{"observed_at":"2026-08-09T14:27:40.974991Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.01971","last_updated":"2024-10-02T19:29:24Z","snapshot_observed_at":"2026-08-06T10:08:39.258355Z","submitted_at":"2024-10-02T19:29:24Z","title":"Run-time Observation Interventions Make Vision-Language-Action Models More Visually Robust","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.01971","snapshot_observed_at":"2026-08-09T14:27:40.978978Z","title":"Run-time observation interventions make vision- language-action models more visually robust","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.978978Z"},"links":{"cited_paper":"/paper/2410.01971","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:8877137c3f25600ada91eeddbcae124c351ce4e0688a35ccda44574fdb6bc054","observation_id":"3ce77235-0230-4751-ad79-d126cc403196","resolution":{"observed_at":"2026-08-09T14:27:40.978978Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.344776Z","title":"LoRA: Low-rank adaptation of large language models","venue":null,"work_id":"2cd901b7-e3d8-43cd-9857-e030befd9b39","year":2022},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.983204Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:7dc93ef02b8bd97b388ad8847a944a61cab5fbc1e36e40eb4381ce3b0382e49b","observation_id":"995493b8-dd95-4762-a7cc-8689aa7f4ef1","resolution":{"observed_at":"2026-08-09T14:27:42.349383Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.08782","last_updated":"2024-10-01T08:54:53Z","snapshot_observed_at":"2026-07-06T17:01:35.415575Z","submitted_at":"2023-12-14T10:02:55Z","title":"Toward General-Purpose Robots via Foundation Models: A Survey and Meta-Analysis","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.08782","snapshot_observed_at":"2026-08-09T14:27:40.986804Z","title":"Toward general- purpose robots via foundation models: A survey and meta-analysis","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.986804Z"},"links":{"cited_paper":"/paper/2312.08782","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:30be45133b68d95e7a8a48d51993f82d189c0f541c810ed8c945a2a62d0bf4d8","observation_id":"bae29f8b-e45c-4959-8d52-57c7a8c6fc1b","resolution":{"observed_at":"2026-08-09T14:27:40.986804Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.19112","last_updated":"2025-01-08T06:45:02Z","snapshot_observed_at":"2026-07-06T20:13:17.794908Z","submitted_at":"2024-12-26T08:11:41Z","title":"Future Success Prediction in Open-Vocabulary Object Manipulation Tasks Based on End-Effector Trajectories","version":2},"cited_work":{"arxiv_id":"2412.19112","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.19112","snapshot_observed_at":"2026-08-09T14:27:41.689865Z","title":"Future Success Prediction in Open-Vocabulary Object Manipulation Tasks Based on End-Effector Trajectories","venue":"cs.RO","work_id":"ea4640f5-f89b-4952-8b5e-e1dfbff2819c","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.990564Z"},"links":{"cited_paper":"/paper/2412.19112","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:d449054a52a6969a1ede0bda8530f9ff4c10ab7e252e98656e905152fa6ed6ab","observation_id":"08f89959-cb68-415c-ae06-42db6a63530d","resolution":{"observed_at":"2026-08-09T14:27:41.694875Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.330640Z","title":"Behavior generation with latent actions","venue":null,"work_id":"a40b42c3-fbcd-472d-8432-93bf01103701","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.994374Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:8369d0146016eda0b9c695cf25096bea4e1e2c74081648ca7df32154c383f869","observation_id":"5fd43d77-c435-4628-94fd-96ecc74c377e","resolution":{"observed_at":"2026-08-09T14:27:42.335295Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.316343Z","title":"Model-based runtime monitoring with interac- tive imitation learning","venue":null,"work_id":"82912a23-8204-4ffc-970f-3b444f91d815","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:40.997791Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:78298d2c060eb256ca693f9568d40e201ca883f7a8825c18de53f820a1221143","observation_id":"ada5ad8c-03a0-4f70-a6c6-a8f565440a57","resolution":{"observed_at":"2026-08-09T14:27:42.320888Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.301866Z","title":"Multi-task interactive robot fleet learning with visual world models","venue":null,"work_id":"d480540f-32bb-45af-af9d-b44ccb26b666","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.001820Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:baa005f2073b8e2212bba6e31ba5d41d9bef179446ceb85c716a14b78194cfb0","observation_id":"4e3c5044-ce33-49b8-9e34-78ccd7ac6f90","resolution":{"observed_at":"2026-08-09T14:27:42.307056Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.286662Z","title":"Reflect: Summarizing robot experiences for failure explanation and correction","venue":null,"work_id":"9ef34cbd-8518-41b0-abc9-9637ab626f78","year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.006179Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:63a88fba789a3e13aa3818d30eeb4d22d7bccecdca87c577ed80ccc7adbd0502","observation_id":"db3ba704-e8e2-4b2a-9c7d-6537307fdffe","resolution":{"observed_at":"2026-08-09T14:27:42.291819Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.272047Z","title":"Steering your generalists: Improving robotic foundation models via value guidance","venue":null,"work_id":"dc9a1f76-2a19-4119-8301-4b4c3b79ed39","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.010291Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:cd4c60b9215ad0b34a8d26b320e991c1c3771c9fc865ac016f35f3b80d28ac50","observation_id":"aa27cf27-99c3-45e1-a031-7dd0b76381a2","resolution":{"observed_at":"2026-08-09T14:27:42.276220Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.258902Z","title":"Algorithms for inverse reinforcement learning","venue":null,"work_id":"a1e568d6-8c09-4fe1-a4fb-e24367e33bdb","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.014278Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:fac76b2afdc16bab5c9e9d96afded6b3b86938ec3e60e80f7d98e646a62a8e60","observation_id":"290f5333-c69b-4b1e-8b17-512f6fd39f3c","resolution":{"observed_at":"2026-08-09T14:27:42.263294Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-09T14:27:41.018516Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.018516Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:40669609cc08120e2a907c99105fc1ab14ef92d3ba26c5aaf52b27f051c87a46","observation_id":"33f88971-d952-4f22-9ced-e1124a191cbd","resolution":{"observed_at":"2026-08-09T14:27:41.018516Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.245441Z","title":"Dinov2: Learning robust visual features without supervision","venue":null,"work_id":"ab9a39c9-1825-43b9-8e20-4744557c99e9","year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.023135Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:ee43a9b9fc7591552fc7595e01f68c7e450b3c2710d6a1c10f169211d72249e3","observation_id":"2cdaa777-3366-49a6-a27b-eaa3290f47c2","resolution":{"observed_at":"2026-08-09T14:27:42.249505Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.232437Z","title":"Learning to search: Functional gradient techniques for imitation learning","venue":null,"work_id":"e75d078c-d27a-40c1-aa24-49c02c841007","year":2009},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.027135Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:14baa36b578b9c513d7daed8a1e5bbde591302e75b5ca0c6613c4c5f2f5da8ef","observation_id":"b995d317-1f22-406a-8eac-4771d8292d77","resolution":{"observed_at":"2026-08-09T14:27:42.236428Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08848","last_updated":"2024-06-05T00:17:20Z","snapshot_observed_at":"2026-07-06T17:29:47.564468Z","submitted_at":"2024-02-13T23:29:09Z","title":"Hybrid Inverse Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08848","snapshot_observed_at":"2026-08-09T14:27:41.031294Z","title":"Hybrid inverse rein- forcement learning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.031294Z"},"links":{"cited_paper":"/paper/2402.08848","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:9193b9d27b705232f3578d025db3bc6ef3ad6fd6d0fca280f0e9700506c5a3b5","observation_id":"d5c73c1f-6834-4082-b65e-c70b808a7ace","resolution":{"observed_at":"2026-08-09T14:27:41.031294Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.218764Z","title":"Multimodal diffusion transformer: Learning versatile behavior from multimodal goals","venue":null,"work_id":"a3bfcebd-441c-4f93-affc-3bcc5ac6f65a","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.035562Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:929c0142c9e38f520c363b3fd2499c84671a4f1dc9b61b50547c72444a04dbba","observation_id":"51b14143-d27a-4c57-8211-92751485e476","resolution":{"observed_at":"2026-08-09T14:27:42.223132Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.040059Z","title":"Efficient reductions for imitation learning","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.040059Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:614acb77e9a63c8b37d4d6aabfbd94440e675dc8b7d3596f3833455503a18782","observation_id":"4db534f7-4a9a-4530-baf3-cbb4b00151f4","resolution":{"observed_at":"2026-08-09T14:27:41.040059Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.044257Z","title":"A reduction of imitation learning and structured prediction to no-regret online learning","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.044257Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:6771a0ddee48b9223a0d2d608d8a715d03675a87bb31dd27288c42fc5c0eae59","observation_id":"d4b4c681-031b-4c51-b24c-b93c314caa19","resolution":{"observed_at":"2026-08-09T14:27:41.044257Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.186546Z","title":"Motionlm: Multi-agent mo- tion forecasting as language modeling","venue":null,"work_id":"825e428e-a30c-4d54-afbc-f7d0bf226486","year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.048406Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:de2750a05d273ae8a9585a13b5152b8007b757915af8aeedb7eafeee54bc1e40","observation_id":"dd7f8b25-e4c9-4a17-9b8c-c0c4df64830f","resolution":{"observed_at":"2026-08-09T14:27:42.190859Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.173292Z","title":"Shafiullah, Siyuan","venue":null,"work_id":"87c3065a-4ecc-40ee-8762-0c4076d3ba0e","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.052730Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:e24e455b97a56375340e84d7da798f33862fabcdd7f4973ef144247866baca45","observation_id":"56d26b11-195b-433e-bb05-7dced5f8c6fb","resolution":{"observed_at":"2026-08-09T14:27:42.177466Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1604.06915","last_updated":"2016-04-23T15:13:43Z","snapshot_observed_at":"2026-08-06T03:30:43.330377Z","submitted_at":"2016-04-23T15:13:43Z","title":"On the Sample Complexity of End-to-end Training vs. Semantic Abstraction Training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1604.06915","snapshot_observed_at":"2026-08-09T14:27:41.057081Z","title":"On the sample complexity of end-to-end training vs","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.057081Z"},"links":{"cited_paper":"/paper/1604.06915","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:b06c8751cfb25dbfdbf6c392393f53d07fea6e58fc2d4fa9dd96234c166525fb","observation_id":"830ee759-6941-415a-89d7-067b2ea4ccce","resolution":{"observed_at":"2026-08-09T14:27:41.057081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.156265Z","title":"Real-time anomaly detection and reactive planning with large lan- guage models","venue":null,"work_id":"c0531e08-f84c-4fba-bd6f-b8e8b39d6250","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.061526Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:203495957d5e215bc51397ff6d0765e96698fdc021f3bbd937d2832ffb4d06fd","observation_id":"137493bb-d8fb-40e7-8cd6-87423367d4b3","resolution":{"observed_at":"2026-08-09T14:27:42.161527Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.06718","last_updated":"2023-03-11T11:47:54Z","snapshot_observed_at":"2026-08-02T06:10:04.416926Z","submitted_at":"2022-10-13T04:19:05Z","title":"Hybrid RL: Using Both Offline and Online Data Can Make RL Efficient","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.06718","snapshot_observed_at":"2026-08-09T14:27:41.065666Z","title":"Hybrid rl: Using both offline and online data can make rl efficient","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.065666Z"},"links":{"cited_paper":"/paper/2210.06718","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:ceb0a9b77f54458a6c49e60345df72009e29a5c032baee1305c1df422fe9d7b1","observation_id":"4ee7bc27-b8b3-44e2-be21-0238b4b01c53","resolution":{"observed_at":"2026-08-09T14:27:41.065666Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.139462Z","title":"Of moments and matching: A game- theoretic framework for closing the imitation gap","venue":null,"work_id":"f562430a-a59c-476f-a3ec-644e071b8f28","year":2021},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.070027Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:260a99ac6d3dfdb15a5af3ae16ab45b21c76adb415094faecec00c6b2edc7a76","observation_id":"94c2e2c4-c045-46fb-ae27-15595efc8776","resolution":{"observed_at":"2026-08-09T14:27:42.144725Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.124409Z","title":"Inverse reinforcement learning without reinforcement learning","venue":null,"work_id":"4b42f7fa-cbd7-4878-9d1a-acf02816a3d0","year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.073969Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:31bcbc34fe8f640ab3867ac1a727075b3097b1c6ac38c1b1949ce85b7a8fb02e","observation_id":"0cbeca50-09bc-41b5-abb5-1905d698c784","resolution":{"observed_at":"2026-08-09T14:27:42.129334Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.077554Z","title":"All roads lead to likelihood: The value of reinforcement learning in fine- tuning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.077554Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:3d13f95568825bd17e0a80e9faddf18a5636b00840e3fb2a4f5e715edf4421e6","observation_id":"e2e98ee7-1f1e-44b6-836b-cf99b9424741","resolution":{"observed_at":"2026-08-09T14:27:41.077554Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.109457Z","title":"Open X-Embodiment: Robotic learning datasets and RT-X models","venue":null,"work_id":"387b91fa-3679-4a52-8db9-97c7c58e2edb","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.081120Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:0eb8a1deef4690c53462e35e4b660f0dfc1911af7451c6b48942a7a7e3b609a8","observation_id":"11ef39d3-0b49-42a4-b3d9-4367073cd1c9","resolution":{"observed_at":"2026-08-09T14:27:42.114493Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.094769Z","title":"The virtues of laziness in model-based rl: A unified objective and algorithms","venue":null,"work_id":"a3464423-b4cb-4c39-bb32-1ae625a0b2c0","year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.084844Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:d48fe186813bbe521efbcfb82d1bdc7f2a6515b36cdca6fd0f358f3d7c9a5f4a","observation_id":"290dbfa2-7fc9-4760-81a0-7a1a4188cc67","resolution":{"observed_at":"2026-08-09T14:27:42.099643Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.435228Z","title":"Vincent, Haruki Nishimura, Masha Itkina, Paarth Shah, Mac Schwager, and Thomas Kollar","venue":null,"work_id":"6882ebb2-6084-41b7-ad62-da238cfa2e01","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.088345Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:9e85399a42655b43beb78c1c25a10fb67804745c6d89e449a636f5130acd928c","observation_id":"3de94b86-33af-427f-be63-d3c9f4388217","resolution":{"observed_at":"2026-08-09T14:27:41.440863Z","resolver_source":"arxiv_id_nonexistent","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.16627","last_updated":"2025-03-26T02:40:00Z","snapshot_observed_at":"2026-07-06T19:56:43.637339Z","submitted_at":"2024-11-25T18:03:50Z","title":"Inference-Time Policy Steering through Human Interactions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.16627","snapshot_observed_at":"2026-08-09T14:27:41.091991Z","title":"Inference- time policy steering through human interactions","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.091991Z"},"links":{"cited_paper":"/paper/2411.16627","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:55ac46461193e4235f82bb3089d8401e49ee7e419b676553238afa37574cabab","observation_id":"15ffe3ec-6a88-4c52-a879-af091929eb3d","resolution":{"observed_at":"2026-08-09T14:27:41.091991Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.079909Z","title":"I can tell what i am doing: Toward real-world natural language grounding of robot experiences","venue":null,"work_id":"31b0e839-c2da-4497-949e-2064555d951d","year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.095730Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:6cc7cd39a272c74aa9ab1c46ed319b2b632c2cc5b2437df7d89c43dab892e640","observation_id":"121590dc-976e-49aa-9b31-4996f59e3dc5","resolution":{"observed_at":"2026-08-09T14:27:42.084987Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.099593Z","title":"ivideogpt: Inter- active videogpts are scalable world models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.099593Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:ef91ee8826e3afe2169410b6f6ee591114ad099b7729962af2506f635924d412","observation_id":"590503f7-1599-4a35-aa29-bb2b1d56f542","resolution":{"observed_at":"2026-08-09T14:27:41.099593Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.056688Z","title":"Daydreamer: World models for physical robot learning","venue":null,"work_id":"8ac5ef59-28f7-43e8-8c8c-384789da0bc2","year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.103763Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:582ddc5d921c806f601cb252ffed0a930f40e78d92f5e884349b09ba6b4e31f4","observation_id":"6c651721-e6bb-439b-9278-ed7ceb05b698","resolution":{"observed_at":"2026-08-09T14:27:42.061410Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.13705","last_updated":"2023-04-23T19:10:53Z","snapshot_observed_at":"2026-08-03T01:22:01.078078Z","submitted_at":"2023-04-23T19:10:53Z","title":"Learning Fine-Grained Bimanual Manipulation with Low-Cost Hardware","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.13705","snapshot_observed_at":"2026-08-09T14:27:41.107736Z","title":"Zhao, Vikash Kumar, Sergey Levine, and Chelsea Finn","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.107736Z"},"links":{"cited_paper":"/paper/2304.13705","citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:777b52d89776b08adf425f4e3753c3cad096ba16afb79f458a8f3ba4fb81b1fc","observation_id":"b6e45317-302e-4b13-8fca-0ab7eaba9a8c","resolution":{"observed_at":"2026-08-09T14:27:41.107736Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.044138Z","title":"Maximum entropy inverse reinforcement learning","venue":null,"work_id":"ef457c34-cb3c-458d-a79f-c1ff98809005","year":2008},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.111624Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:4290addeec94580e0b0a475017f2e36ef91256142909caf30fecd946daf3fec8","observation_id":"6c8988b6-4fe9-4f9a-9c16-0a1e7e019331","resolution":{"observed_at":"2026-08-09T14:27:42.048246Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.030852Z","title":"Rollouts from Demonstration Dataset Cup Task Bag Task Fork Task Fig","venue":null,"work_id":"ae4119a0-c937-4ce5-9376-bbe76ec9ef40","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.115271Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:d014c298f17a3cfd9761a5b26f80c09e2e6ad5d9f6bfbdbe01691e67ef342d1d","observation_id":"4246158a-a597-4fd8-abe4-ef8ae97f2088","resolution":{"observed_at":"2026-08-09T14:27:42.034947Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.017731Z","title":"Each mode has 25 demonstrations","venue":null,"work_id":"caf7666e-2503-496b-8eec-f04c065563ac","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.119280Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:92cea0af42673e6a55138442a11b24b406db2fdb362d0cdfdea1efeedf0c2e13","observation_id":"3e5e1362-f299-4b1d-803a-3c2d98623638","resolution":{"observed_at":"2026-08-09T14:27:42.021710Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:42.004068Z","title":"The effectiveness of world models has been demonstrated across various embodied domains [25, 50]","venue":null,"work_id":"c146a6d4-5681-41c3-81e1-f12ff2735b1f","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.123721Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:5ac1da09cacda346b1f183f9f72c7a0ef1d317bde312e75fdfbe21c9218545d2","observation_id":"834a5348-24e4-47be-a5cc-e02f8ffaccab","resolution":{"observed_at":"2026-08-09T14:27:42.008450Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.989840Z","title":"The robot aims to grasp the cup. Describe the behavior","venue":null,"work_id":"1732d447-56aa-4f94-984a-3572d234678c","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.127519Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:61c2d46d1aa1bb117eee794d468b0dbbdf2c4ae42bcd6f12358fe913d105765e","observation_id":"685454a3-965f-43a9-9232-448851f66816","resolution":{"observed_at":"2026-08-09T14:27:41.995019Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.977721Z","title":"We use Llama-3.2-11B-Vision-Instruct model as our VLM backbone","venue":null,"work_id":"f1a47de5-db27-4eb6-9387-df9c64b46d03","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.132604Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:36d21dab4adafd929348640ad6cdc1527d9f8b8354b6f4abe7e53280eda78c68","observation_id":"d472b462-4da8-4c97-a9f0-fd799ade0acb","resolution":{"observed_at":"2026-08-09T14:27:41.981561Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.964235Z","title":"handle, rim","venue":null,"work_id":"ee1b8639-1484-4882-bbc5-d2304b562779","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.137508Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:caaee62dfe015421dfa6b4900eaf40bfe3d2c67f0ec8e88f6cd0b806d1d488e8","observation_id":"275ea7af-e10b-4a5f-a2f7-a3e44ef48adf","resolution":{"observed_at":"2026-08-09T14:27:41.969456Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.950090Z","title":"The first three dimensions are the x, y, z positions of the robot gripper","venue":null,"work_id":"48a341df-ec72-4a2f-bd38-53e8c1d9aab0","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.141637Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:9b53c805cd234c66ba733002730e2124081727b5f3b1ebfe95e5fa02d2a66262","observation_id":"3503ccaa-341b-4174-99e7-f7b9518f4a5b","resolution":{"observed_at":"2026-08-09T14:27:41.954803Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.935812Z","title":"Do not include any additional text, explanations, or information","venue":null,"work_id":"621749a8-dd44-4a7b-910d-52eb95658ba4","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.145754Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:6080066b92792c14f0e4bfc32d80831a65c2c7f7e5cb9f2bb52fea585256c0fa","observation_id":"8ced2037-f41c-4b47-a20f-0af641754e1e","resolution":{"observed_at":"2026-08-09T14:27:41.940597Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.920969Z","title":"This baseline is an ablated version of FORE- W ARNwithout the explicit world model","venue":null,"work_id":"33390d0b-644b-44c7-a486-b80f10d96722","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.149920Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:504b6341ea4ac4785a1f8a9dcdbbd03e6bde25eb4324fe2b994a949f45d26150","observation_id":"d24188e7-655d-4f36-8500-9d5ecd9db2e0","resolution":{"observed_at":"2026-08-09T14:27:41.925803Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.906615Z","title":null,"venue":null,"work_id":"c8af411b-bc65-441a-98d3-83f12fbec47a","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.154451Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:63c7b20fa3a48413e8e17bd7396f6384590e51245ff394f040708ea88b77b7aa","observation_id":"6166e67c-4580-40e8-b0ab-67439d791d10","resolution":{"observed_at":"2026-08-09T14:27:41.911276Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.892299Z","title":null,"venue":null,"work_id":"85a74252-d25c-479b-aa40-248e075e29af","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.158740Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:da8c1b36d91c529867d71b363e534471b485aff0a7734ff0024d66c5c13173c0","observation_id":"f8802b02-3b33-4eeb-b3e9-86dc5aebab7e","resolution":{"observed_at":"2026-08-09T14:27:41.897180Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.876176Z","title":"handle, inner surface, etc","venue":null,"work_id":"56d3babd-fb8b-4d72-838c-afd2e53694ec","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.162836Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:52ab8cdc268a39b399d6a32b7eab94772a0e781fdec53713788e02bfad176ba3","observation_id":"97a04988-ccc1-4909-bf11-400e36f3b93b","resolution":{"observed_at":"2026-08-09T14:27:41.881476Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.860432Z","title":"If the cup is not grasped in the robot's gripper, the sentence should describe the failure","venue":null,"work_id":"cdc3963c-2a9a-4e18-a999-e096f08e5d76","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.166912Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:634180fa632f6c31bdc9495cd3f7d0e23b7c1068e1ef3b9d094f78c43da30371","observation_id":"9b0645c8-aea2-4033-ab4c-1326999f8c46","resolution":{"observed_at":"2026-08-09T14:27:41.865428Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.846404Z","title":"15 demonstrates the setup of our real-world experiments","venue":null,"work_id":"08c22b59-4c5c-497b-b886-d4bc45d3a1ab","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.171561Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:1f92dd94d4bc2644951d1a2639de3302e615c85f54846eb52bb45943d16e1b1d","observation_id":"48f0af90-85f3-480e-a7c5-60cb0ec5776c","resolution":{"observed_at":"2026-08-09T14:27:41.850605Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.832841Z","title":"We include additional qualitative examples for Cup and Bag tasks in Fig 16","venue":null,"work_id":"79f692a8-8489-4e6c-a8c4-904ea8763aac","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.175887Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:dc767cd89d2c2a154b217ac726edaf33278e5b984ccb59b87f4131acc36b3777","observation_id":"b1cfd582-4974-45d4-a5a4-21d5acda715e","resolution":{"observed_at":"2026-08-09T14:27:41.837085Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.819273Z","title":"V-A and queries the VLM again to decide if the behavior is a success or failure within the context of the task description ℓ","venue":null,"work_id":"46e08947-55b3-4baa-8180-a264253341f5","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.180649Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:722d2890a7d834ac23158ddf9c8929d4d46eaf4895aba47f5323cbb816270be8","observation_id":"d6e2a521-79e8-4b0a-85b6-23c558f2e3fa","resolution":{"observed_at":"2026-08-09T14:27:41.823502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T14:27:41.805967Z","title":"grasp- ing the cup","venue":null,"work_id":"6f715447-9154-4a89-926d-1a6df3410da1","year":null},"citing_paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment","version":3},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-09T14:27:41.185161Z"},"links":{"citing_paper":"/paper/2502.01828"},"observation_digest":"sha256:54e5db5efb42e9f70f80ea35f58a157f136a5fd7cc554f77e2764a24909dc450","observation_id":"1e289fe1-11f1-4f50-8886-14a0e0b6cb36","resolution":{"observed_at":"2026-08-09T14:27:41.809872Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.01828","last_updated":"2025-05-02T17:53:34Z","latest_version":3,"primary_category":"cs.RO","snapshot_observed_at":"2026-08-09T14:17:11.433619Z","submitted_at":"2025-02-03T21:11:02Z","title":"From Foresight to Forethought: VLM-In-the-Loop Policy Steering via Latent Alignment"},"reference_resolution":{"displayed":69,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":22,"verified_exact":2,"verified_fuzzy":45},"total_outbound_references":69},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 69 of 69 outbound references and 15 inbound Pith citation observations for arXiv:2502.01828."}