{"as_of":"2026-08-10T01:41:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:9d74c98e0409b7de0ef9d8803c2d0a9e5d5ad9056aa76053afbf6d1437571508","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":29,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":29,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":29,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":29,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T22:44:50.810323Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T07:26:54.507995Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2401.00870","last_updated":"2026-04-08T14:54:42Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-30T01:26:42Z","title":"ConfusionPrompt: Practical Private Inference for Online Large Language Models","version":5},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-24T04:57:55.197897Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2401.00870"},"observation_digest":"sha256:97c83ff58a53a9e8366e8645f5f9f4cc3cc1de3c5d203b43a154d741bcdb6dd5","observation_id":"14ef7f5e-d878-4474-bbb2-24627362f663","resolution":{"observed_at":"2026-05-24T04:58:54.849006Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2405.14093","last_updated":"2026-05-01T01:50:44Z","snapshot_observed_at":"2026-08-04T06:47:25.827167Z","submitted_at":"2024-05-23T01:43:54Z","title":"A Survey on Vision-Language-Action Models for Embodied AI","version":8},"reference_index":107,"source":"pdf_text","source_observed_at":"2026-05-24T01:25:10.150459Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2405.14093"},"observation_digest":"sha256:f31eaba90d26119c8694c9b78ebf9d8507198c2623499ec47867e20fcebe4bef","observation_id":"28ec5370-4b3d-4212-9945-02add761aeb4","resolution":{"observed_at":"2026-05-24T01:25:54.588575Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-09T22:44:50.810323Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large lan- guage model,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.18733","last_updated":"2025-01-30T20:19:01Z","snapshot_observed_at":"2026-08-09T22:39:04.137566Z","submitted_at":"2025-01-30T20:19:01Z","title":"Integrating LMM Planners and 3D Skill Policies for Generalizable Manipulation","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-09T22:44:50.810323Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2501.18733"},"observation_digest":"sha256:a2350ded2c2e43ba611938b95a74a8c369f7658754a34f09d5d8ba681c170045","observation_id":"89f0551f-a3db-4646-a5b3-761676c6817c","resolution":{"observed_at":"2026-08-09T22:44:50.810323Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-08T00:03:16.363277Z","title":"In- struct2act: Mapping multi-modality instructions to robotic actions with large language model,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.08643","last_updated":"2025-02-18T16:45:59Z","snapshot_observed_at":"2026-08-09T11:31:44.840528Z","submitted_at":"2025-02-12T18:57:22Z","title":"A Real-to-Sim-to-Real Approach to Robotic Manipulation with VLM-Generated Iterative Keypoint Rewards","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-08T00:03:16.363277Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2502.08643"},"observation_digest":"sha256:61866ddd86116579f1115281638b1251ef4e1e4e4fef606ca1a7c1267c14a688","observation_id":"4373253c-8b85-4c23-9d85-3f86929dab2c","resolution":{"observed_at":"2026-08-08T00:03:16.363277Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-07T23:18:39.880512Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.08922","last_updated":"2025-02-13T03:15:31Z","snapshot_observed_at":"2026-08-09T00:20:18.332726Z","submitted_at":"2025-02-13T03:15:31Z","title":"Self-Consistency of the Internal Reward Models Improves Self-Rewarding Language Models","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T23:18:39.880512Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2502.08922"},"observation_digest":"sha256:d63f0a9aa4941e8d40bab699199e66ed6500341d78a373367ae44b4fc48167f7","observation_id":"89abcfe5-4980-4830-a70f-a54d6375978a","resolution":{"observed_at":"2026-08-07T23:18:39.880512Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-07T06:03:23.742554Z","title":"Huang, Z","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.06066","last_updated":"2025-06-06T13:20:30Z","snapshot_observed_at":"2026-08-09T17:37:17.945972Z","submitted_at":"2025-06-06T13:20:30Z","title":"Conversational Interfaces for Parametric Conceptual Architectural Design: Integrating Mixed Reality with LLM-driven Interaction","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T06:03:23.742554Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2506.06066"},"observation_digest":"sha256:08b33d96d6030822eb5936e4f376b3d9d96a89a54d00fc72512a989ace9e66c3","observation_id":"8ca5bd2e-b098-457d-9d5d-a490089a89d2","resolution":{"observed_at":"2026-08-07T06:03:23.742554Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-06T22:34:06.759535Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.21250","last_updated":"2025-06-26T13:35:53Z","snapshot_observed_at":"2026-08-09T18:24:37.934285Z","submitted_at":"2025-06-26T13:35:53Z","title":"ACTLLM: Action Consistency Tuned Large Language Model","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T22:34:06.759535Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2506.21250"},"observation_digest":"sha256:4644d4a2579bf6090d32e3d601c8f63f52fb2f9eb8d7ce084942c10ca2e8e8e8","observation_id":"c0190b68-6726-4d51-925c-eb772a4a5cb2","resolution":{"observed_at":"2026-08-06T22:34:06.759535Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2507.01925","last_updated":"2025-07-02T17:34:52Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-02T17:34:52Z","title":"A Survey on Vision-Language-Action Models: An Action Tokenization Perspective","version":1},"reference_index":129,"source":"pdf_text","source_observed_at":"2026-05-17T14:08:34.893876Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2507.01925"},"observation_digest":"sha256:de9bc5a948106c203a5f067c71c9dbda2991de9c71e39e82af73f3ca8924d9ad","observation_id":"41837580-5943-489d-86cf-d0ebd4a8e2da","resolution":{"observed_at":"2026-05-17T14:08:35.437833Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-06T17:43:53.122220Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.10087","last_updated":"2025-07-14T09:13:07Z","snapshot_observed_at":"2026-08-09T16:03:14.420559Z","submitted_at":"2025-07-14T09:13:07Z","title":"Foundation Model Driven Robotics: A Comprehensive Review","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T17:43:53.122220Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2507.10087"},"observation_digest":"sha256:aa866999149d8a707f1b8b0895ae43df973d77ec211e4d10a069d84ed26df33d","observation_id":"f149ce8b-d424-4047-8791-14a0fd8573fe","resolution":{"observed_at":"2026-08-06T17:43:53.122220Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2508.05635","last_updated":"2025-11-04T12:01:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-08-07T17:59:44Z","title":"Genie Envisioner: A Unified World Foundation Platform for Robotic Manipulation","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-15T21:28:41.904725Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2508.05635"},"observation_digest":"sha256:a1923c3e83ef09fac02ebf522c356007abc5f2c19ada9113fb760c4292127c01","observation_id":"fc477058-e9f4-4f03-aa57-ac6113c7ec82","resolution":{"observed_at":"2026-05-15T21:28:42.012107Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-05T23:17:30.341166Z","title":"Huang, Z","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.05636","last_updated":"2025-08-07T17:59:59Z","snapshot_observed_at":"2026-08-07T21:51:21.989392Z","submitted_at":"2025-08-07T17:59:59Z","title":"FaceAnonyMixer: Cancelable Faces via Identity Consistent Latent Space Mixing","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-05T23:17:30.341166Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2508.05636"},"observation_digest":"sha256:06f563ebeaae2903c7b5dea17069900fbee647c680f313b6ce1d5b7cfe9e584e","observation_id":"ba2340f1-f0a0-4a76-8b2a-df76d4981e15","resolution":{"observed_at":"2026-08-05T23:17:30.341166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-05T20:31:42.886164Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.10399","last_updated":"2025-08-14T06:56:16Z","snapshot_observed_at":"2026-08-09T21:25:06.059787Z","submitted_at":"2025-08-14T06:56:16Z","title":"Large Model Empowered Embodied AI: A Survey on Decision-Making and Embodied Learning","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-05T20:31:42.886164Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2508.10399"},"observation_digest":"sha256:6e2abb5bf05e1818dbbcf6eeeebe5ca28d4568684220681e656d69f3b55322e9","observation_id":"50b1e20c-be23-4d8f-95b1-8a6def55c66b","resolution":{"observed_at":"2026-08-05T20:31:42.886164Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2508.13073","last_updated":"2025-09-01T08:10:01Z","snapshot_observed_at":"2026-08-07T15:40:31.068428Z","submitted_at":"2025-08-18T16:45:48Z","title":"Large VLM-based Vision-Language-Action Models for Robotic Manipulation: A Survey","version":2},"reference_index":159,"source":"pdf_text","source_observed_at":"2026-05-17T20:28:15.818016Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2508.13073"},"observation_digest":"sha256:9b6b08620c2d7541964017e6f913bf26a2d2a07dc6da6e9e4929dc1bcf20df1a","observation_id":"439c4702-a860-4347-a2b3-4074bbf90a83","resolution":{"observed_at":"2026-05-17T20:28:16.007179Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2511.01594","last_updated":"2026-04-07T16:40:59Z","snapshot_observed_at":"2026-07-06T22:34:48.431368Z","submitted_at":"2025-11-03T13:58:37Z","title":"MARS: Multi-Agent Robotic System with Multimodal Large Language Models for Assistive Intelligence","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-18T01:05:47.657235Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2511.01594"},"observation_digest":"sha256:a41af057ca2493bb80c208c21b5e5cf235b4d0e10a645ee20502d711ca59c3ae","observation_id":"ed0b6296-7bb2-4a49-8dcb-3f090240dd80","resolution":{"observed_at":"2026-05-18T01:10:34.249058Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2512.17435","last_updated":"2026-04-30T09:47:33Z","snapshot_observed_at":"2026-08-02T23:32:24.828654Z","submitted_at":"2025-12-19T10:40:16Z","title":"ImagineNav++: Prompting Vision-Language Models as Embodied Navigator through Scene Imagination","version":3},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-05-16T21:07:54.497999Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2512.17435"},"observation_digest":"sha256:e9cf57b73745c8329a2491b047249b0fba428c044f518cc012ef990e064f15b4","observation_id":"ff20d6fc-18c7-4d48-bbdb-df59b914bd95","resolution":{"observed_at":"2026-05-16T21:08:32.486743Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-03T08:15:20.295381Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.17717","last_updated":"2026-06-09T20:25:14Z","snapshot_observed_at":"2026-08-08T23:34:38.789355Z","submitted_at":"2026-01-25T06:40:25Z","title":"A Survey on Evaluating Quality and Trustworthiness in LLM-Generated Data","version":3},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-08-03T08:15:20.295381Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2601.17717"},"observation_digest":"sha256:ab38a8f35dd2a88de1ea91aca70bb8d8e9cc0e8583736bfb51453f6283e95b30","observation_id":"288dafc4-074a-4583-abe1-aaee24236c04","resolution":{"observed_at":"2026-08-03T08:15:20.295381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2604.13533","last_updated":"2026-04-22T06:55:18Z","snapshot_observed_at":"2026-07-31T16:36:40.643204Z","submitted_at":"2026-04-15T06:29:02Z","title":"Evolvable Embodied Agent for Robotic Manipulation via Long Short-Term Reflection and Optimization","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-10T13:04:48.932945Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2604.13533"},"observation_digest":"sha256:134ccfe2116fd7851c0e351fd5d6ddb7e9865913f6b2d1ef17f36b7d40983c1f","observation_id":"995f3c54-2c1d-45fa-a8e2-970b786bc6ab","resolution":{"observed_at":"2026-05-10T13:05:24.469130Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2606.05004","last_updated":"2026-06-03T15:23:06Z","snapshot_observed_at":"2026-08-05T17:12:55.961654Z","submitted_at":"2026-06-03T15:23:06Z","title":"SharedRequest: Privacy-Preserving Model-Agnostic Inference for Large Language Models","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-06-28T05:29:47.894226Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2606.05004"},"observation_digest":"sha256:a7e752fe3518a2f5a874df78a6944b6bbdc4b3fc123da0c82d313dac0da94461","observation_id":"be0a79a6-6112-47a0-aa88-ecbeb602750f","resolution":{"observed_at":"2026-07-02T09:26:51.544337Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2606.07999","last_updated":"2026-06-06T06:33:51Z","snapshot_observed_at":"2026-08-01T19:10:51.691151Z","submitted_at":"2026-06-06T06:33:51Z","title":"Efficient Skill Grounding via Code Refactoring with Small Language Models","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-06-27T19:55:20.212198Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2606.07999"},"observation_digest":"sha256:110026a100b69411884798444051cc6e78f058c8ce5dbbd28b49d4d59b22491f","observation_id":"9809110e-0ecf-49d2-a9b3-76dfeb8145ba","resolution":{"observed_at":"2026-07-02T21:07:24.153749Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2606.13097","last_updated":"2026-07-29T12:55:33Z","snapshot_observed_at":"2026-08-06T14:19:46.913836Z","submitted_at":"2026-06-11T09:25:27Z","title":"Functional Cache Grafting: Robust and Rapid Code-Policy Synthesis for Embodied Agents","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-27T05:26:48.431739Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2606.13097"},"observation_digest":"sha256:b03b7b69ec5781ef9df9f98f392af091b54a16b5f799a7b75c7fc99fa1c6db38","observation_id":"0e5fb998-6181-44f1-8c8e-c8d955e8c208","resolution":{"observed_at":"2026-06-27T05:30:35.769413Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-02T11:44:03.545449Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.13097","last_updated":"2026-07-29T12:55:33Z","snapshot_observed_at":"2026-08-06T14:19:46.913836Z","submitted_at":"2026-06-11T09:25:27Z","title":"Functional Cache Grafting: Robust and Rapid Code-Policy Synthesis for Embodied Agents","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-02T11:44:03.545449Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2606.13097"},"observation_digest":"sha256:4c82f513253af3c8cf5b4b1d025df082e7fdf30b9a9de938bce4e5174a6f2ae8","observation_id":"34c3500e-747b-4446-a4f6-566ad15a71d2","resolution":{"observed_at":"2026-08-02T11:44:03.545449Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2606.17030","last_updated":"2026-06-17T13:54:57Z","snapshot_observed_at":"2026-08-01T21:48:31.832288Z","submitted_at":"2026-06-15T17:52:31Z","title":"Qwen-RobotWorld Technical Report: Unifying Embodied World Modeling through Language-Conditioned Video Generation","version":3},"reference_index":230,"source":"arxiv_source","source_observed_at":"2026-06-27T04:19:26.332718Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2606.17030"},"observation_digest":"sha256:48ab64c187374df61b0e1bdb277e52de86b7f278eb13d9940a17689979e3a3c6","observation_id":"f29e2a24-7b6b-4337-81ef-652bb752d701","resolution":{"observed_at":"2026-07-03T17:18:43.761394Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":"2305.11176","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-10T07:26:54.507995Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":"cs.RO","work_id":"ec87e5e6-72f8-42dd-af75-48d4cdb39f43","year":2023},"citing_paper":{"arxiv_id":"2607.08448","last_updated":"2026-07-15T16:40:12Z","snapshot_observed_at":"2026-08-04T18:37:38.875340Z","submitted_at":"2026-07-09T13:08:54Z","title":"Harness VLA: Steering Frozen VLAs into Reliable Manipulation Primitives via Memory-Guided Agents","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-07-10T07:18:57.823444Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2607.08448"},"observation_digest":"sha256:de19af74a8d3eb6884d4abf482ecc090840307347aa9bdc6bb860574628a4711","observation_id":"1645445c-a951-4492-a399-749d9d4b7ff7","resolution":{"observed_at":"2026-07-10T07:26:54.509395Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-02T07:56:56.620153Z","title":"Huang, Z","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.08448","last_updated":"2026-07-15T16:40:12Z","snapshot_observed_at":"2026-08-04T18:37:38.875340Z","submitted_at":"2026-07-09T13:08:54Z","title":"Harness VLA: Steering Frozen VLAs into Reliable Manipulation Primitives via Memory-Guided Agents","version":3},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-02T07:56:56.620153Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2607.08448"},"observation_digest":"sha256:3fe1838dee1aee4fcbbe8c466086cbbe3c5893c5d534b0a54f245ab9857859dd","observation_id":"a68f4efd-fefc-40af-aaf3-647509499690","resolution":{"observed_at":"2026-08-02T07:56:56.620153Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-07-14T03:31:19.309532Z","title":"arXiv:2305.11176 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11738","last_updated":"2026-07-13T16:00:03Z","snapshot_observed_at":"2026-08-06T12:36:13.363835Z","submitted_at":"2026-07-13T16:00:03Z","title":"Qwen-Audio-VAE Technical Report","version":1},"reference_index":237,"source":"arxiv_source","source_observed_at":"2026-07-14T03:31:19.309532Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2607.11738"},"observation_digest":"sha256:d3cf55314bb178eb77ec00a3b1e95f1bf06a20fc2e24782882117008fa677388","observation_id":"c237d157-fea3-4b06-8ba0-d60383f9904d","resolution":{"observed_at":"2026-07-14T03:31:19.309532Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-02T09:54:52.432986Z","title":"arXiv preprint arXiv:2305.11176 (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.16247","last_updated":"2026-06-26T14:26:24Z","snapshot_observed_at":"2026-08-08T01:50:47.251899Z","submitted_at":"2026-06-26T14:26:24Z","title":"Self-Evolving Just-In-Time Memory for Proactive Embodied Safety","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-02T09:54:52.432986Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2607.16247"},"observation_digest":"sha256:2ea415e5e27716a2a0c23d6d16c9ac5a679bcf3f18bf2a3c2ecb7a3982894b00","observation_id":"4c58c8a4-f760-43e9-8d79-5fa90acfec7c","resolution":{"observed_at":"2026-08-02T09:54:52.432986Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-04T19:45:33.996455Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.01851","last_updated":"2026-08-03T07:58:35Z","snapshot_observed_at":"2026-08-07T23:22:57.641630Z","submitted_at":"2026-08-03T07:58:35Z","title":"Weights or Skills? A Survey of Robot-Learning Techniques: from Action-Predicting Weights to Robots that Write their Own Skills","version":1},"reference_index":97,"source":"pdf_text","source_observed_at":"2026-08-04T19:45:33.996455Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2608.01851"},"observation_digest":"sha256:6aeb57a3530d0b6b1acee96544a058089062135297fd44897e27e5ebb1740f58","observation_id":"e3d210c8-b04d-4049-83ea-fe0b50d20bb5","resolution":{"observed_at":"2026-08-04T19:45:33.996455Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-05T15:25:40.330321Z","title":"doi:10.48550/arXiv.2305.11176 , abstract =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.03644","last_updated":"2026-08-04T13:29:00Z","snapshot_observed_at":"2026-08-08T18:58:23.959751Z","submitted_at":"2026-08-04T13:29:00Z","title":"Is Inter-Seed Cross-Play Enough? Evaluating the Robustness of Zero-Shot Coordination Algorithms to Implementation Details","version":1},"reference_index":220,"source":"arxiv_source","source_observed_at":"2026-08-05T15:25:40.330321Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2608.03644"},"observation_digest":"sha256:93c795bf851580adda8d6fd3e025e0ee278d0b46dc39db25fe9a91c7879df561","observation_id":"766e97bd-0a28-4853-bb37-341cae407934","resolution":{"observed_at":"2026-08-05T15:25:40.330321Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11176","snapshot_observed_at":"2026-08-05T05:44:44.230086Z","title":"Instruct2act: Mapping multi-modality instructions to robotic actions with large language model","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.03924","last_updated":"2026-08-04T16:56:09Z","snapshot_observed_at":"2026-08-08T09:16:47.522854Z","submitted_at":"2026-08-04T16:56:09Z","title":"ETA: A New Agentic Paradigm for Embodied Tasks","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-05T05:44:44.230086Z"},"links":{"cited_paper":"/paper/2305.11176","citing_paper":"/paper/2608.03924"},"observation_digest":"sha256:24225e3425a7c41afa374a9fd067aad7654daeb9917c15d849b6a15ad801af39","observation_id":"54530e77-b581-4687-a48e-ad8354d3d278","resolution":{"observed_at":"2026-08-05T05:44:44.230086Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2305.11176/citation-record","integrity":"/paper/2305.11176/integrity","json":"/paper/2305.11176/citation-record.json","paper":"/paper/2305.11176"},"outbound":[],"paper":{"arxiv_id":"2305.11176","last_updated":"2023-05-24T04:17:34Z","latest_version":3,"primary_category":"cs.RO","snapshot_observed_at":"2026-07-06T15:29:20.405296Z","submitted_at":"2023-05-18T17:59:49Z","title":"Instruct2Act: Mapping Multi-modality Instructions to Robotic Actions with Large Language Model"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 29 inbound Pith citation observations for arXiv:2305.11176."}