{"as_of":"2026-08-13T02:33:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:55323a6ee97072dc97ea9b598dddfe2f09e85b0080336133cd10edd5caf2c0a9","coverage":[{"denominator":39,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":39,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T21:18:52.322374Z","state":"measured"},{"denominator":39,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":39,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2509.08126/citation-record","integrity":"/paper/2509.08126/integrity","json":"/paper/2509.08126/citation-record.json","paper":"/paper/2509.08126"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:55.990524Z","title":"Adversarial object rearrangement in constrained environments with heterogeneous graph neural networks,","venue":null,"work_id":"43a98141-8c12-499c-89e9-7984bf5697be","year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:49.826222Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:a04b0e3b64c7f5a6a9a851f715e9d03e23f93d312c5d11ef8a2123014d076f5a","observation_id":"0864f882-6c82-4e02-85c6-9c3f686b90b2","resolution":{"observed_at":"2026-08-04T21:18:56.058531Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:55.836997Z","title":"Self-supervised interactive object segmentation through a singulation-and-grasping approach,","venue":null,"work_id":"d3f0a0e6-6864-4c68-a391-0b13a3d4128e","year":2022},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:49.874819Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:929bbb0eb6f5adc403902e52b9e1c7d2ece66200a991bea8238c66758f6cd013","observation_id":"0af25705-9f85-4a6c-b07a-7f65b5cff7d6","resolution":{"observed_at":"2026-08-04T21:18:55.900683Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:55.645705Z","title":"Iosg: Image-driven object searching and grasping,","venue":null,"work_id":"048b133b-789b-4a96-8ec3-2a0369b195ec","year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:49.954341Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:571499b2fd0e1368331f905371995bf894c8234bc1ac4ad4d26950f6eef5c170","observation_id":"9e8a4860-f8fd-4078-9e37-5154bf936593","resolution":{"observed_at":"2026-08-04T21:18:55.712776Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:55.459906Z","title":"Antipodal robotic grasping using generative residual convolutional neural network,","venue":null,"work_id":"507ec5ed-c3e4-4932-997f-85d8cbb6aaa6","year":2020},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.023541Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:97ebf7d3c906821f111783f58df77b080ed37065de6b12b17e0ebfaf04e796d1","observation_id":"4110869f-f6c2-427a-8de8-3443b9d7ff44","resolution":{"observed_at":"2026-08-04T21:18:55.543245Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:55.310932Z","title":"Language-driven grasp detection,","venue":null,"work_id":"360a2319-caa9-4614-af50-f3486fc91812","year":2024},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.084659Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:95c5c62d8bdf8c76c90bbee5ab9e77974452ed7e0b5c046b384ec3547d4d6878","observation_id":"e7a3b780-0f3f-4a02-8f33-246c7f0e4254","resolution":{"observed_at":"2026-08-04T21:18:55.390359Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:55.160992Z","title":"Attribute-based robotic grasping with data-efficient adaptation,","venue":null,"work_id":"fff5ea83-949d-4073-9aad-17e4d6f62ded","year":2024},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.146980Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:5d593767f19cc0f364cbc48175f93c9949faf2a3db6b01f5fd00c4dd807b66bc","observation_id":"482a48c9-0bf1-487b-9916-6253d08ed256","resolution":{"observed_at":"2026-08-04T21:18:55.204165Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:55.008117Z","title":"A joint modeling of vision-language-action for target- oriented grasping in clutter,","venue":null,"work_id":"89441967-0384-4f44-9fe6-4b5487e66382","year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.239490Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:2de794255714a91df4f182c8ac6612f778c9bcaba703b1d94a50cb6bfad0d84e","observation_id":"b66829ee-45ca-426b-8b4f-f276b69a32d8","resolution":{"observed_at":"2026-08-04T21:18:55.061406Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:54.844740Z","title":"Visually grounding language instruction for history-dependent manipulation,","venue":null,"work_id":"ab4eb2e6-5a32-46ac-a2cd-9d428c9de4ea","year":2022},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.296436Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:3bf5c3650181b4d86b6d756a093aa503eb5911c4afa73e461452ea5adc0bd717","observation_id":"0efbf667-4d6c-4a6b-8bd6-37b00db5e25a","resolution":{"observed_at":"2026-08-04T21:18:54.903284Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.06798","last_updated":"2024-10-15T05:15:38Z","snapshot_observed_at":"2026-08-12T07:05:31.275511Z","submitted_at":"2024-02-09T21:48:19Z","title":"Reasoning Grasping via Multimodal Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.06798","snapshot_observed_at":"2026-08-04T21:18:50.362985Z","title":"Reasoning grasping via multi- modal large language model,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.362985Z"},"links":{"cited_paper":"/paper/2402.06798","citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:acfb58f11f9181c1f7e0e55c52469069955da2e3842aa5b6c4dee97becb30ef2","observation_id":"26f4e140-14d6-4a00-a112-57117dff0897","resolution":{"observed_at":"2026-08-04T21:18:50.362985Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.09246","last_updated":"2024-09-05T19:46:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-13T15:46:55Z","title":"OpenVLA: An Open-Source Vision-Language-Action Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.09246","snapshot_observed_at":"2026-08-04T21:18:50.418258Z","title":"Open- vla: An open-source vision-language-action model,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.418258Z"},"links":{"cited_paper":"/paper/2406.09246","citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:f1b8f31826bb4c98d06de6738e41abb9acaf66c283cca828ded2b9f6b2fc43e8","observation_id":"aec0dd06-a3f9-47a0-af39-5fed4c227fa8","resolution":{"observed_at":"2026-08-04T21:18:50.418258Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.12286","last_updated":"2025-05-01T09:13:09Z","snapshot_observed_at":"2026-08-12T23:23:50.199119Z","submitted_at":"2024-11-19T07:12:48Z","title":"GLOVER: Generalizable Open-Vocabulary Affordance Reasoning for Task-Oriented Grasping","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.12286","snapshot_observed_at":"2026-08-04T21:18:50.490457Z","title":"Glover: Generaliz- able open-vocabulary affordance reasoning for task-oriented grasping,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.490457Z"},"links":{"cited_paper":"/paper/2411.12286","citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:113b4b0fb05b9dff98965b3eb43faac51699a9bee55aa95b7b1c1904f33a40e6","observation_id":"73e154f8-05f8-45f2-92f9-a5ddb6d9cfc3","resolution":{"observed_at":"2026-08-04T21:18:50.490457Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.08693","last_updated":"2025-03-06T19:29:03Z","snapshot_observed_at":"2026-08-07T02:44:43.738657Z","submitted_at":"2024-07-11T17:31:01Z","title":"Robotic Control via Embodied Chain-of-Thought Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.08693","snapshot_observed_at":"2026-08-04T21:18:50.546350Z","title":"Robotic control via embodied chain-of-thought reasoning,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.546350Z"},"links":{"cited_paper":"/paper/2407.08693","citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:83e246fd85c8265aaf53894b012b09ebb2a9f06198263d52527322386c8b2a86","observation_id":"c0b3509a-8df7-49f1-86f4-38d5bd8c90f0","resolution":{"observed_at":"2026-08-04T21:18:50.546350Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:54.728315Z","title":"Towards open-world grasping with large vision-language models,","venue":null,"work_id":"5eada7b9-f917-4080-a9ab-75e4811b191c","year":2024},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.638721Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:15d5399858a3bedc101e0f317a9742ccb42ffd246ab3125732f4f953b6c2e3b4","observation_id":"060ade24-d912-4887-b42a-d99fed50650f","resolution":{"observed_at":"2026-08-04T21:18:54.754007Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:50.715515Z","title":"Thinkgrasp: A vision-language system for strategic part grasping in clutter,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.715515Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:43012482a3d836016731c94035931c3d7803e531ac5a1230f64bef98c2736cca","observation_id":"6e0b0447-23c5-47b9-a87c-52ae062bf551","resolution":{"observed_at":"2026-08-04T21:18:50.715515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:54.600070Z","title":"Learning transferable visual models from natural language supervision,","venue":null,"work_id":"e0590718-dc39-45cb-9211-0caef47ee9f8","year":2021},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.761077Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:eab6c0146c8cf9dc2838b21f1153f93d94aa587d94499182c4546317df11fba4","observation_id":"e5d2e0dd-de07-4140-ac6b-369b2a74f175","resolution":{"observed_at":"2026-08-04T21:18:54.640338Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:50.829390Z","title":"Blip: Bootstrapping language- image pre-training for unified vision-language understanding and generation,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.829390Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:9ee50e7ed1e1d8121b50a2d12eb590818155a858ee919e2f87b53ddc973319d6","observation_id":"88ddba33-6586-46f8-a905-1396b1b23b1a","resolution":{"observed_at":"2026-08-04T21:18:50.829390Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.19457","last_updated":"2025-02-07T05:10:48Z","snapshot_observed_at":"2026-08-12T22:35:24.525972Z","submitted_at":"2024-09-28T21:11:25Z","title":"A Parameter-Efficient Tuning Framework for Language-guided Object Grounding and Robot Grasping","version":4},"cited_work":{"arxiv_id":"2409.19457","doi":null,"metadata_source":"pith","pith_arxiv_id":"2409.19457","snapshot_observed_at":"2026-08-04T21:18:52.418096Z","title":"A Parameter-Efficient Tuning Framework for Language-guided Object Grounding and Robot Grasping","venue":"cs.RO","work_id":"e403c2d9-9bef-4928-991c-fb4f15c8863f","year":2024},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.911494Z"},"links":{"cited_paper":"/paper/2409.19457","citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:7963c194d335e46128e5223b6c66f3eb70f1599387c83a1493803243872b0580","observation_id":"69a00805-c536-48d7-8f6f-fd8f79ce11bb","resolution":{"observed_at":"2026-08-04T21:18:52.451521Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:54.472307Z","title":"Language-guided robot grasping: Clip-based referring grasp synthesis in clutter,","venue":null,"work_id":"60dc74cc-e43a-4045-b436-f35e71292d53","year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:50.987346Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:e40a362f2a83254bc9a3a8402321b00ccd8780f23c2b67d02d9e034a38218912","observation_id":"6120b516-c6cd-4c7f-9b0c-65474dcc6663","resolution":{"observed_at":"2026-08-04T21:18:54.531099Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:54.309548Z","title":"Grounding language with visual affordances over unstructured data,","venue":null,"work_id":"7f02e672-8f72-43e7-af1b-ba0e05daa8e0","year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.043897Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:005b9d983fa675584a6278d4dc63e00757ddd5e7a98e5641b0e2de821e1b6534","observation_id":"bd47a1fe-e4b2-49ca-9b85-43d25fc8374d","resolution":{"observed_at":"2026-08-04T21:18:54.395612Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:54.141808Z","title":"Segmentation from natural language expressions,","venue":null,"work_id":"cf5c6b34-b017-4a04-9648-b3b44cb471b9","year":2016},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.075275Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:70058acaf75f04266b7df6fdc704629d84e471f8dc4b39e44841b3e1d62b82f9","observation_id":"1a0afcf0-f84a-4603-b7d4-20855b322f1b","resolution":{"observed_at":"2026-08-04T21:18:54.228833Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:54.038822Z","title":"Referring image segmentation via recurrent refinement networks,","venue":null,"work_id":"14533efb-9528-4a59-851f-60955299dd3d","year":2018},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.140132Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:f3bbc9579ec1c2b94aad972936329a5d8bd7a90f10dfcdeb716dfd9f69884401","observation_id":"a306f667-fbcd-4660-9f74-fc6f38c3a9d0","resolution":{"observed_at":"2026-08-04T21:18:54.091892Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:53.928479Z","title":"Recurrent multimodal interaction for referring image segmentation,","venue":null,"work_id":"96616739-d5ad-4008-b989-02b4eb80c701","year":2017},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.198043Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:86388a9632f93aba91db7d6d8b65fd909188f847771fe3ff3ffa0ad5f97de4ef","observation_id":"544df011-427d-48fd-ab38-e5430b06b6f2","resolution":{"observed_at":"2026-08-04T21:18:53.990084Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:53.756979Z","title":"Bridging vision and language encoders: Parameter-efficient tuning for referring image segmentation,","venue":null,"work_id":"b1992a1c-8524-4b5a-a371-4f289927970f","year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.275247Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:7cd329a6974d4fd44b51f606b7132fd69a8180d393edae8c38f9969a2ac16913","observation_id":"345dbc6b-4faa-4e41-b980-ee08f4d48b6d","resolution":{"observed_at":"2026-08-04T21:18:53.822649Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:53.593098Z","title":"Barleria: An efficient tuning framework for referring image segmentation,","venue":null,"work_id":"07e77199-80db-4f51-83b0-016172d662c4","year":null},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.341023Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:cdad84dc2f848a6290186e86d96e4907f9ada9fc07b9d90c2666635b779d67ed","observation_id":"e0481d41-f2ce-4c68-b850-48480735f832","resolution":{"observed_at":"2026-08-04T21:18:53.656554Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:51.400655Z","title":"Cris: Clip-driven referring image segmentation,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.400655Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:d6298dde3424c9ae32b626717a033381b73a0d594d045f23aff90c093572cd68","observation_id":"70867d91-0ff0-4201-92bf-ba9dc69a95d1","resolution":{"observed_at":"2026-08-04T21:18:51.400655Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:53.412351Z","title":"Lavt: Language-aware vision transformer for referring image segmentation,","venue":null,"work_id":"b1d20e53-cf2e-4976-9fa3-4e5abbc6d911","year":2022},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.473575Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:1d790f2c761f54f5c86cb5afd209fd8b50fb25f208703dd70635b047dc0e0061","observation_id":"ecdc2c22-6814-436b-a423-0473c7cb7a95","resolution":{"observed_at":"2026-08-04T21:18:53.472820Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:53.330095Z","title":"Contrastive grouping with transformer for referring image segmentation,","venue":null,"work_id":"043311a4-6862-4ba4-bc7c-e20397c138a2","year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.543082Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:88003daa01eb913ac873b641eb1e35f5244e2474a060a74878a3cbde858f4459","observation_id":"cfba2855-794d-4247-b40b-7b162ee1b2ac","resolution":{"observed_at":"2026-08-04T21:18:53.361440Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:53.198974Z","title":"Beyond one-to-one: Rethinking the referring image segmentation,","venue":null,"work_id":"01c1b499-1c5a-4866-8796-29b8ae74f921","year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.585629Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:95a338c927f9126fc6472012967448ef1d5d21b0cd1deafca9fa59f8696c614e","observation_id":"ca104c85-2aab-4f40-b19c-c413103c314b","resolution":{"observed_at":"2026-08-04T21:18:53.265494Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:51.650913Z","title":"Swin transformer: Hierarchical vision transformer using shifted windows,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.650913Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:827c651fbf89f556968699289c75a4a076e0c4552a1223be208934ae40654c8e","observation_id":"7042b5c3-c0a9-4081-b500-9ac9c7ee02a4","resolution":{"observed_at":"2026-08-04T21:18:51.650913Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:51.721714Z","title":"Visual instruction tuning,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.721714Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:2c3300b3438df31c86cab6430da941d194f6bd8f81f9a7cadaeaf0716dee6cb4","observation_id":"83e6cb22-e098-4824-91e3-188219ba9a6a","resolution":{"observed_at":"2026-08-04T21:18:51.721714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-04T21:18:51.795036Z","title":"Llama: Open and efficient foundation language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.795036Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:4a68b58af174d79a7ea1f0e47d6676805e531e01d04bf04f5516ebb3a2d95976","observation_id":"ba582a63-8693-4bfe-bf3c-202c38716dfa","resolution":{"observed_at":"2026-08-04T21:18:51.795036Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:53.021447Z","title":"Cliport: What and where pathways for robotic manipulation,","venue":null,"work_id":"5e382165-7bec-42f4-8988-d7eccb64542d","year":2022},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.849521Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:a39d968e2084737d4123b1ecc745c22841131adaab0e026fc33b69a5e2fa5e73","observation_id":"8c39a48a-965b-4c86-9d3f-34d4ece7a7c6","resolution":{"observed_at":"2026-08-04T21:18:53.071323Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1810.04805","last_updated":"2019-05-24T20:37:26Z","snapshot_observed_at":"2026-07-30T09:12:38.100527Z","submitted_at":"2018-10-11T00:50:01Z","title":"BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1810.04805","snapshot_observed_at":"2026-08-04T21:18:51.940944Z","title":"Bert: Pre-training of deep bidirectional transformers for language understanding,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.940944Z"},"links":{"cited_paper":"/paper/1810.04805","citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:929e15e480fc2909a0020d2d4b0b997c10e0dc998441e0b1131f78f6c8936c09","observation_id":"d8f1c0dd-53c6-4a94-bc64-d9b27eef7fcd","resolution":{"observed_at":"2026-08-04T21:18:51.940944Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:51.996152Z","title":"Deep residual learning for image recognition,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:51.996152Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:9e92981a5b306fa5bd1b5fbe4ec4d30e8688128df8af349744e68356e2e46ae1","observation_id":"8fc73e5e-4e31-41fa-9639-11ffd5c1ba2c","resolution":{"observed_at":"2026-08-04T21:18:51.996152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:52.835189Z","title":"V-rep: A versatile and scalable robot simulation framework,","venue":null,"work_id":"e614fc91-ceec-485c-82d9-a2c6892a4a3d","year":2013},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:52.074546Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:be0d26b8dfa45385cf89f8671ae6cadb905c1a6c3cbb8fabf84b5e01ea972138","observation_id":"6eab9420-b1da-4f2c-ba31-a0af5d3b8f63","resolution":{"observed_at":"2026-08-04T21:18:52.893555Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:52.133839Z","title":"The ycb object and model set: Towards common benchmarks for manipulation research,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:52.133839Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:baa476d4473839f5ec6a35a298d9c581d89d5ec4275505df134db7fd9ed7ad4b","observation_id":"bd4cdfd4-0b00-4c33-9c97-c6386962dc7c","resolution":{"observed_at":"2026-08-04T21:18:52.133839Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:52.702925Z","title":"Vl-grasp: a 6-dof interactive grasp policy for language-oriented objects in cluttered indoor scenes,","venue":null,"work_id":"3a800ed2-bb9c-432a-8c30-32640134d863","year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:52.201302Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:36b2d5124f036436eb7a79d00ceb75b61cd6e5801aa909ac1591f7fc7ac4ad0b","observation_id":"1fce998e-260e-465e-9555-330492cfe572","resolution":{"observed_at":"2026-08-04T21:18:52.752470Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:52.268691Z","title":"Film: Visual reasoning with a general conditioning layer,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:52.268691Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:9a1e29cb0e6257642161dbb027b76446a9afa82058507b6069bf590c19c5cc93","observation_id":"b3b3bd3c-6539-4c1c-9b4e-ce902eb17a69","resolution":{"observed_at":"2026-08-04T21:18:52.268691Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T21:18:52.559730Z","title":"Vima: General robot manipulation with multimodal prompts,","venue":null,"work_id":"62d48fd1-ba58-42d5-aabe-95cf11725f39","year":2023},"citing_paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-04T21:18:52.322374Z"},"links":{"citing_paper":"/paper/2509.08126"},"observation_digest":"sha256:ce2c44ec40fc10a3f7123ccdbd800afd6a68fa3332ab34e12601c0420975bbd6","observation_id":"3b190952-349a-4234-b5e5-6ca3d43ccc41","resolution":{"observed_at":"2026-08-04T21:18:52.599568Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2509.08126","last_updated":"2025-09-09T20:07:51Z","latest_version":1,"primary_category":"cs.RO","snapshot_observed_at":"2026-08-09T17:56:41.614907Z","submitted_at":"2025-09-09T20:07:51Z","title":"Attribute-based Object Grounding and Robot Grasp Detection with Spatial Reasoning"},"reference_resolution":{"displayed":39,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":14,"verified_exact":1,"verified_fuzzy":24},"total_outbound_references":39},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 39 of 39 outbound references and 0 inbound Pith citation observations for arXiv:2509.08126."}