{"as_of":"2026-08-16T01:36:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:946217a5a86d6f0ba45e98eb6eeac72ba5db64cd01b4da8d6e6176ad3611ed66","coverage":[{"denominator":62,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":62,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T20:59:21.425916Z","state":"measured"},{"denominator":63,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":63,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-15T06:32:42.880941+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-08T08:15:20.055650Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-11T20:41:13.790288Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"cited_work":{"arxiv_id":"2511.17384","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2511.17384","snapshot_observed_at":"2026-07-07T03:17:14.809235Z","title":"Nikolai Helwig, Eliseo Pignanelli, and Andreas Schütze","venue":null,"work_id":"7dadd250-96ad-4c0d-a705-4312987cb149","year":2026},"citing_paper":{"arxiv_id":"2604.23446","last_updated":"2026-04-25T21:11:41Z","snapshot_observed_at":"2026-07-06T23:09:38.799345Z","submitted_at":"2026-04-25T21:11:41Z","title":"IndustryAssetEQA: A Neurosymbolic Operational Intelligence System for Embodied Question Answering in Industrial Asset Maintenance","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-08T08:15:20.055650Z"},"links":{"cited_paper":"/paper/2511.17384","citing_paper":"/paper/2604.23446"},"observation_digest":"sha256:7da7818ff6e75db9055b2da3427acea6641d3941a9e2efc1e8ee1c5b45a5d82f","observation_id":"791a46fe-c4ea-4b72-bf9b-21dc7da14591","resolution":{"observed_at":"2026-07-07T03:17:14.809235Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2511.17384/citation-record","integrity":"/paper/2511.17384/integrity","json":"/paper/2511.17384/citation-record.json","paper":"/paper/2511.17384"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:14.715333Z","title":"Vision-and-language navigation: Interpreting visually-grounded navigation instructions in real environments","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:14.715333Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:4f98de261155507697bca051c8d6715621cee0ef2c77dd77c99a59f8d33c98dd","observation_id":"268aea4b-6a76-4493-b6a7-3b2bea92801e","resolution":{"observed_at":"2026-08-03T20:59:14.715333Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:14.763941Z","title":"Claude sonnet 4.5, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:14.763941Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:24adaa389c375759257f2d80fd32775e633ce840160c9511d62a0b85e4ef552b","observation_id":"fff39733-19ee-4d9f-b245-319af0a0ffbd","resolution":{"observed_at":"2026-08-03T20:59:14.763941Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:14.841075Z","title":"Scanqa: 3d question answering for spatial scene understanding","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:14.841075Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:f62d19a8a39797f1a1a6cc78d936a6e4da8b9b2088da64ae56bf9426d7874ff2","observation_id":"a715cad2-004a-4c2d-9afa-530ae07f1f30","resolution":{"observed_at":"2026-08-03T20:59:14.841075Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:14.904407Z","title":"The r2r framework: Publishing and discovering mappings on the web.COLD, 665:97–108, 2010","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:14.904407Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:eeef0ca69aee31a4f0954326fd83b026e951a444283aa90bef6beb55e4487ad6","observation_id":"e1316510-d3ef-43f6-927f-e27919579d7c","resolution":{"observed_at":"2026-08-03T20:59:14.904407Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:15.029846Z","title":"Depth pro: Sharp monocular metric depth in less than a second","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:15.029846Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:f744e2a9309515e76c813463a01e750041551e0a0dc638db957320f41593687f","observation_id":"d46d9a2b-bba0-4882-8bb3-1bcbd2896d7e","resolution":{"observed_at":"2026-08-03T20:59:15.029846Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:15.122121Z","title":"Spatialbot: Pre- cise spatial understanding with vision language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:15.122121Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:53ebc10e66a9b86d2f0c5d4bef5fcd2f8611902a539c2f10aca306e20b3714a6","observation_id":"45290bed-d13d-4b59-865c-81cd9e2338da","resolution":{"observed_at":"2026-08-03T20:59:15.122121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:15.198676Z","title":"Partnr: A benchmark for planning and rea- soning in embodied multi-agent tasks","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:15.198676Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:c27e64873dd2e7b2725d9b9922c24b080f10fd2cef19b8f6bc792f80dfd770ec","observation_id":"72cb38fa-cc5a-4756-b80f-673c8889a044","resolution":{"observed_at":"2026-08-03T20:59:15.198676Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:15.275007Z","title":"Neural topological slam for vi- sual navigation","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:15.275007Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:29d0e0bf80913dfe838fe2236b34f01402f0d0311bbea5d503e772f30742f1e9","observation_id":"476e92a7-c27a-4919-8020-43ca8735e671","resolution":{"observed_at":"2026-08-03T20:59:15.275007Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:15.365432Z","title":"Spatialvlm: Endow- ing vision-language models with spatial reasoning capabili- ties","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:15.365432Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:1d5f673a9efd2497dcb30a8644b821f8d0fa05ffb136ffb20a4c7680f7c9d454","observation_id":"2a3d542c-e622-4417-884b-5ff4825f8d78","resolution":{"observed_at":"2026-08-03T20:59:15.365432Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:15.442995Z","title":"Spatial- rgpt: Grounded spatial reasoning in vision-language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:15.442995Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:e213648561708a57486356af5e94b5bf275c372fa9399179102dfa7dc85ae79f","observation_id":"cefecfe8-cdc9-40e2-8f00-7ee8d7c0df13","resolution":{"observed_at":"2026-08-03T20:59:15.442995Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.20263","last_updated":"2025-08-08T23:10:26Z","snapshot_observed_at":"2026-08-12T22:14:34.765661Z","submitted_at":"2024-10-26T19:48:47Z","title":"EfficientEQA: An Efficient Approach to Open-Vocabulary Embodied Question Answering","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.20263","snapshot_observed_at":"2026-08-03T20:59:15.554071Z","title":"Efficienteqa: An efficient approach for open vocabulary embodied question answering.arXiv preprint arXiv:2410.20263, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:15.554071Z"},"links":{"cited_paper":"/paper/2410.20263","citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:84c3eb4b93b6d01da33c3c90a8be5e82ff5219da5ed5408353486dc7b5ed9066","observation_id":"9f1d3e99-879f-44e2-b031-4adefe539988","resolution":{"observed_at":"2026-08-03T20:59:15.554071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:15.645170Z","title":"Lota-bench: Benchmarking language- oriented task planners for embodied agents","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:15.645170Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:7c2f28c4597cca915af79ef0a640a1c9824c3b46e1895e37c68ed0b172aa0584","observation_id":"0e6dd7fd-825c-49d4-b05b-7a182c8e543b","resolution":{"observed_at":"2026-08-03T20:59:15.645170Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:15.719156Z","title":"Embodied question answer- ing","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:15.719156Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:b21547c8300173de53ecd6bade3eb592d9216d1dea5741fad9951624c0789574","observation_id":"7b1e12fb-aa81-4e6d-accc-3d3433d8ed9d","resolution":{"observed_at":"2026-08-03T20:59:15.719156Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:15.824408Z","title":"Mm-spatial: Exploring 3d spatial understanding in multimodal llms","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:15.824408Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:379e8658437d4faec9b9dbd88c14d8d3e39292aa4a5946fd98f72545fcb2a01e","observation_id":"3cacc8f1-7bf6-4548-ad4c-bb7f74e49430","resolution":{"observed_at":"2026-08-03T20:59:15.824408Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:16.041496Z","title":"EmbSpatial-bench: Benchmarking spatial un- derstanding for embodied tasks with large vision-language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:16.041496Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:38f2ebb7c41362e116f13ffb9ea3b304427aa44126559d96aa1528a9520e63c5","observation_id":"2e83de45-62df-4af9-a141-c834b9dfef76","resolution":{"observed_at":"2026-08-03T20:59:16.041496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:16.193029Z","title":"Blink: Multimodal large language models can see but not perceive","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:16.193029Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:6e3672dbaf2e539a59511ebaade467b00edad65060fc444d08b4f935fc275e22","observation_id":"98aa84ae-30bd-422c-b69d-63ade9c05ff9","resolution":{"observed_at":"2026-08-03T20:59:16.193029Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:16.345965Z","title":"Gemini 2.5 Flash Preview Model Card","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:16.345965Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:f366293c0a93946caa7ba8c7f8d4ee2b31158b0cb04bf792e90dd028a4fa5b06","observation_id":"ec07af9e-1545-4e62-a124-8dc157f90bb5","resolution":{"observed_at":"2026-08-03T20:59:16.345965Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:16.528121Z","title":"Autotag & tagmap: Llm-powered moodle plugins for peda- gogical alignment checks.SN Computer Science, 6(7):1–11,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:16.528121Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:82d4e36d4e728980f70c93ffa66fc28c6f8847641302a4d26983f335ec6521b4","observation_id":"6b753aa9-19db-4d46-b6d6-6a26785b6a5d","resolution":{"observed_at":"2026-08-03T20:59:16.528121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:16.586819Z","title":"3d-llm: In- jecting the 3d world into large language models.Advances in Neural Information Processing Systems, 36:20482–20494,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:16.586819Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:637998c8b8c53ad22586c5f283b41bbd79c631f09b6545a82730ac46828e23f8","observation_id":"ecebcd8e-2084-4abe-8dd4-e8ed58015860","resolution":{"observed_at":"2026-08-03T20:59:16.586819Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:16.681143Z","title":"Enhancing visualization and interaction of complex spatial data through augmented reality.The International Journal of Advanced Manufacturing Technology, 134(11):5891–5906, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:16.681143Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:b43477995782740d801bfd6cd127d321e87c2e7395e8c94beaa1249b3791210b","observation_id":"b6fca1c9-e57a-4da3-8d2a-30822944f937","resolution":{"observed_at":"2026-08-03T20:59:16.681143Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:16.789318Z","title":"Visual language maps for robot navigation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:16.789318Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:850b1394f200c48f8aa409df2c128fe0458693badcc1ebf762deabda31e83ec1","observation_id":"37175fff-ffaf-4c68-8a87-3a45f0e2aa43","resolution":{"observed_at":"2026-08-03T20:59:16.789318Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:16.938067Z","title":"Gqa: A new dataset for real-world visual reasoning and compositional question answering","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:16.938067Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:ee68ef43bfd531c4ccc0fac5e132a97618c2f96c2c3ee0fd62102b9378c990a3","observation_id":"5e29cb65-16c6-4860-9b01-b26b23e2ddf5","resolution":{"observed_at":"2026-08-03T20:59:16.938067Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:17.069487Z","title":"Clevr: A diagnostic dataset for compositional language and elementary visual reasoning","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:17.069487Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:bc46d4fd7c5f5477beb5548200aa292af275b94d3ac5cf950687688c51c3277d","observation_id":"2df73521-921e-4e4f-beea-9acaca451b97","resolution":{"observed_at":"2026-08-03T20:59:17.069487Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.09246","last_updated":"2024-09-05T19:46:34Z","snapshot_observed_at":"2026-08-15T04:55:15.159462Z","submitted_at":"2024-06-13T15:46:55Z","title":"OpenVLA: An Open-Source Vision-Language-Action Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.09246","snapshot_observed_at":"2026-08-03T20:59:17.162745Z","title":"Openvla: An open-source vision-language-action model.arXiv preprint arXiv:2406.09246, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:17.162745Z"},"links":{"cited_paper":"/paper/2406.09246","citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:fcb0a9e401ef53df3f1cfe8a99ac6a844b6b88218b9a6b2b83b4025945f35a48","observation_id":"08ff3179-7c2d-4476-b1ff-039788fa47d5","resolution":{"observed_at":"2026-08-03T20:59:17.162745Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:17.304591Z","title":"Llava-next: Stronger llms supercharge multimodal capa- bilities in the wild, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:17.304591Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:22eadc1d4a83942eba5d1d400206d5243d2dcef5e8824944f5b5baa386b63f95","observation_id":"c194c02a-a13c-46eb-8516-888ead71a700","resolution":{"observed_at":"2026-08-03T20:59:17.304591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:17.444715Z","title":"Behavior-1k: A benchmark for embodied ai with 1,000 ev- eryday activities and realistic simulation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:17.444715Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:35ace9d7eec0da54faddf272611ca33fb65a70317670fdab402a22aa90107c65","observation_id":"60cb29bc-4cf0-4d6c-8195-ffdf541b2936","resolution":{"observed_at":"2026-08-03T20:59:17.444715Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20640","last_updated":"2025-05-27T02:36:17Z","snapshot_observed_at":"2026-08-09T03:09:00.133004Z","submitted_at":"2025-05-27T02:36:17Z","title":"IndustryEQA: Pushing the Frontiers of Embodied Question Answering in Industrial Scenarios","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.20640","snapshot_observed_at":"2026-08-03T20:59:17.546623Z","title":"Industryeqa: Push- ing the frontiers of embodied question answering in indus- trial scenarios.arXiv preprint arXiv:2505.20640, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:17.546623Z"},"links":{"cited_paper":"/paper/2505.20640","citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:9208f72b2f14ce6f0416d89bb6cfcde810044e135a067007dcb2527c61904a6e","observation_id":"36373aeb-9c68-4af6-8c82-30c4f13623dc","resolution":{"observed_at":"2026-08-03T20:59:17.546623Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.02765","last_updated":"2025-01-06T05:15:59Z","snapshot_observed_at":"2026-08-14T12:32:34.328935Z","submitted_at":"2025-01-06T05:15:59Z","title":"Visual Large Language Models for Generalized and Specialized Applications","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.02765","snapshot_observed_at":"2026-08-03T20:59:17.714781Z","title":"Visual large language models for generalized and specialized applications.arXiv preprint arXiv:2501.02765, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:17.714781Z"},"links":{"cited_paper":"/paper/2501.02765","citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:d769c1be489ef3856cdf6c6f7bbf5202452ea8d3a2cd8c5ddb42f01a0fa10bf0","observation_id":"cda5e9a0-1daf-4993-89dd-7463904e7c07","resolution":{"observed_at":"2026-08-03T20:59:17.714781Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:17.888807Z","title":"Toa: Task-oriented active vqa.Advances in Neural Information Processing Systems, 36:54061–54074, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:17.888807Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:63a0ebeb63fdf9a96c9508c5db557eeae79a7ffa513777d34bb0adfe84b56a04","observation_id":"d236cf6a-e475-4c54-aae7-7975b8ff6e5e","resolution":{"observed_at":"2026-08-03T20:59:17.888807Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:18.085542Z","title":"Reasoning paths with reference objects elicit quanti- tative spatial reasoning in large vision-language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:18.085542Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:279c379e360049a99cd1df3a0f243cca15e164cf00d84111f8be9b1a90ba2f78","observation_id":"52fa8124-bae6-4ffa-8ba4-bb8d0e84e668","resolution":{"observed_at":"2026-08-03T20:59:18.085542Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:18.251368Z","title":"Navcot: Boosting llm-based vision-and- language navigation via learning disentangled reasoning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:18.251368Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:1298844ce099e377bbc0f88097980c5a1fc3e16985678320e49fbbb3b76f7e36","observation_id":"c88596c5-ca77-47bf-8c04-50372fe13a43","resolution":{"observed_at":"2026-08-03T20:59:18.251368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:18.417328Z","title":"3dsrbench: A comprehensive 3d spatial reasoning benchmark","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:18.417328Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:77f3e482e1acc3cadd1dc69d9ad652dd6e0f27f8087f9fc3eaff5767e88452ac","observation_id":"d26b32e1-6cd7-43ea-bea9-83ddca0f94b1","resolution":{"observed_at":"2026-08-03T20:59:18.417328Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:18.531424Z","title":"Openeqa: Embodied question answering in the era of foun- dation models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:18.531424Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:a469cbd6cd1bdf628dacf0e9e0dd3c44bf037135b2cecc0cdc961bca2c9ab262","observation_id":"7a529095-9256-4f69-bc28-ef47feddd9f6","resolution":{"observed_at":"2026-08-03T20:59:18.531424Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:18.645377Z","title":"Llama 4 model card, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:18.645377Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:e491f1c2c6013a3607e18fa92c14683fba7d670004ccfb3af9c14fd8ec940b86","observation_id":"1eca95f3-b0d0-4ee2-bde1-a378358b7035","resolution":{"observed_at":"2026-08-03T20:59:18.645377Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.14444","last_updated":"2025-09-02T16:12:36Z","snapshot_observed_at":"2026-07-06T22:15:27.120497Z","submitted_at":"2025-08-20T06:00:57Z","title":"NVIDIA Nemotron Nano 2: An Accurate and Efficient Hybrid Mamba-Transformer Reasoning Model","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.14444","snapshot_observed_at":"2026-08-03T20:59:18.876864Z","title":"Efficient hybrid mamba- transformer reasoning model.arXiv preprint arXiv:2508.14444, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:18.876864Z"},"links":{"cited_paper":"/paper/2508.14444","citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:817287aa5a41127a5b3f02ad373f140c34fafaf1852846ba039c5242af7d6346","observation_id":"9f214e3c-ab9c-47e3-9bc7-8acec09677b4","resolution":{"observed_at":"2026-08-03T20:59:18.876864Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:19.012998Z","title":"Gpt-4o system card, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:19.012998Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:afb55524b0c50c21e8156776acc183051d79ec0b87f1220957d44adffc27c7c3","observation_id":"15f8da26-30c1-40b3-9448-e090af8c6b63","resolution":{"observed_at":"2026-08-03T20:59:19.012998Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:19.090182Z","title":"Is map- ping necessary for realistic pointgoal navigation? InCVPR, pages 17232–17241, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:19.090182Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:310267190bb7d60f88b33b299e3bd7950a2a2f07149790736da8852f2b26dd72","observation_id":"fb2208dd-9e1b-43b9-a9c9-64af267eae17","resolution":{"observed_at":"2026-08-03T20:59:19.090182Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:19.226498Z","title":"Reverie: Remote embodied visual referring expres- sion in real indoor environments","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:19.226498Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:0eda61b660a3df59dbe2f0b03a8d2d603b1f05f564bb7ea01a504ec4d8423885","observation_id":"997268c5-8c2b-4b46-b378-cc198d0708d8","resolution":{"observed_at":"2026-08-03T20:59:19.226498Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:19.309828Z","title":"Habitat: A plat- form for embodied ai research","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:19.309828Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:941e3da1a391f6d16833db26638b2cdece50ed6ae6b101ea428c5efc153e0715","observation_id":"fc97cf25-f771-4dc5-8510-788cf0a31a20","resolution":{"observed_at":"2026-08-03T20:59:19.309828Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:19.394606Z","title":"Learning to navigate using mid-level visual priors","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:19.394606Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:a02018fbf72c086f24eeaf91464290ff012a49542d7a2c840c8f6898adf667ec","observation_id":"8e917696-02e1-4163-a172-54f82124f09e","resolution":{"observed_at":"2026-08-03T20:59:19.394606Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:19.493879Z","title":"Alfred: A benchmark for interpreting grounded instructions for everyday tasks","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:19.493879Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:ca2a7391de14922b9fe92e39ff6212951135257a3b8b9c9450bc5fbc2089bd39","observation_id":"9926cc71-3b36-4f0c-a24f-756db9d3b9b3","resolution":{"observed_at":"2026-08-03T20:59:19.493879Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:19.588095Z","title":"Robospatial: Teaching spatial understanding to 2d and 3d vision-language models for robotics","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:19.588095Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:1edabdc1dddbdc3ca60e76b94c56200f33726bde91d5f26b6ea81e5ae6e232cb","observation_id":"0977eb82-99cb-4903-9faf-db975a2bab58","resolution":{"observed_at":"2026-08-03T20:59:19.588095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:19.714172Z","title":"Habitat 2.0: Training home assistants to rearrange their habitat","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:19.714172Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:f072ee465db3afaa25b4ebe5115c58ff70a9784c44e3dabb961818ac6bf8c197","observation_id":"29b6991b-1ce4-4ec4-aea9-940d66551d1e","resolution":{"observed_at":"2026-08-03T20:59:19.714172Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:19.785212Z","title":"Cambrian-1: A fully open, vision-centric explo- ration of multimodal llms","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:19.785212Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:16f3e4ddabc3f12449e0b273a68b96b586653f3301c6a6d7fd6a4ecc3bbec86f","observation_id":"3eefe272-89c0-44a8-8feb-7659bd3d526e","resolution":{"observed_at":"2026-08-03T20:59:19.785212Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:19.976111Z","title":"Hierarchical open- vocabulary 3d scene graphs for language-grounded robot navigation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:19.976111Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:13c5512c3b1182717c455a61a73eacb198185e8c58dab898f31daaa864adb311","observation_id":"6c36bd83-474b-4fa9-996b-2095b57bcc2f","resolution":{"observed_at":"2026-08-03T20:59:19.976111Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.030245Z","title":"Vsp: Diagnosing the dual challenges of perception and reasoning in spatial planning tasks for mllms","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.030245Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:9c411f046d858d22db2c1d8cfb600c2c40a9f3cc9259345baac01406d9a225ca","observation_id":"3f197a8c-8d56-4e54-ac48-ccbadf420a10","resolution":{"observed_at":"2026-08-03T20:59:20.030245Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.079159Z","title":"The rise and potential of large language model based agents: A survey.SCIS, 68(2):121101, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.079159Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:894e1ebda36925f5809fdf1bf2f7ec1034bae218b6058af4db4c9edee4adada7","observation_id":"5a8bd369-503a-41de-897e-9d9468ff45b2","resolution":{"observed_at":"2026-08-03T20:59:20.079159Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.204740Z","title":"Gibson env: Real-world percep- tion for embodied agents","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.204740Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:ea4ab4ee7740351fc6fbd3b8738bedba476a776cefb9a876d3bba7eb968a1cbd","observation_id":"333a32d1-fe5d-4c6c-83ea-67a61c2dccd5","resolution":{"observed_at":"2026-08-03T20:59:20.204740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.396297Z","title":"Expand vsr benchmark for vllm to expertize in spatial rules","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.396297Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:7c900acf50f1497fdde3253e1180287c17c54b511b32725e2c19dbdf891c191a","observation_id":"99632ebf-e066-4e7c-acbf-d36d4527fcdd","resolution":{"observed_at":"2026-08-03T20:59:20.396297Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.488266Z","title":"Point2graph: An end-to-end point cloud- based 3d open-vocabulary scene graph for robot navigation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.488266Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:21e5fdaa428cf35de0a54c565563eb39907c235cd9288ecdfdf72a39266a193a","observation_id":"d0ee4011-2ca3-4fac-a51f-8e67ffc2c750","resolution":{"observed_at":"2026-08-03T20:59:20.488266Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-08-03T20:59:20.497223Z","title":"Qwen3 technical report.arXiv preprint arXiv:2505.09388, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.497223Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:d6313dfbaf02c95789ab039d44bae3ca7b873fb43db1b3c390dfd748cd07e411","observation_id":"167388c3-ba6e-42ae-a5f7-69a2391588ef","resolution":{"observed_at":"2026-08-03T20:59:20.497223Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.586060Z","title":"Thinking in space: How mul- timodal large language models see, remember, and recall spaces","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.586060Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:ba5386fc1a06cd729fb45955883d876fc7bc38ec3a6d5a49e43b340c5c31b8eb","observation_id":"d69784f3-25bc-4ada-934d-12dce8081bd1","resolution":{"observed_at":"2026-08-03T20:59:20.586060Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.658315Z","title":"Rila: Re- flective and imaginative language agent for zero-shot seman- tic audio-visual navigation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.658315Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:4edcb1b26045fccbe9cd0db83aa5589a1820934dfe839e18d57b478da25b322e","observation_id":"9f350c7a-08d3-4c90-b2ac-93747a099db6","resolution":{"observed_at":"2026-08-03T20:59:20.658315Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.752621Z","title":"Sg-nav: Online 3d scene graph prompting for llm-based zero-shot object navigation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.752621Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:20dd73d790b9cd774293fd6a21d676c17952b3cacdcfd1b19ec508a4fd4ae11b","observation_id":"df90fd26-8854-4519-8e06-005483e7803e","resolution":{"observed_at":"2026-08-03T20:59:20.752621Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.844060Z","title":"L3mvn: Leveraging large language models for visual target naviga- tion","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.844060Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:60d3d1125fd6fbb4c66f121c424071f269e6f4730bde500d38de2bfc9ff09180","observation_id":"7cd8029e-c90d-4f3f-8bbf-38d1ff3e3d9e","resolution":{"observed_at":"2026-08-03T20:59:20.844060Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.908530Z","title":"A continual learning approach for em- bodied question answering with generative adversarial imi- tation learning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.908530Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:f3dd7e67fe099e85f7727478abdcc1e68130eea440519e6fc024163a81ffb3aa","observation_id":"ca3c587d-bf40-4c45-85cb-59ba2f46d33e","resolution":{"observed_at":"2026-08-03T20:59:20.908530Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:20.985488Z","title":"Open3d-vqa: A benchmark for embodied spatial concept reasoning with multimodal large language model in open space","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:20.985488Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:e68b269aded136b25a0db8b24ae0e345bcb01e618a79a21f363469feb944c4c0","observation_id":"cbe7c04a-71e2-4313-801b-0ee0d80291c1","resolution":{"observed_at":"2026-08-03T20:59:20.985488Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:21.043075Z","title":"Dsi- bench: A benchmark for dynamic spatial intelligence.arXiv preprint arXiv:2510.18873, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:21.043075Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:b8fbc1f1f6d58c54c2a8f6adca4307600fb5e542d249de9450400d0f81fc3d0e","observation_id":"4f3cceb3-5137-41f3-b24c-51aa056d3546","resolution":{"observed_at":"2026-08-03T20:59:21.043075Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.12532","last_updated":"2025-05-22T00:44:13Z","snapshot_observed_at":"2026-08-14T15:27:38.086844Z","submitted_at":"2025-02-18T04:36:15Z","title":"CityEQA: A Hierarchical LLM Agent on Embodied Question Answering Benchmark in City Space","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.12532","snapshot_observed_at":"2026-08-03T20:59:21.103190Z","title":"Cityeqa: A hierarchical llm agent on embodied question answering benchmark in city space.arXiv preprint arXiv:2502.12532, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:21.103190Z"},"links":{"cited_paper":"/paper/2502.12532","citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:75517c73e39583a0d8c5508633a57c0a1ba660d03747d19ecae6e63d046df799","observation_id":"34d72e47-e672-477d-9f5a-7cd5feacd7c2","resolution":{"observed_at":"2026-08-03T20:59:21.103190Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:21.199741Z","title":"3d- vla: a 3d vision-language-action generative world model","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:21.199741Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:1aa573332f7ffcdf13b56aeee57f26694657e4022ac1b4b3cfcc94cb36ffefec","observation_id":"34ba3d3d-319a-41ee-9b4f-20f237f6b381","resolution":{"observed_at":"2026-08-03T20:59:21.199741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:21.305893Z","title":"Towards learning a generalist model for embodied navigation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:21.305893Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:d632a12415b53431973002975aabcd90f1a0edadf7d1332b5662211676acd465","observation_id":"5b4e9aab-9716-4780-974b-9cd6bff2d5bc","resolution":{"observed_at":"2026-08-03T20:59:21.305893Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T20:59:21.425916Z","title":"reasoning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-03T20:59:21.425916Z"},"links":{"citing_paper":"/paper/2511.17384"},"observation_digest":"sha256:06a5f6f697fbf8e329e5e4cedcb759d325652565641a1d2ee46e4c04d50dc2a4","observation_id":"caf8d0e8-7d2c-4ddd-9129-4645877eec85","resolution":{"observed_at":"2026-08-03T20:59:21.425916Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2511.17384","last_updated":"2026-07-03T23:12:55Z","latest_version":2,"primary_category":"cs.RO","snapshot_observed_at":"2026-08-14T14:14:30.533453Z","submitted_at":"2025-11-21T16:48:49Z","title":"IndustryNav: Exploring Spatial Reasoning of Embodied Agents in Dynamic Industrial Navigation"},"reference_resolution":{"displayed":62,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":62,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":62},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 62 of 62 outbound references and 1 inbound Pith citation observation for arXiv:2511.17384."}