{"as_of":"2026-08-16T05:30:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:fcad60ab343f965e3c80e83f56704620087570d22ce5b0dd5b6521625e394fc0","coverage":[{"denominator":53,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":53,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T16:09:09.577056Z","state":"measured"},{"denominator":55,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":55,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-15T06:32:42.880941+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-29T04:26:43.018373Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-03T06:07:41.698096Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"cited_work":{"arxiv_id":"2509.08757","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.08757","snapshot_observed_at":"2026-07-03T06:07:41.698096Z","title":"arXiv preprint arXiv:2509.08757 , year=","venue":null,"work_id":"dbae9ae0-7d2f-4440-bfc7-92816c1a72c7","year":2025},"citing_paper":{"arxiv_id":"2606.10495","last_updated":"2026-06-09T07:18:01Z","snapshot_observed_at":"2026-08-13T05:21:01.334976Z","submitted_at":"2026-06-09T07:18:01Z","title":"Act on What You See: Unlocking Safe Social Navigation in Vision-Language-Action Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-06-27T12:52:38.764427Z"},"links":{"cited_paper":"/paper/2509.08757","citing_paper":"/paper/2606.10495"},"observation_digest":"sha256:e2e5cad59937d85317fc13e355f55e7042422948870b452977f20eae5b5fe04b","observation_id":"fbe228f8-93eb-4a05-9d7c-aa608b56c5de","resolution":{"observed_at":"2026-07-03T06:07:41.699543Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"cited_work":{"arxiv_id":"2509.08757","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.08757","snapshot_observed_at":"2026-07-03T06:07:41.698096Z","title":"arXiv preprint arXiv:2509.08757 , year=","venue":null,"work_id":"dbae9ae0-7d2f-4440-bfc7-92816c1a72c7","year":2025},"citing_paper":{"arxiv_id":"2606.27826","last_updated":"2026-08-10T02:19:01Z","snapshot_observed_at":"2026-08-14T05:08:22.666221Z","submitted_at":"2026-06-26T08:10:39Z","title":"NormAct: Benchmarking Embodied Agents' Proactive Compliance with Unspoken Social Norms","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-06-29T04:26:43.018373Z"},"links":{"cited_paper":"/paper/2509.08757","citing_paper":"/paper/2606.27826"},"observation_digest":"sha256:c928abd17e3eb40b5dac659a2a521d41e70bdde4aaecaedde247b2d853a2f145","observation_id":"2858a009-7d4f-4402-86ab-5a3ad7b4ba71","resolution":{"observed_at":"2026-07-01T16:55:51.087519Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2509.08757/citation-record","integrity":"/paper/2509.08757/integrity","json":"/paper/2509.08757/citation-record.json","paper":"/paper/2509.08757"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:09.417890Z","title":"Mavrogiannis, F","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.417890Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:baad33a5e94d25be4323c5b048444e18b22b759a3cc4433afb87ccc74b7cfa49","observation_id":"aeeffb35-f605-4aa8-94bc-8c6cf62c638d","resolution":{"observed_at":"2026-08-15T16:09:09.417890Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.16740","last_updated":"2023-09-19T20:02:06Z","snapshot_observed_at":"2026-08-13T11:06:59.881605Z","submitted_at":"2023-06-29T07:31:43Z","title":"Principles and Guidelines for Evaluating Social Robot Navigation Algorithms","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.16740","snapshot_observed_at":"2026-08-15T16:09:09.421600Z","title":"Francis, C","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.421600Z"},"links":{"cited_paper":"/paper/2306.16740","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:b42b008af84577ab0db9342f632a686eb75d410bc309260cffccba9608829879","observation_id":"b0405286-1fd5-4acf-88c1-cf25d8ab3a48","resolution":{"observed_at":"2026-08-15T16:09:09.421600Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:09.425337Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.425337Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:164ac65eae19d43717477f6978c6b2f82e45036f2ca22b74ba079f837d077279","observation_id":"6447513a-75ca-4557-8dce-c835044cb74c","resolution":{"observed_at":"2026-08-15T16:09:09.425337Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2203.15041","last_updated":"2022-06-08T20:24:44Z","snapshot_observed_at":"2026-08-13T16:13:51.227367Z","submitted_at":"2022-03-28T19:09:11Z","title":"Socially Compliant Navigation Dataset (SCAND): A Large-Scale Dataset of Demonstrations for Social Navigation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.15041","snapshot_observed_at":"2026-08-15T16:09:09.428939Z","title":"Karnan, A","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.428939Z"},"links":{"cited_paper":"/paper/2203.15041","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:a2c5b1a0495892365f13405513d01dedc584c290ecc6372b3caa3223d9eb1429","observation_id":"7d3a114c-99c5-4471-ae2d-e1ef0cd124d1","resolution":{"observed_at":"2026-08-15T16:09:09.428939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08485","last_updated":"2023-12-11T17:46:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-17T17:59:25Z","title":"Visual Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08485","snapshot_observed_at":"2026-08-15T16:09:09.432524Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.432524Z"},"links":{"cited_paper":"/paper/2304.08485","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:95a3ed47b240f085804fbd381996481e2bb2b875542328146cee44477987ba48","observation_id":"63eb2d17-a6fc-454a-94da-62ae741c4591","resolution":{"observed_at":"2026-08-15T16:09:09.432524Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21276","last_updated":"2024-10-25T17:43:01Z","snapshot_observed_at":"2026-08-15T14:02:47.366139Z","submitted_at":"2024-10-25T17:43:01Z","title":"GPT-4o System Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21276","snapshot_observed_at":"2026-08-15T16:09:09.435558Z","title":"Hurst, A","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.435558Z"},"links":{"cited_paper":"/paper/2410.21276","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:04e36a9132ec30d7eed7f8fc283b31373319281d7f8808818327c4112d54ddbe","observation_id":"b65b594a-f1ed-459e-af26-17f707c11d5f","resolution":{"observed_at":"2026-08-15T16:09:09.435558Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05530","last_updated":"2024-12-16T17:39:39Z","snapshot_observed_at":"2026-08-14T18:15:53.516440Z","submitted_at":"2024-03-08T18:54:20Z","title":"Gemini 1.5: Unlocking multimodal understanding across millions of tokens of context","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.05530","snapshot_observed_at":"2026-08-15T16:09:09.438716Z","title":"Team and P","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.438716Z"},"links":{"cited_paper":"/paper/2403.05530","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:0464608337e8f3979595a6520be958a64c2b66b1810879df3fe905474fe2ea7c","observation_id":"c518c8fa-d854-4769-b304-346d5a20b67f","resolution":{"observed_at":"2026-08-15T16:09:09.438716Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.03000","last_updated":"2024-10-10T06:19:33Z","snapshot_observed_at":"2026-08-15T00:12:08.778411Z","submitted_at":"2024-07-03T10:59:06Z","title":"VIVA: A Benchmark for Vision-Grounded Decision-Making with Human Values","version":2},"cited_work":{"arxiv_id":"2407.03000","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.03000","snapshot_observed_at":"2026-08-15T16:09:09.993204Z","title":"VIVA: A Benchmark for Vision-Grounded Decision-Making with Human Values","venue":"cs.CL","work_id":"6f2db67d-31cd-47af-938d-c6764cac736d","year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.441727Z"},"links":{"cited_paper":"/paper/2407.03000","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:bcb494c513f60ea22b68d2ff97490e14930fb0f48c0ab7f0e810b6a992c23515","observation_id":"4b70a516-edb5-46ef-bb02-681972e566a3","resolution":{"observed_at":"2026-08-15T16:09:09.996392Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.00210","last_updated":"2024-11-25T21:05:42Z","snapshot_observed_at":"2026-08-15T13:47:54.200227Z","submitted_at":"2024-03-30T01:17:40Z","title":"VLM-Social-Nav: Socially Aware Robot Navigation through Scoring using Vision-Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.00210","snapshot_observed_at":"2026-08-15T16:09:09.445325Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.445325Z"},"links":{"cited_paper":"/paper/2404.00210","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:2e564f28fb891f642987d7b9c5b73c3afb311b99ebf8fa913f701249bb98aa91","observation_id":"f92f7ba9-d034-4a04-b88e-06956f66d3d7","resolution":{"observed_at":"2026-08-15T16:09:09.445325Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.06468","last_updated":"2025-04-18T03:53:04Z","snapshot_observed_at":"2026-08-14T23:08:11.906639Z","submitted_at":"2024-10-09T01:41:49Z","title":"Does Spatial Cognition Emerge in Frontier Models?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.06468","snapshot_observed_at":"2026-08-15T16:09:09.448556Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.448556Z"},"links":{"cited_paper":"/paper/2410.06468","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:ba7d53467dc450c28eeb2f706edb07774ee8065c420a0a68154ad1f325c200f8","observation_id":"9801f013-3d3e-4a5e-b857-dbc8819f20a8","resolution":{"observed_at":"2026-08-15T16:09:09.448556Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.317258Z","title":"Kessler, J","venue":null,"work_id":"4669214d-3b02-4ee3-b490-5d961d63938f","year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.451625Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:734a5b12c860ba2b35e63014fa04c8e1b015b96e91db4355983475be05dc086a","observation_id":"940f8321-1c09-4c98-a324-73122060ddd1","resolution":{"observed_at":"2026-08-15T16:09:10.320106Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1016/j.neuron.2023.03.001","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:09.611961Z","title":null,"venue":null,"work_id":"05cd834e-119b-4f29-b1bd-7e9ce587b879","year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.454559Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:0e6951bf588fe15ceca7847b8e7ee55d9961daf58bbe18280080384d520d2291","observation_id":"2c796588-7cd6-470d-bb62-24960ed44cc4","resolution":{"observed_at":"2026-08-15T16:09:09.615963Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.308664Z","title":"Moussa ¨ıd, N","venue":null,"work_id":"14bc6594-e059-44f7-96dc-d754cf392511","year":2010},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.458490Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:d12e06bfa37e60b65326536bb907a23e8a1792c70d97cc6c52a61cef4a9859c4","observation_id":"8b8b21a3-c83d-49f9-9908-63725576a1b0","resolution":{"observed_at":"2026-08-15T16:09:10.311629Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.300302Z","title":"Moussa ¨ıd, D","venue":null,"work_id":"9865b90b-1afc-4c94-bc2a-fdc8b339de85","year":2009},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.461437Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:d8e8d9edebc0f9cb5ed1838df3ef266b3e2139ed37691f3fbd85f2843101286c","observation_id":"08dfe599-780d-46fe-b297-67c9d7828899","resolution":{"observed_at":"2026-08-15T16:09:10.303145Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.291606Z","title":"Karnan, A","venue":null,"work_id":"331b4b5e-757e-4e7a-ae82-7a22992aa461","year":2022},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.464332Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:94ae6028fd87d86b45534732007d663a2a469bec95c1dda194513d7b75d73a0c","observation_id":"d44df9f5-8b7b-41ed-bf99-c6aff76b541f","resolution":{"observed_at":"2026-08-15T16:09:10.294422Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.282867Z","title":"Openai o3 and o4-mini system card, 2025","venue":null,"work_id":"b6047014-bb1d-48ba-85fe-1be505e5eb82","year":2025},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.467296Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:9ff9c6eabd9d2845eeaa561f60477c67052d1523597d98a2f15d4fc58eec59b0","observation_id":"da918603-7765-472f-b99a-aa9383a5e0e9","resolution":{"observed_at":"2026-08-15T16:09:10.285651Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.274807Z","title":"Zhang, B","venue":null,"work_id":"cc52fdb3-d64f-4407-a5c2-53eff435d8d3","year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.470269Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:96e0989e6d92910fa69303d30efd19a8dcd6000e3e2c493a7a5c1c3240d8c5fe","observation_id":"78f21009-35a9-4062-8d13-c75239333bbb","resolution":{"observed_at":"2026-08-15T16:09:10.277324Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.07872","last_updated":"2024-02-12T18:33:47Z","snapshot_observed_at":"2026-08-14T01:10:13.014483Z","submitted_at":"2024-02-12T18:33:47Z","title":"PIVOT: Iterative Visual Prompting Elicits Actionable Knowledge for VLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.07872","snapshot_observed_at":"2026-08-15T16:09:09.473481Z","title":"Nasiriany, F","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.473481Z"},"links":{"cited_paper":"/paper/2402.07872","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:90b60b0021e1d1b810320302c1c059ec8b0f9b89a040403a0a6ead3051039d87","observation_id":"4d0442ee-7f3e-4452-8e7e-353bd32c35f6","resolution":{"observed_at":"2026-08-15T16:09:09.473481Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02193","last_updated":"2024-10-03T04:14:21Z","snapshot_observed_at":"2026-08-13T01:35:10.442787Z","submitted_at":"2024-10-03T04:14:21Z","title":"Guiding Long-Horizon Task and Motion Planning with Vision Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.02193","snapshot_observed_at":"2026-08-15T16:09:09.476780Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.476780Z"},"links":{"cited_paper":"/paper/2410.02193","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:f2afc9ecdf8fb2339b5b3cf3c02b0c7057652ac81fb67c5db1cd4d910172caea","observation_id":"601d362f-0b3e-4ad0-8332-2785e10d9e4f","resolution":{"observed_at":"2026-08-15T16:09:09.476780Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.266398Z","title":"Chang, S","venue":null,"work_id":"47b17b47-3203-4288-bf4a-4de1fc8bb1fc","year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.479935Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:5f79d41215b81eec7c99f4df35018279c6d83308e0a93668646025459ae98820","observation_id":"dda746fd-342c-4d41-9b89-db53510c5bb0","resolution":{"observed_at":"2026-08-15T16:09:10.269544Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.257743Z","title":"Chang, S","venue":null,"work_id":"3402f882-b4df-4e78-bcc9-44364ec9b82f","year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.483013Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:816022e038e41f21423e5780cb140a0eb6027e66986bb0009fdaaeee3f8eaab9","observation_id":"a963da45-ead0-47c4-a031-b43c32c45466","resolution":{"observed_at":"2026-08-15T16:09:10.260478Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.12537","last_updated":"2023-08-24T03:47:27Z","snapshot_observed_at":"2026-08-13T10:28:37.624357Z","submitted_at":"2023-08-24T03:47:27Z","title":"HuBo-VLM: Unified Vision-Language Model designed for HUman roBOt interaction tasks","version":1},"cited_work":{"arxiv_id":"2308.12537","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.12537","snapshot_observed_at":"2026-08-15T16:09:09.943115Z","title":"HuBo-VLM: Unified Vision-Language Model designed for HUman roBOt interaction tasks","venue":"cs.RO","work_id":"963858bf-9544-4f8f-86d4-de386d5280a6","year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.485689Z"},"links":{"cited_paper":"/paper/2308.12537","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:67132795504b132cdda7497a79d5796d1da55fffe1eaeb7a9f16e69f8cfd8fd8","observation_id":"e915f67c-5c55-495e-a218-3aee85f81264","resolution":{"observed_at":"2026-08-15T16:09:09.947474Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15637","last_updated":"2024-03-22T22:27:42Z","snapshot_observed_at":"2026-08-13T20:44:10.000116Z","submitted_at":"2024-03-22T22:27:42Z","title":"CoNVOI: Context-aware Navigation using Vision Language Models in Outdoor and Indoor Environments","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.15637","snapshot_observed_at":"2026-08-15T16:09:09.488652Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.488652Z"},"links":{"cited_paper":"/paper/2403.15637","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:b35f15f34a6ad9c9a8a3cda3e277d04f881984089e711359ae54d0d692838b8e","observation_id":"ea2644f2-7191-43c9-bd58-a81a872a6bf7","resolution":{"observed_at":"2026-08-15T16:09:09.488652Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.07775","last_updated":"2024-07-12T14:37:08Z","snapshot_observed_at":"2026-08-12T23:24:55.990893Z","submitted_at":"2024-07-10T15:49:07Z","title":"Mobility VLA: Multimodal Instruction Navigation with Long-Context VLMs and Topological Graphs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.07775","snapshot_observed_at":"2026-08-15T16:09:09.491569Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.491569Z"},"links":{"cited_paper":"/paper/2407.07775","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:06fd02393a36797f300b968fd5c04ad8898d69be820bf5e56ecfdd1f877cefdf","observation_id":"3ce68ac1-087b-499d-8521-a083d36df4b1","resolution":{"observed_at":"2026-08-15T16:09:09.491569Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.16484","last_updated":"2024-10-02T19:50:54Z","snapshot_observed_at":"2026-08-12T22:38:15.520691Z","submitted_at":"2024-09-24T22:15:24Z","title":"BehAV: Behavioral Rule Guided Autonomy Using VLMs for Robot Navigation in Outdoor Scenes","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.16484","snapshot_observed_at":"2026-08-15T16:09:09.494683Z","title":"Weerakoon, M","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.494683Z"},"links":{"cited_paper":"/paper/2409.16484","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:92483b98ea23b1cd3e60cd60ca2ab09293d20bb4697e4420bce67be8b18e2f67","observation_id":"d333312d-080a-4f0f-a57a-f8b7b728bb0c","resolution":{"observed_at":"2026-08-15T16:09:09.494683Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.03603","last_updated":"2024-10-04T17:03:14Z","snapshot_observed_at":"2026-08-14T01:01:26.534012Z","submitted_at":"2024-10-04T17:03:14Z","title":"LeLaN: Learning A Language-Conditioned Navigation Policy from In-the-Wild Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.03603","snapshot_observed_at":"2026-08-15T16:09:09.498085Z","title":"Hirose, C","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.498085Z"},"links":{"cited_paper":"/paper/2410.03603","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:5563ff11c88ebe39bb3a551630bdd57afcffcf2111d1857495802d2838a92e6d","observation_id":"9cd0b670-40b6-4e2c-af3e-efc537703cbf","resolution":{"observed_at":"2026-08-15T16:09:09.498085Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.12168","last_updated":"2024-01-22T18:01:01Z","snapshot_observed_at":"2026-08-15T15:25:08.731632Z","submitted_at":"2024-01-22T18:01:01Z","title":"SpatialVLM: Endowing Vision-Language Models with Spatial Reasoning Capabilities","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.12168","snapshot_observed_at":"2026-08-15T16:09:09.501365Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.501365Z"},"links":{"cited_paper":"/paper/2401.12168","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:7f01c32b7c1c88c5d1d7f0473dd2a767374c4cd790d888ccd68c00712210dbd6","observation_id":"d3ee2306-22c3-456b-98d5-857937422e67","resolution":{"observed_at":"2026-08-15T16:09:09.501365Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:09.504246Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.504246Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:fcb9888ca8c4aeb61d423693e1e695cbe82fdc1c1ace4b10d8c1fd6fa5a08d93","observation_id":"89a97022-28bf-4161-a2ef-ded6c3b78f39","resolution":{"observed_at":"2026-08-15T16:09:09.504246Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.249309Z","title":"Helbing and P","venue":null,"work_id":"554ab082-921b-40a4-af80-6f4df2dc9d8c","year":1995},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.507014Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:b50622ee45b7c35206ebda75e4625d1f063f26b0d00869794174d0f009e3821d","observation_id":"fa04703e-6e6f-4489-be34-9c4b4af5529f","resolution":{"observed_at":"2026-08-15T16:09:10.252382Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:09.509689Z","title":"Mumm and B","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.509689Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:19131d6363fc23932bc733da115274f43521d260ef2c4f75115c7771b87ac6d2","observation_id":"c06f0127-7fd2-4568-8d28-cf44e1db2e45","resolution":{"observed_at":"2026-08-15T16:09:09.509689Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.240160Z","title":"Hirose, D","venue":null,"work_id":"eb7b9ee5-78e3-42ae-8691-70c8142b248b","year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.512211Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:be31ff4d508f2013bc83684948584838cc1f02c84819213c83f840d8a82f40b0","observation_id":"d8463d73-6f02-486d-996d-8b2e0a43ad17","resolution":{"observed_at":"2026-08-15T16:09:10.243934Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.231148Z","title":"Zhu and T","venue":null,"work_id":"0b7508c9-a16b-4e0c-8d93-b520edfd5179","year":2021},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.514737Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:52cfce6c28727f872c20871fca7763dad3f47a93a79eae0b62450ffd110e45e6","observation_id":"1fdd5333-24bc-4874-8058-9f9997c31c70","resolution":{"observed_at":"2026-08-15T16:09:10.234039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.221499Z","title":null,"venue":null,"work_id":"5fa601d5-803a-4ae8-8d7c-3ae05dc83190","year":2017},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.517444Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:58513080dae47896c21f6198cc711c93e3aedebc77ac109f9efba1007ac00ee3","observation_id":"6c481bd7-9ed3-4b18-8fcd-d1851f72381d","resolution":{"observed_at":"2026-08-15T16:09:10.225022Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.212143Z","title":null,"venue":null,"work_id":"b97a8e45-8e9f-46fe-aec4-15552cee4d66","year":2019},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.520178Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:b5d1fbc87c0af2690c6a9fa8484360b2aa824a5f271ed7ee10d6d741cadcc6c2","observation_id":"51229579-36ac-41cf-ac1e-cb54ed388937","resolution":{"observed_at":"2026-08-15T16:09:10.215425Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.203426Z","title":null,"venue":null,"work_id":"9479c7df-9b0e-41cf-b684-457737575640","year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.522973Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:db9f5525261131354259192c5d873565fbd757b160b46af12a19cf00dc67d96e","observation_id":"8d8ea0fc-275c-44fc-86b4-c7be98603cb8","resolution":{"observed_at":"2026-08-15T16:09:10.206057Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:09.525808Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.525808Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:38d908c6ceda2de4e30e36a8b02bb726387386c660950dd873033a1a0e46560c","observation_id":"f4f0d0a0-3120-4bcd-89aa-8dee6699e61f","resolution":{"observed_at":"2026-08-15T16:09:09.525808Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.194548Z","title":null,"venue":null,"work_id":"85c8a36b-da55-4d9f-8494-5f979a8fe266","year":null},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.528785Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:f7ff1d063a5db018510d8e9b607338149238ea31b1e3551306613de002bbb8f4","observation_id":"fd844bbb-920d-427a-9379-d233d16c1814","resolution":{"observed_at":"2026-08-15T16:09:10.197515Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:09.531559Z","title":"Hirose, D","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.531559Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:eda767fad2aad1358ec5724d17f97b643f857741694b4ee8acda504558930979","observation_id":"80e85337-c56c-435c-9b07-671a0c1d66a9","resolution":{"observed_at":"2026-08-15T16:09:09.531559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-08-15T16:09:09.535002Z","title":"Payandeh, D","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.535002Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:2ad9712c16c979ff37954e4940356485f8795dafd9806551c9c7f56cca0cb6b4","observation_id":"6012a5e8-fb67-4766-b76b-7354c207465f","resolution":{"observed_at":"2026-08-15T16:09:09.535002Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.13682","last_updated":"2024-09-20T17:50:07Z","snapshot_observed_at":"2026-08-12T22:40:48.823902Z","submitted_at":"2024-09-20T17:50:07Z","title":"ReMEmbR: Building and Reasoning Over Long-Horizon Spatio-Temporal Memory for Robot Navigation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.13682","snapshot_observed_at":"2026-08-15T16:09:09.537952Z","title":"Anwar, J","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.537952Z"},"links":{"cited_paper":"/paper/2409.13682","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:95df36d8b0d88a5dced178fba96d97cb44e2af002ee6f51bfc4bfa4ab8d03cd8","observation_id":"87ff0b86-196d-4309-9e13-314e84b87190","resolution":{"observed_at":"2026-08-15T16:09:09.537952Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.184852Z","title":"Zhang, C","venue":null,"work_id":"011b9c54-2fcb-45b3-90e0-b6b0568c009f","year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.540949Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:df1a163485875c1ae330c29dcc46c40f05c7206066632d79f728c605f293312b","observation_id":"b9cf18b3-78fc-4394-8b97-4d9947d560d2","resolution":{"observed_at":"2026-08-15T16:09:10.187874Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09167","last_updated":"2025-01-15T21:36:19Z","snapshot_observed_at":"2026-08-15T14:06:19.800460Z","submitted_at":"2025-01-15T21:36:19Z","title":"Embodied Scene Understanding for Vision Language Models via MetaVQA","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09167","snapshot_observed_at":"2026-08-15T16:09:09.543760Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.543760Z"},"links":{"cited_paper":"/paper/2501.09167","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:3a51557b8a285926ffdf8b96190f64c78d10ad647f1824855d47bc6d6eb0167f","observation_id":"b2f2e624-7885-4af5-81a5-fe6d55540101","resolution":{"observed_at":"2026-08-15T16:09:09.543760Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.05956","last_updated":"2024-10-25T19:31:37Z","snapshot_observed_at":"2026-08-13T00:10:47.548058Z","submitted_at":"2024-05-09T17:52:42Z","title":"Probing Multimodal LLMs as World Models for Driving","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.05956","snapshot_observed_at":"2026-08-15T16:09:09.546597Z","title":"Sreeram, T.-H","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.546597Z"},"links":{"cited_paper":"/paper/2405.05956","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:4d56022ba43208b3152370445ec08289e1bb9f0f04df08e63516c784bb5ed865","observation_id":"ac0f48aa-5854-459d-94a0-07b368231f2e","resolution":{"observed_at":"2026-08-15T16:09:09.546597Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.174926Z","title":"Chow*, J","venue":null,"work_id":"c2350710-0d97-4c35-ae7b-a3f299284eb9","year":2025},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.549885Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:b3b52b1fab7d9e35502af5ed2186e1e5c87e5858deba4765585527e6268b7d16","observation_id":"23449f9c-1257-4a6d-9b06-e633ddb65d8b","resolution":{"observed_at":"2026-08-15T16:09:10.177665Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.11441","last_updated":"2023-11-06T07:39:49Z","snapshot_observed_at":"2026-08-12T21:25:32.312122Z","submitted_at":"2023-10-17T17:51:31Z","title":"Set-of-Mark Prompting Unleashes Extraordinary Visual Grounding in GPT-4V","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.11441","snapshot_observed_at":"2026-08-15T16:09:09.553196Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.553196Z"},"links":{"cited_paper":"/paper/2310.11441","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:8775a9d7566b75a3cd51a92c18dd39e3dded3234593fe7a500609175478e891d","observation_id":"734bc7b8-7c67-4efd-86e5-5288670a3688","resolution":{"observed_at":"2026-08-15T16:09:09.553196Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.164854Z","title":"Rajasegaran, G","venue":null,"work_id":"766cf6ed-c9f5-496e-a357-e220124c2f3a","year":2022},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.556022Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:96a4bb4b33a418cd04648b4378a9f0635d74c128cb79d32b2ccc3036bd5a1f28","observation_id":"4ae1749d-d475-474b-b47d-5f17c7fda854","resolution":{"observed_at":"2026-08-15T16:09:10.168523Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.03921","last_updated":"2025-06-26T06:42:04Z","snapshot_observed_at":"2026-08-07T17:26:27.344738Z","submitted_at":"2025-03-05T21:42:46Z","title":"CREStE: Scalable Mapless Navigation with Internet Scale Priors and Counterfactual Guidance","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.03921","snapshot_observed_at":"2026-08-15T16:09:09.558906Z","title":"Zhang, H","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.558906Z"},"links":{"cited_paper":"/paper/2503.03921","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:3f897df3d7351c2fd478316379aa75af42af1c14e932b43c10afbe765da74c1e","observation_id":"6ae716ce-d9ea-469b-a809-a75646e658af","resolution":{"observed_at":"2026-08-15T16:09:09.558906Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.155200Z","title":"Tadic, A","venue":null,"work_id":"7f37f79c-a12e-4d7a-bee7-63126c7c6815","year":2022},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.561984Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:96518457bdf55b644eba3143c01ab317a2a8ad5b1c73c00bd022216bc3d07588","observation_id":"ea3826ca-6b2a-448c-9d62-ff031293c203","resolution":{"observed_at":"2026-08-15T16:09:10.158595Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.146110Z","title":"Aharony, A","venue":null,"work_id":"01a996cc-71d2-4a92-ab8e-a68320cd9311","year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.564967Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:06193cd7eca8fbaab0558400d9301b57b0053fe140aa53fcdf4cfd1a77e02ff6","observation_id":"aac568b5-1052-4584-9b1a-9bb4db5b7ba1","resolution":{"observed_at":"2026-08-15T16:09:10.149038Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:10.135524Z","title":"Prolific.https://www.prolific.com, 2014.Accessed on [date accessed]","venue":null,"work_id":"95164929-0336-4335-9c9e-d1d5a87cf020","year":2014},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.567758Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:73022dab275d9355b5e211743ae752e282f82849e9b97deecf6530e8710c491a","observation_id":"d00309d1-a46e-4af6-967a-91d8a4fe1937","resolution":{"observed_at":"2026-08-15T16:09:10.139095Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T16:09:09.570918Z","title":"Antol, A","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.570918Z"},"links":{"citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:f456fd6215f4369b2ee88520cb72437ad04c10a21a5cb8db93cf85a28c2f884f","observation_id":"9616f61f-1b87-4fa3-91c1-7fa11f57e707","resolution":{"observed_at":"2026-08-15T16:09:09.570918Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-08-15T12:33:55.451951Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-15T16:09:09.573838Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.573838Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:52355d4853b8f186d89277c3f5270384b763d949c91f452146ac7d1fe33a02e8","observation_id":"cf365ace-2bfc-4fa6-9aa0-1e64b1e83e24","resolution":{"observed_at":"2026-08-15T16:09:09.573838Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01584","last_updated":"2024-10-15T01:16:20Z","snapshot_observed_at":"2026-08-16T01:32:18.967483Z","submitted_at":"2024-06-03T17:59:06Z","title":"SpatialRGPT: Grounded Spatial Reasoning in Vision Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01584","snapshot_observed_at":"2026-08-15T16:09:09.577056Z","title":"avoiding","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-15T16:09:09.577056Z"},"links":{"cited_paper":"/paper/2406.01584","citing_paper":"/paper/2509.08757"},"observation_digest":"sha256:155507b0d91398856000b448ccbce1d5e84df3fb6ae3b0bf6daa006d912bbaff","observation_id":"701aa8ba-57e0-404b-aacb-18f69344c1be","resolution":{"observed_at":"2026-08-15T16:09:09.577056Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2509.08757","last_updated":"2025-09-10T16:47:00Z","latest_version":1,"primary_category":"cs.RO","snapshot_observed_at":"2026-08-15T15:57:48.357122Z","submitted_at":"2025-09-10T16:47:00Z","title":"SocialNav-SUB: Benchmarking VLMs for Scene Understanding in Social Robot Navigation"},"reference_resolution":{"displayed":53,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":32,"verified_exact":3,"verified_fuzzy":17},"total_outbound_references":53},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 53 of 53 outbound references and 2 inbound Pith citation observations for arXiv:2509.08757."}