{"as_of":"2026-08-09T20:06:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:70ad16670ac1c0d4c1deeeeb44f7d792b27db7a0d49abcf356b45e183d45839f","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":55,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":55,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":55,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":55,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T17:12:16.596574Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T15:39:56.857028Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2412.14171","last_updated":"2025-07-02T21:00:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-18T18:59:54Z","title":"Thinking in Space: How Multimodal Large Language Models See, Remember, and Recall Spaces","version":2},"reference_index":107,"source":"pdf_text","source_observed_at":"2026-05-22T09:27:43.919941Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2412.14171"},"observation_digest":"sha256:ebd24fb50e30cff9549e3df51b646e12171f67aabc41b1c3b51fb076aea0f8c8","observation_id":"ea7272f7-bad2-4ed9-978f-d9fab1682bfe","resolution":{"observed_at":"2026-05-22T09:27:44.102329Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2501.15830","last_updated":"2025-05-19T02:40:18Z","snapshot_observed_at":"2026-07-06T20:26:31.558337Z","submitted_at":"2025-01-27T07:34:33Z","title":"SpatialVLA: Exploring Spatial Representations for Visual-Language-Action Model","version":5},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-05-12T06:12:19.643111Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2501.15830"},"observation_digest":"sha256:b7ec4b22de324bf45b1f7206138d94413a6be5ce27a3b20ea5a8a4d5d720170b","observation_id":"fbe72696-13d7-474c-bec0-7d43145b806e","resolution":{"observed_at":"2026-05-12T06:12:19.853033Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-09T17:12:16.596574Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.00954","last_updated":"2025-05-27T19:23:41Z","snapshot_observed_at":"2026-08-09T17:04:36.999228Z","submitted_at":"2025-02-02T23:11:42Z","title":"Hypo3D: Exploring Hypothetical Reasoning in 3D","version":3},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-09T17:12:16.596574Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2502.00954"},"observation_digest":"sha256:c7fe763d3ebf38a0e5681b503b106f4bce0d50cb0cd92dbcd40fe11bcc148075","observation_id":"087e4f40-97e0-4330-ac64-ee25e0a78d3b","resolution":{"observed_at":"2026-08-09T17:12:16.596574Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-08T04:54:02.129840Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.08503","last_updated":"2025-06-06T01:51:44Z","snapshot_observed_at":"2026-08-08T14:13:09.142446Z","submitted_at":"2025-02-12T15:34:45Z","title":"Revisiting 3D LLM Benchmarks: Are We Really Testing 3D Capabilities?","version":3},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-08T04:54:02.129840Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2502.08503"},"observation_digest":"sha256:871e1dd9b2da62e59d76b1d5f8c878713523a9a6edf363f26825e0720845b90b","observation_id":"923e75a7-626e-4223-bbb4-39188fb0d440","resolution":{"observed_at":"2026-08-08T04:54:02.129840Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-07T13:46:14.423259Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21079","last_updated":"2025-05-27T12:03:30Z","snapshot_observed_at":"2026-08-08T23:14:14.149778Z","submitted_at":"2025-05-27T12:03:30Z","title":"Uni3D-MoE: Scalable Multimodal 3D Scene Understanding via Mixture of Experts","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T13:46:14.423259Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2505.21079"},"observation_digest":"sha256:bb5ca8b2136f0054c54d2ac171c035610cde1056db10069233f05e006bc9cbdc","observation_id":"b26fd254-490e-46e6-b794-35981a26791a","resolution":{"observed_at":"2026-08-07T13:46:14.423259Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2505.23747","last_updated":"2026-05-19T02:23:16Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-29T17:59:04Z","title":"Spatial-MLLM: Boosting MLLM Capabilities in Visual-based Spatial Intelligence","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-16T08:34:36.824053Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2505.23747"},"observation_digest":"sha256:271d7253ccd60ec94ed3acb576e81be4f898c9673690b04246204cfa70a03c8d","observation_id":"6ac13126-0ca5-4894-912f-2fb74e07e39c","resolution":{"observed_at":"2026-05-16T08:34:36.868541Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2505.23747","last_updated":"2026-05-19T02:23:16Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-29T17:59:04Z","title":"Spatial-MLLM: Boosting MLLM Capabilities in Visual-based Spatial Intelligence","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-22T00:59:13.826054Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2505.23747"},"observation_digest":"sha256:28b9540378ef9909e161abbca50f598e6c08b192d072b83cf04e2779aa17e154","observation_id":"506780a9-b225-41c3-82f0-71e1a9ab68f3","resolution":{"observed_at":"2026-05-22T01:00:51.268241Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-07T12:35:20.069930Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24257","last_updated":"2025-05-30T06:32:26Z","snapshot_observed_at":"2026-08-09T16:57:27.213555Z","submitted_at":"2025-05-30T06:32:26Z","title":"Out of Sight, Not Out of Context? Egocentric Spatial Reasoning in VLMs Across Disjoint Frames","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:20.069930Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2505.24257"},"observation_digest":"sha256:5e0e40c693313121b292296ff4bd973bd966b8046de8347b46cc6980ff1b1873","observation_id":"e9eba971-990b-4037-98c2-afd389d58a14","resolution":{"observed_at":"2026-08-07T12:35:20.069930Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-07T12:16:16.730134Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125 , 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.00123","last_updated":"2025-05-30T18:00:34Z","snapshot_observed_at":"2026-08-09T02:39:20.854550Z","submitted_at":"2025-05-30T18:00:34Z","title":"Visual Embodied Brain: Let Multimodal Large Language Models See, Think, and Control in Spaces","version":1},"reference_index":111,"source":"pdf_text","source_observed_at":"2026-08-07T12:16:16.730134Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2506.00123"},"observation_digest":"sha256:d50ce22833901dc0fb794552bfd62e1f9e45ae0dc978dca031b2df3e1a4d7f36","observation_id":"7046edbd-9db3-4a06-8615-8c6489c47695","resolution":{"observed_at":"2026-08-07T12:16:16.730134Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-07T10:25:48.356054Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.05318","last_updated":"2025-06-06T07:09:06Z","snapshot_observed_at":"2026-08-08T07:19:59.563841Z","submitted_at":"2025-06-05T17:56:12Z","title":"Does Your 3D Encoder Really Work? When Pretrain-SFT from 2D VLMs Meets 3D VLMs","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T10:25:48.356054Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2506.05318"},"observation_digest":"sha256:debe0942921d918ff972b7a2af011fcf702c480e60c00e80202d626d57cc1eee","observation_id":"264df855-651c-4abf-a942-66b432d95dc4","resolution":{"observed_at":"2026-08-07T10:25:48.356054Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-07T10:50:52.588064Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.05414","last_updated":"2025-06-04T19:11:20Z","snapshot_observed_at":"2026-08-09T06:09:58.091457Z","submitted_at":"2025-06-04T19:11:20Z","title":"SAVVY: Spatial Awareness via Audio-Visual LLMs through Seeing and Hearing","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T10:50:52.588064Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2506.05414"},"observation_digest":"sha256:f3821ccee3460380f77a2001bf8ff0a4e112fbb8d4a4c2a2be237b0ee366bef2","observation_id":"7ab239f0-022e-42d6-bc5d-4a37fbd8672b","resolution":{"observed_at":"2026-08-07T10:50:52.588064Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-07T10:20:37.480142Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.05689","last_updated":"2025-06-06T02:35:26Z","snapshot_observed_at":"2026-08-09T04:48:26.376897Z","submitted_at":"2025-06-06T02:35:26Z","title":"Pts3D-LLM: Studying the Impact of Token Structure for 3D Scene Understanding With Large Language Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T10:20:37.480142Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2506.05689"},"observation_digest":"sha256:496e10549e3017e4d5cb5dd9aa9421432d9d0fa4244798a203c00e1fc170d323","observation_id":"50cb826f-7ff1-4bd9-9c9f-50f230d11047","resolution":{"observed_at":"2026-08-07T10:20:37.480142Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-07T00:57:27.599419Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12374","last_updated":"2025-06-24T10:01:18Z","snapshot_observed_at":"2026-08-07T07:38:26.007252Z","submitted_at":"2025-06-14T07:11:44Z","title":"AntiGrounding: Lifting Robotic Actions into VLM Representation Space for Decision Making","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T00:57:27.599419Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2506.12374"},"observation_digest":"sha256:c44618e013669a02ebc7e8ee94cb9c062e397fe82d8a4e355a5ab6b4565d44c9","observation_id":"c8499945-4b1d-4d4d-89e3-eaac8adcd815","resolution":{"observed_at":"2026-08-07T00:57:27.599419Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-06T21:52:12.924934Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.23329","last_updated":"2025-06-29T17:02:57Z","snapshot_observed_at":"2026-08-07T17:07:15.412812Z","submitted_at":"2025-06-29T17:02:57Z","title":"IR3D-Bench: Evaluating Vision-Language Model Scene Understanding as Agentic Inverse Rendering","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T21:52:12.924934Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2506.23329"},"observation_digest":"sha256:fd23226bc5e944e43fbb23eb491d7aff54f7b598feccbab11af820720a3593ca","observation_id":"baa76daf-d0cc-4046-867a-c8c4e18cb39f","resolution":{"observed_at":"2026-08-06T21:52:12.924934Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-06T10:49:47.146212Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.23478","last_updated":"2025-07-31T11:59:06Z","snapshot_observed_at":"2026-08-07T03:56:49.161093Z","submitted_at":"2025-07-31T11:59:06Z","title":"3D-R1: Enhancing Reasoning in 3D VLMs for Unified Scene Understanding","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-06T10:49:47.146212Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2507.23478"},"observation_digest":"sha256:b9c1c880a25105363731379b6980bbbdb77d716e905405a17bdaa74b106d6b8e","observation_id":"b4048b41-fcc1-4dee-a24f-dc6722516e0a","resolution":{"observed_at":"2026-08-06T10:49:47.146212Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-05T22:18:05.150726Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.07253","last_updated":"2025-08-10T09:12:29Z","snapshot_observed_at":"2026-08-08T15:34:10.323551Z","submitted_at":"2025-08-10T09:12:29Z","title":"PySeizure: A single machine learning classifier framework to detect seizures in diverse datasets","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T22:18:05.150726Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2508.07253"},"observation_digest":"sha256:a127ff4062413f8e50d948d3a158e4c61531bba0bb4a23817e24406c8df06118","observation_id":"3cb81618-21a3-405b-8dac-28d034d2a138","resolution":{"observed_at":"2026-08-05T22:18:05.150726Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-05T21:16:50.649376Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.09071","last_updated":"2025-08-13T16:47:50Z","snapshot_observed_at":"2026-08-07T16:52:12.636532Z","submitted_at":"2025-08-12T16:46:05Z","title":"GeoVLA: Empowering 3D Representations in Vision-Language-Action Models","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T21:16:50.649376Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2508.09071"},"observation_digest":"sha256:ca589589bf2adc933c948b97e6604eeab404ad949b61738d020336bb1b3a47e9","observation_id":"8fb6f571-f63e-4466-985d-9a52c58cb6d2","resolution":{"observed_at":"2026-08-05T21:16:50.649376Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-04T11:38:13.467839Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.03896","last_updated":"2026-06-11T03:17:22Z","snapshot_observed_at":"2026-08-09T06:57:01.817389Z","submitted_at":"2025-10-04T18:33:27Z","title":"GAE: Unleashing Physical Potential of VLM with Generalizable Action Expert","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-04T11:38:13.467839Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2510.03896"},"observation_digest":"sha256:b2701cc55f32c4fbe601d08e51bc20c1053e1b2f25768c41d33433b4d333a7d8","observation_id":"d947ec6d-b7b5-4312-a67d-8132fdde69d8","resolution":{"observed_at":"2026-08-04T11:38:13.467839Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2511.16567","last_updated":"2026-05-06T15:47:48Z","snapshot_observed_at":"2026-07-06T22:36:30.947164Z","submitted_at":"2025-11-20T17:22:51Z","title":"POMA-3D: The Point Map Way to 3D Scene Understanding","version":3},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-17T20:27:27.347592Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2511.16567"},"observation_digest":"sha256:02392879f77c493790cd8b3949354c7a1d44f5c6249417493aca4c20a31b8b11","observation_id":"8b515b32-dab1-44d5-ace8-129bf2380628","resolution":{"observed_at":"2026-05-17T20:30:11.559924Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2511.19972","last_updated":"2026-05-07T02:58:13Z","snapshot_observed_at":"2026-08-06T20:13:12.622191Z","submitted_at":"2025-11-25T06:31:57Z","title":"Boosting Reasoning in Large Multimodal Models via Activation Replay","version":3},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-05-17T05:05:48.682057Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2511.19972"},"observation_digest":"sha256:0ff7ed43a64c0df8b912f2139ea95cb14b6528da67643aa95dfeed34ac29b303","observation_id":"9ece6334-1468-428b-868d-cd6c9d1295ac","resolution":{"observed_at":"2026-05-17T05:09:03.922650Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-03T20:15:37.516145Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2511.20644","last_updated":"2026-07-09T17:59:10Z","snapshot_observed_at":"2026-08-09T00:51:01.709796Z","submitted_at":"2025-11-25T18:59:02Z","title":"Vision-Language Memory for Spatial Reasoning","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-03T20:15:37.516145Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2511.20644"},"observation_digest":"sha256:ba9e730ba4d3899b61562d1eda46de1ae71bbe9f613253c5cacb4ccb362c2725","observation_id":"7ea85dfe-8b0f-466b-b79d-b8a9c07fb2f3","resolution":{"observed_at":"2026-08-03T20:15:37.516145Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2511.21471","last_updated":"2026-05-07T07:59:46Z","snapshot_observed_at":"2026-08-02T23:27:20.204280Z","submitted_at":"2025-11-26T15:04:18Z","title":"SpatialBench: Benchmarking Multimodal Large Language Models for Spatial Cognition","version":4},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-05-17T04:54:59.903644Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2511.21471"},"observation_digest":"sha256:6cb5cb7a211259e7aa7d7d11bfdb1fb515bd88dc5d41ebd596787d35d8f317da","observation_id":"6e8a4b7d-4e8b-426a-888d-96b9dceb4bfa","resolution":{"observed_at":"2026-05-17T04:59:04.240779Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2512.10719","last_updated":"2026-07-12T18:02:23Z","snapshot_observed_at":"2026-08-03T17:07:38.286115Z","submitted_at":"2025-12-11T14:59:07Z","title":"SpaceDrive: Infusing Spatial Awareness into VLM-based Autonomous Driving","version":2},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-05-22T12:17:54.325055Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2512.10719"},"observation_digest":"sha256:7c8a152b797aa2bd5b5a164aa8ab3aeb149fd81a32204b4a8a064b91d063b45c","observation_id":"fef0e861-6800-4855-80ee-a03f7a27a63b","resolution":{"observed_at":"2026-05-22T12:21:31.269741Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-03T17:07:45.189496Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2512.10719","last_updated":"2026-07-12T18:02:23Z","snapshot_observed_at":"2026-08-03T17:07:38.286115Z","submitted_at":"2025-12-11T14:59:07Z","title":"SpaceDrive: Infusing Spatial Awareness into VLM-based Autonomous Driving","version":3},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-03T17:07:45.189496Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2512.10719"},"observation_digest":"sha256:5213c3e8e199f1662984e428d6950ffbb6ccb8bc1e08a884ac76168df86b6728","observation_id":"0c2d3562-4638-4175-90ac-aafe8eccde47","resolution":{"observed_at":"2026-08-03T17:07:45.189496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-03T13:46:30.457026Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2512.23020","last_updated":"2026-07-07T09:58:55Z","snapshot_observed_at":"2026-08-06T21:09:09.898549Z","submitted_at":"2025-12-28T17:44:20Z","title":"OpenGround: Planning-based Online Perception for Open-World 3D Visual Grounding","version":3},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-03T13:46:30.457026Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2512.23020"},"observation_digest":"sha256:6225cb2c9729161b2fc72e49b728febede19b9db5e0ece3358f6799d500a8d01","observation_id":"725628d7-5b50-4f3c-9e4c-88091d41cc87","resolution":{"observed_at":"2026-08-03T13:46:30.457026Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2603.01400","last_updated":"2026-04-09T18:07:09Z","snapshot_observed_at":"2026-08-07T14:16:18.120869Z","submitted_at":"2026-03-02T03:06:40Z","title":"Token Reduction via Local and Global Contexts Optimization for Efficient Video Large Language Models","version":2},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-05-15T18:25:21.621268Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2603.01400"},"observation_digest":"sha256:302ffa94c87008c4561b4fd1f5e0c073369a12a42263df0dcf6a0a733be95d3f","observation_id":"bf521ba6-1604-4758-b59a-e5446d40d410","resolution":{"observed_at":"2026-05-15T18:26:26.916148Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2603.08592","last_updated":"2026-04-26T17:36:17Z","snapshot_observed_at":"2026-07-31T15:55:15.225728Z","submitted_at":"2026-03-09T16:42:43Z","title":"Boosting MLLM Spatial Reasoning with Geometrically Referenced 3D Scene Representations","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-15T14:31:03.909336Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2603.08592"},"observation_digest":"sha256:3ca99a2758da40454fe74c85aacef3e57962aa8b00baccec8f9dc437a4f51cb7","observation_id":"46529768-f4c3-4547-991f-32b41b4863c2","resolution":{"observed_at":"2026-05-15T14:35:56.050224Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-02T18:06:09.995051Z","title":"arXiv preprint arXiv:2409.18125 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.16461","last_updated":"2026-07-20T02:21:38Z","snapshot_observed_at":"2026-08-08T13:28:18.188071Z","submitted_at":"2026-03-17T12:43:48Z","title":"GAP-MLLM: Geometry-Aligned Pre-training for Activating 3D Spatial Perception in Multimodal Large Language Models","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-02T18:06:09.995051Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2603.16461"},"observation_digest":"sha256:48bed314e6cf69574356e0045cc40fdf7dc8acd6dcc1ef2bc53afa37ef02c53b","observation_id":"78606104-cf6a-4120-8cd3-97f27074cf27","resolution":{"observed_at":"2026-08-02T18:06:09.995051Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-13T23:42:44.515158Z","title":"arXiv preprint arXiv:2409.18125 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.16506","last_updated":"2026-03-18T04:22:15Z","snapshot_observed_at":"2026-08-08T08:49:07.375086Z","submitted_at":"2026-03-17T13:36:30Z","title":"VIEW2SPACE: Studying Multi-View Visual Reasoning from Sparse Observations","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-07-13T23:42:44.515158Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2603.16506"},"observation_digest":"sha256:832a14ca1397ce0984f6b3e88cc5d7a617ea55d401bbc1a61d4e2e729fa41d59","observation_id":"57ad9ea2-ad06-4ac9-a736-0307cecf070c","resolution":{"observed_at":"2026-07-13T23:42:44.515158Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2603.17980","last_updated":"2026-05-07T05:29:31Z","snapshot_observed_at":"2026-07-06T22:49:37.944352Z","submitted_at":"2026-03-18T17:42:49Z","title":"Feeling the Space: Egomotion-Aware Video Representation for Efficient and Accurate 3D Scene Understanding","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-05-15T09:30:03.668178Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2603.17980"},"observation_digest":"sha256:1ea0f09601acf1d8caf7d95d778b5c5cdc50d3a87e6b015e64b61afabc22185b","observation_id":"829037dc-1aad-4ab7-8b1b-df4face8d24e","resolution":{"observed_at":"2026-05-15T09:30:22.395942Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2603.23404","last_updated":"2026-04-20T17:18:17Z","snapshot_observed_at":"2026-07-06T22:50:23.931381Z","submitted_at":"2026-03-24T16:38:09Z","title":"Unleashing Spatial Reasoning in Multimodal Large Language Models via Textual Representation Guided Reasoning","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-15T00:08:22.362342Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2603.23404"},"observation_digest":"sha256:68e8e5f75e298dd1d05c6e87060271c1db1c0f4277b16a33b63de2efe49f1e15","observation_id":"c438ce0f-c7a5-40bd-9c12-975277b5176a","resolution":{"observed_at":"2026-05-15T00:09:35.116047Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2604.03296","last_updated":"2026-03-28T00:54:19Z","snapshot_observed_at":"2026-07-06T22:52:29.705312Z","submitted_at":"2026-03-28T00:54:19Z","title":"3D-IDE: 3D Implicit Depth Emergent","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-14T22:34:04.833557Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2604.03296"},"observation_digest":"sha256:d60e1cd7f18cc08ed950fc2854d7e2c5130e08dd909fbb71d93d7ef10f11c585","observation_id":"0691fced-5885-4c15-adcb-4003ed9e02a6","resolution":{"observed_at":"2026-05-14T22:38:11.453485Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2604.03318","last_updated":"2026-05-25T15:52:36Z","snapshot_observed_at":"2026-07-13T14:39:49.573224Z","submitted_at":"2026-04-01T15:28:13Z","title":"EgoMind: Activating Spatial Cognition through Linguistic Reasoning in MLLMs","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-05-13T22:41:09.840792Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2604.03318"},"observation_digest":"sha256:e7e7f2765fc6063002d24fb1250c96db48bd1f7c288eb7ad6acbcc6ecf735c53","observation_id":"3a769354-f6c4-494b-89d9-0e9f76277df1","resolution":{"observed_at":"2026-05-13T22:43:22.828897Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-13T14:39:55.177552Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2604.03318","last_updated":"2026-05-25T15:52:36Z","snapshot_observed_at":"2026-07-13T14:39:49.573224Z","submitted_at":"2026-04-01T15:28:13Z","title":"EgoMind: Activating Spatial Cognition through Linguistic Reasoning in MLLMs","version":2},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-07-13T14:39:55.177552Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2604.03318"},"observation_digest":"sha256:fe46d6eccb26519b5b644f1b3ec5fdaf6aa800d8098a0f8ac0bd5092baf61f4c","observation_id":"e8389b68-bbd3-450c-a0e2-ab92616c3637","resolution":{"observed_at":"2026-07-13T14:39:55.177552Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2604.17472","last_updated":"2026-04-19T14:53:38Z","snapshot_observed_at":"2026-07-06T23:04:32.734175Z","submitted_at":"2026-04-19T14:53:38Z","title":"UniMesh: Unifying 3D Mesh Understanding and Generation","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-05-10T06:55:42.679323Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2604.17472"},"observation_digest":"sha256:6fef92661b1f79d35ba04632de9cffdad26c46e57c83f10296502997aef736db","observation_id":"2b917ffe-6cf8-440d-804d-e651014702dc","resolution":{"observed_at":"2026-05-10T06:56:47.338138Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2605.09218","last_updated":"2026-05-09T23:35:27Z","snapshot_observed_at":"2026-08-03T01:38:17.494762Z","submitted_at":"2026-05-09T23:35:27Z","title":"Flame3D: Zero-shot Compositional Reasoning of 3D Scenes with Agentic Language Models","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-05-12T02:09:09.684786Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2605.09218"},"observation_digest":"sha256:83b3d8930ea6797d885c7c5a51e2d6ef6be9731d52c8d99c97734348a48b2e91","observation_id":"a819ce1a-ddab-48f3-8388-30975d06eab5","resolution":{"observed_at":"2026-05-12T02:11:15.854529Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2605.09719","last_updated":"2026-05-10T19:38:29Z","snapshot_observed_at":"2026-08-02T06:20:11.964388Z","submitted_at":"2026-05-10T19:38:29Z","title":"Distilling 3D Spatial Reasoning into a Lightweight Vision-Language Model with CoT","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-12T02:42:09.922300Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2605.09719"},"observation_digest":"sha256:c9def1d2c2624b887a472dff00e313dfbf627de1b598d7968f10220662f692a4","observation_id":"05ddfdf3-00fb-458a-aa00-d13e6dd877d0","resolution":{"observed_at":"2026-05-12T07:31:25.837388Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2605.15876","last_updated":"2026-05-20T09:56:58Z","snapshot_observed_at":"2026-08-08T12:17:20.913823Z","submitted_at":"2026-05-15T11:54:17Z","title":"Unlocking Dense Metric Depth Estimation in VLMs","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-05-20T19:20:04.468206Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2605.15876"},"observation_digest":"sha256:99694ecc123025b95562143ef1eb49d903662f6ffb39078c2286b70a8770c3cd","observation_id":"b2e9ba17-1ab5-4936-a7cf-ef96753ab33d","resolution":{"observed_at":"2026-05-20T19:23:40.971811Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2605.15876","last_updated":"2026-05-20T09:56:58Z","snapshot_observed_at":"2026-08-08T12:17:20.913823Z","submitted_at":"2026-05-15T11:54:17Z","title":"Unlocking Dense Metric Depth Estimation in VLMs","version":3},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-05-21T07:54:52.926995Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2605.15876"},"observation_digest":"sha256:3d51fd15ec1fff36ae4eef3f2469d3f95ab6fe1f13b8cbd63fd9a0060505e553","observation_id":"c7c1b832-f5bc-48bb-9d9c-d284f8a40baf","resolution":{"observed_at":"2026-05-21T07:59:51.117080Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2605.20165","last_updated":"2026-05-19T17:50:25Z","snapshot_observed_at":"2026-07-06T23:30:49.211483Z","submitted_at":"2026-05-19T17:50:25Z","title":"CaMo: Camera Motion Grounded Evaluation and Training for Vision-Language Models","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-05-20T05:27:30.938311Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2605.20165"},"observation_digest":"sha256:5bf08c2393f2365c80af08c1db87e258d71693f14c170e53fbaaa79b03e031fe","observation_id":"c568eb3c-3fee-4e39-bcc1-a6c77c8e4a7e","resolution":{"observed_at":"2026-05-20T05:28:04.610018Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2605.24642","last_updated":"2026-05-23T16:18:41Z","snapshot_observed_at":"2026-07-30T07:41:42.487489Z","submitted_at":"2026-05-23T16:18:41Z","title":"Understanding the Impact of Geometric Foundation Models on Vision-Language-Action Models","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-06-30T13:40:59.330788Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2605.24642"},"observation_digest":"sha256:77744f80fc21c23fac8d9891518e2fb130f7fe45759366e9d82275472bdc7b85","observation_id":"d8ccdd82-7f95-41e5-8e78-e71a7b07fca7","resolution":{"observed_at":"2026-06-30T13:44:40.748826Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2605.25326","last_updated":"2026-05-25T01:16:19Z","snapshot_observed_at":"2026-08-07T14:40:29.243934Z","submitted_at":"2026-05-25T01:16:19Z","title":"Perceive-then-Plan: Layout-as-Policy for Monocular 3D Scene Layout Estimation","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-06-29T23:11:59.282627Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2605.25326"},"observation_digest":"sha256:cc89a387ccb8361733dfe196213ec744a17a1c104ac283e6625ed11f2223522d","observation_id":"4350b7bb-d452-4ce5-aed8-9627f153003b","resolution":{"observed_at":"2026-06-29T23:14:01.279614Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2605.26230","last_updated":"2026-05-25T18:01:05Z","snapshot_observed_at":"2026-08-08T08:15:58.432143Z","submitted_at":"2026-05-25T18:01:05Z","title":"Geometry-Aware Representation Denoising for Robust Multi-view 3D Reconstruction","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-06-29T23:08:52.333329Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2605.26230"},"observation_digest":"sha256:51b0b896ed0ff189c4aa46542a8a3b478da8682642e77c4469afa26cb85110ab","observation_id":"19ace1ec-6239-45a8-9243-d82e6f3fa22b","resolution":{"observed_at":"2026-06-29T23:14:01.824595Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2605.30231","last_updated":"2026-05-28T17:00:52Z","snapshot_observed_at":"2026-07-06T23:39:34.232661Z","submitted_at":"2026-05-28T17:00:52Z","title":"Beyond 3D VQAs: Injecting 3D Spatial Priors into Vision-Language Models for Enhanced Geometric Reasoning","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-06-29T07:47:52.739735Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2605.30231"},"observation_digest":"sha256:4d556c6fdc45378d01256ade261ca17f71adf236fc1b5b66090761b0fcfd859b","observation_id":"c55ddd02-4122-41fc-81b9-b0161d46f9b0","resolution":{"observed_at":"2026-06-29T07:53:13.622650Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2606.06476","last_updated":"2026-06-04T17:56:36Z","snapshot_observed_at":"2026-07-06T23:46:20.554117Z","submitted_at":"2026-06-04T17:56:36Z","title":"Thinking with Imagination: Agentic Visual Spatial Reasoning with World Simulators","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-28T02:25:30.998989Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2606.06476"},"observation_digest":"sha256:d8b399aec4eec61812656309aebb7603eb89b06ee51da5bd2ef555baea508e68","observation_id":"a27de77b-f796-415c-8585-73b8247538b8","resolution":{"observed_at":"2026-07-02T12:06:56.162004Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2606.09669","last_updated":"2026-06-08T15:51:51Z","snapshot_observed_at":"2026-08-02T20:17:41.727375Z","submitted_at":"2026-06-08T15:51:51Z","title":"SpatialWorld: Benchmarking Interactive Spatial Reasoning of Multimodal Agents in Real-World Tasks","version":1},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-06-27T16:35:14.099586Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2606.09669"},"observation_digest":"sha256:4fec0ae1269eec67798fac20d18a01beba5c6571d4f7c030306f861159d381ce","observation_id":"75ee56c0-637b-4033-af2f-4716bda53048","resolution":{"observed_at":"2026-07-03T01:27:30.682419Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2606.17539","last_updated":"2026-06-16T05:32:39Z","snapshot_observed_at":"2026-07-06T23:53:07.149182Z","submitted_at":"2026-06-16T05:32:39Z","title":"Reinforcing Dual-Path Reasoning in Spatial Vision Language Models","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-06-27T01:42:30.005911Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2606.17539"},"observation_digest":"sha256:24d9a0654f34183cc55b50da4682df7b71bf1ba507e0931b877cf2a39733e58d","observation_id":"7eb1e127-ffc6-40f8-ab3c-a070e54e1fa4","resolution":{"observed_at":"2026-07-03T20:08:55.586365Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2606.19776","last_updated":"2026-06-18T04:24:28Z","snapshot_observed_at":"2026-08-03T12:16:21.711624Z","submitted_at":"2026-06-18T04:24:28Z","title":"Occ-VLM: Occupancy Grounded Vision Language Model for Indoor Scene Understanding","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-06-26T18:40:20.588652Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2606.19776"},"observation_digest":"sha256:377b75274bb3c379acbe1b5af5cdf029fe7c573f61f80d8d16160ac705d69298","observation_id":"2d49d67e-8080-4b9b-8b8c-83bc07d1cde2","resolution":{"observed_at":"2026-07-04T02:59:25.806993Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2606.19828","last_updated":"2026-06-18T06:12:59Z","snapshot_observed_at":"2026-08-07T17:56:22.312399Z","submitted_at":"2026-06-18T06:12:59Z","title":"3D-PLOT-LLM: Part-Level Object Tokens for 3D Large Language Models","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-06-26T18:25:55.011125Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2606.19828"},"observation_digest":"sha256:c5c8248a44ebd94ae4212d9ab99476737430082bde5bde241d747a2574353408","observation_id":"63d561c5-8a01-4c23-84f7-bc3d0e40f601","resolution":{"observed_at":"2026-07-04T03:09:29.331044Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2606.22476","last_updated":"2026-06-21T12:35:43Z","snapshot_observed_at":"2026-07-06T23:57:18.596834Z","submitted_at":"2026-06-21T12:35:43Z","title":"CVSBench: A Comprehensive Benchmark for Cross-view Spatial Reasoning and Dreaming","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-06-26T10:54:49.774490Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2606.22476"},"observation_digest":"sha256:f6d3cec1008d79a3af42ad6e664cf5e321bb3af25554038fc726d451174a1956","observation_id":"0fcf1d7a-fc6c-4154-9227-0f0858810e3d","resolution":{"observed_at":"2026-07-04T08:49:42.526702Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":"2409.18125","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-07-04T15:39:56.857028Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness","venue":null,"work_id":"25557a62-345c-4dfe-a582-4ff77ad59e6a","year":2024},"citing_paper":{"arxiv_id":"2606.24068","last_updated":"2026-06-23T02:17:08Z","snapshot_observed_at":"2026-08-02T14:48:35.079433Z","submitted_at":"2026-06-23T02:17:08Z","title":"ObsGraph: Hierarchical Observation Representation for Embodied Reasoning and Exploration","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-06-26T01:30:05.294887Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2606.24068"},"observation_digest":"sha256:2b605f0da91ce9b970b9a88514d3d424ec332f37b0451191044edbc1f0499449","observation_id":"5b1597f1-6ccf-4f79-951a-0010d5918c3a","resolution":{"observed_at":"2026-07-04T15:39:56.858556Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-02T06:33:31.366247Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.12477","last_updated":"2026-07-15T01:51:44Z","snapshot_observed_at":"2026-08-09T16:15:33.960649Z","submitted_at":"2026-07-14T08:04:31Z","title":"Self in Space: Benchmarking Self-Awareness and Spatial Cognition in UAV Embodied Intelligence","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-02T06:33:31.366247Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2607.12477"},"observation_digest":"sha256:8219d132165a3ba781317897ecd0601a36a8858c41b5c8fdfa998c2f87a4ea86","observation_id":"c41b482a-6788-499c-a549-4641c73976b2","resolution":{"observed_at":"2026-08-02T06:33:31.366247Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-02T03:31:46.403528Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.13860","last_updated":"2026-07-15T14:05:58Z","snapshot_observed_at":"2026-08-08T22:42:36.856342Z","submitted_at":"2026-07-15T14:05:58Z","title":"Towards Enhancing 3D Spatial Reasoning in Medical Multimodal Large Language Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-02T03:31:46.403528Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2607.13860"},"observation_digest":"sha256:7ff06a390f589cdd439daa1824f236514a192e66e3d98d7efaf85eb7def3c9ce","observation_id":"d4fbc204-770b-4f7d-8a80-f4f86b9f6505","resolution":{"observed_at":"2026-08-02T03:31:46.403528Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-01T13:07:03.267487Z","title":"Llava-3d: A simple yet effective pathway to empowering lmms with 3d-awareness.arXiv preprint arXiv:2409.18125, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.19228","last_updated":"2026-07-21T16:00:01Z","snapshot_observed_at":"2026-08-08T02:56:12.285976Z","submitted_at":"2026-07-21T16:00:01Z","title":"IGGT4D: Streaming 4D Instance-Grounded Geometry Transformer","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-01T13:07:03.267487Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2607.19228"},"observation_digest":"sha256:369e20cf56c50cde42872c1bbc618a4df4de17afacf3a30ac68289b6f90658e3","observation_id":"18e083b6-e336-4298-b764-a396607929ae","resolution":{"observed_at":"2026-08-01T13:07:03.267487Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18125","snapshot_observed_at":"2026-08-01T09:09:24.877243Z","title":"LLaV A-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.20868","last_updated":"2026-07-23T02:49:24Z","snapshot_observed_at":"2026-08-06T15:56:10.459729Z","submitted_at":"2026-07-23T02:49:24Z","title":"ViSTR-Bench: Can MLLMs Reason from Continuous Visual Cues in Dynamic Scenes?","version":1},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-01T09:09:24.877243Z"},"links":{"cited_paper":"/paper/2409.18125","citing_paper":"/paper/2607.20868"},"observation_digest":"sha256:0bfe9acc94f2798c9c433a449f7eda83ccf97576f7e672c99f7b24740bb491b2","observation_id":"8fd70c46-7875-4f27-8e60-647ad55a4dc6","resolution":{"observed_at":"2026-08-01T09:09:24.877243Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2409.18125/citation-record","integrity":"/paper/2409.18125/integrity","json":"/paper/2409.18125/citation-record.json","paper":"/paper/2409.18125"},"outbound":[],"paper":{"arxiv_id":"2409.18125","last_updated":"2025-04-27T06:50:23Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-07T23:35:09.059226Z","submitted_at":"2024-09-26T17:59:11Z","title":"LLaVA-3D: A Simple yet Effective Pathway to Empowering LMMs with 3D-awareness"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 55 inbound Pith citation observations for arXiv:2409.18125."}