{"as_of":"2026-08-09T17:24:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:658cb70c103b3e7909095f62531a29ad0a425f772a23a5805109b2bb76452aa9","coverage":[{"denominator":300,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T13:02:28.913510Z","state":"measured"},{"denominator":109,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":109,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":9,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":9,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T09:05:34.406350Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-01T10:15:44.967658Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"cited_work":{"arxiv_id":"2507.21045","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.21045","snapshot_observed_at":"2026-07-01T10:15:44.967658Z","title":"Reconstructing 4D spatial intelligence: A survey","venue":null,"work_id":"d1dfd87a-04ea-454f-810f-0e3484b07ed5","year":2025},"citing_paper":{"arxiv_id":"2510.17568","last_updated":"2026-08-04T12:59:26Z","snapshot_observed_at":"2026-08-07T23:11:19.285549Z","submitted_at":"2025-10-20T14:17:16Z","title":"PAGE-4D: Disentangled pose and geometry estimation for vggt-4d perception","version":6},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-18T06:19:22.275116Z"},"links":{"cited_paper":"/paper/2507.21045","citing_paper":"/paper/2510.17568"},"observation_digest":"sha256:006d609af8fa23fdd600ba8d62e1aaa4c39e8806f864a8c0b30544a6dda28c3b","observation_id":"72d89ce9-bfa9-4f90-bf78-3f2b92e094c1","resolution":{"observed_at":"2026-05-18T06:20:58.410981Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.21045","snapshot_observed_at":"2026-08-04T09:05:34.406350Z","title":"Reconstructing 4D spatial intelligence: A survey","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.17568","last_updated":"2026-08-04T12:59:26Z","snapshot_observed_at":"2026-08-07T23:11:19.285549Z","submitted_at":"2025-10-20T14:17:16Z","title":"PAGE-4D: Disentangled pose and geometry estimation for vggt-4d perception","version":7},"reference_index":2012,"source":"pdf_text","source_observed_at":"2026-08-04T09:05:34.406350Z"},"links":{"cited_paper":"/paper/2507.21045","citing_paper":"/paper/2510.17568"},"observation_digest":"sha256:93d6e3ef4cc316e414dd25c573ef15eaa2e1577774e08b530a1f74c9b1fb4af3","observation_id":"bfd21bf4-25c7-4aa3-8d80-60f0a11e5ca2","resolution":{"observed_at":"2026-08-04T09:05:34.406350Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.21045","snapshot_observed_at":"2026-08-04T08:45:14.759040Z","title":": Reconstructing 4d spatial intelligence: A survey, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.19255","last_updated":"2026-06-15T21:42:55Z","snapshot_observed_at":"2026-08-09T09:13:24.649522Z","submitted_at":"2025-10-22T05:22:20Z","title":"Advances in 4D Representation: Geometry, Motion, and Interaction","version":3},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-04T08:45:14.759040Z"},"links":{"cited_paper":"/paper/2507.21045","citing_paper":"/paper/2510.19255"},"observation_digest":"sha256:1571e57c90b035d8277969590c6178c6c46e0a12ff3beb9945dbf6c6761ee03e","observation_id":"19c1df8a-7ddf-4f7d-9f4b-fa45b3533431","resolution":{"observed_at":"2026-08-04T08:45:14.759040Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"cited_work":{"arxiv_id":"2507.21045","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.21045","snapshot_observed_at":"2026-07-01T10:15:44.967658Z","title":"Reconstructing 4D spatial intelligence: A survey","venue":null,"work_id":"d1dfd87a-04ea-454f-810f-0e3484b07ed5","year":2025},"citing_paper":{"arxiv_id":"2601.10632","last_updated":"2026-04-10T16:10:59Z","snapshot_observed_at":"2026-07-06T22:41:48.115083Z","submitted_at":"2026-01-15T17:52:29Z","title":"CoMoVi: Co-Generation of 3D Human Motions and Realistic Videos","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-16T13:43:26.460480Z"},"links":{"cited_paper":"/paper/2507.21045","citing_paper":"/paper/2601.10632"},"observation_digest":"sha256:2f2c6a060e2072ce9e8afd0a27e11ddd5481cab4d5f55d68919fd350375d7061","observation_id":"e3cc3bbe-5995-4340-bc0c-6b0b1414eb61","resolution":{"observed_at":"2026-05-16T13:47:57.520596Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"cited_work":{"arxiv_id":"2507.21045","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.21045","snapshot_observed_at":"2026-07-01T10:15:44.967658Z","title":"Reconstructing 4D spatial intelligence: A survey","venue":null,"work_id":"d1dfd87a-04ea-454f-810f-0e3484b07ed5","year":2025},"citing_paper":{"arxiv_id":"2604.07923","last_updated":"2026-07-01T09:11:07Z","snapshot_observed_at":"2026-08-02T22:35:45.631804Z","submitted_at":"2026-04-09T07:45:51Z","title":"Stitch4D: Sparse Multi-Location 4D Urban Reconstruction via Spatio-Temporal Interpolation","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T18:05:19.044269Z"},"links":{"cited_paper":"/paper/2507.21045","citing_paper":"/paper/2604.07923"},"observation_digest":"sha256:1ea06273f205c0fca2fd06e7271024cbe2c87316f7f25e0b269967cc9e619d37","observation_id":"812886aa-9395-4f1a-8867-a7bc8e618169","resolution":{"observed_at":"2026-05-11T05:31:00.499998Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.21045","snapshot_observed_at":"2026-07-13T00:10:18.016312Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.07923","last_updated":"2026-07-01T09:11:07Z","snapshot_observed_at":"2026-08-02T22:35:45.631804Z","submitted_at":"2026-04-09T07:45:51Z","title":"Stitch4D: Sparse Multi-Location 4D Urban Reconstruction via Spatio-Temporal Interpolation","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-13T00:10:18.016312Z"},"links":{"cited_paper":"/paper/2507.21045","citing_paper":"/paper/2604.07923"},"observation_digest":"sha256:09ece0393e3bb6cd32efa5155cc4ba819e161cd3bcbea33710d16cb73c60eb48","observation_id":"943342c0-c503-42a4-a790-556bb57acd82","resolution":{"observed_at":"2026-07-13T00:10:18.016312Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"cited_work":{"arxiv_id":"2507.21045","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.21045","snapshot_observed_at":"2026-07-01T10:15:44.967658Z","title":"Reconstructing 4D spatial intelligence: A survey","venue":null,"work_id":"d1dfd87a-04ea-454f-810f-0e3484b07ed5","year":2025},"citing_paper":{"arxiv_id":"2605.14462","last_updated":"2026-05-14T06:56:57Z","snapshot_observed_at":"2026-08-02T11:43:51.742393Z","submitted_at":"2026-05-14T06:56:57Z","title":"Real2Sim in HOI: Toward Physically Plausible HOI Reconstruction from Monocular Videos","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-15T02:43:51.976969Z"},"links":{"cited_paper":"/paper/2507.21045","citing_paper":"/paper/2605.14462"},"observation_digest":"sha256:c8948fe82f06a429aa09bfc0a10220fc8a416f0196c559e8a5c090e5592ffa00","observation_id":"7a00b81c-1662-4420-93a9-f7de33928076","resolution":{"observed_at":"2026-05-15T02:49:41.510037Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"cited_work":{"arxiv_id":"2507.21045","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.21045","snapshot_observed_at":"2026-07-01T10:15:44.967658Z","title":"Reconstructing 4D spatial intelligence: A survey","venue":null,"work_id":"d1dfd87a-04ea-454f-810f-0e3484b07ed5","year":2025},"citing_paper":{"arxiv_id":"2606.31388","last_updated":"2026-06-30T09:16:21Z","snapshot_observed_at":"2026-08-02T07:18:38.830324Z","submitted_at":"2026-06-30T09:16:21Z","title":"One Video, One World: Turning Monocular Video into Physical 4D Scenes","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-01T05:38:19.541391Z"},"links":{"cited_paper":"/paper/2507.21045","citing_paper":"/paper/2606.31388"},"observation_digest":"sha256:b3f7a34f02b37abe1593ec955ee247354a2afcba08a541a2bdac772326c52c92","observation_id":"8f327591-e523-40e7-9771-937a58c48eb4","resolution":{"observed_at":"2026-07-01T10:15:44.970278Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.21045","snapshot_observed_at":"2026-07-31T01:49:33.542976Z","title":"Reconstructing 4D spatial intelligence: A survey.arXiv 2507.21045, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28625","last_updated":"2026-07-30T17:59:50Z","snapshot_observed_at":"2026-08-07T19:30:56.249621Z","submitted_at":"2026-07-30T17:59:50Z","title":"ACE-Data-0: Human-Centric Ambient Capture as Embodied Data Engine","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-07-31T01:49:33.542976Z"},"links":{"cited_paper":"/paper/2507.21045","citing_paper":"/paper/2607.28625"},"observation_digest":"sha256:aa634c44a5a643dc4b4610bf340b9da042271f0a5be8de373a73defa9a82e942","observation_id":"26bc1c29-aea4-4686-8583-c5ab3534f2c9","resolution":{"observed_at":"2026-07-31T01:49:33.542976Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2507.21045/citation-record","integrity":"/paper/2507.21045/integrity","json":"/paper/2507.21045/citation-record.json","paper":"/paper/2507.21045"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.419368Z","title":"Neural point-based graphics,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.419368Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:1afc61b3a9e69fceac19146bbbc8dd3ad48ac3c1bf6cca1be45599f02f10aa18","observation_id":"c6e7ff7f-c07b-46cf-82f6-1756d3af94b7","resolution":{"observed_at":"2026-08-06T13:02:28.419368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.425203Z","title":"Deep video portraits,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.425203Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:46d01d67aaab095f75916c0ed7816cdd2d8f3732e54142baf90c2b0e3df21953","observation_id":"a1e9439e-57fc-46ea-9c81-33ad30398450","resolution":{"observed_at":"2026-08-06T13:02:28.425203Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.430574Z","title":"Fov-nerf: Foveated neural radiance fields for virtual reality,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.430574Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:bd7176d83f80ef641571c5da44af589838afc36a51ac12a107e6637f0a854a2a","observation_id":"c1a7844c-9afb-4a4f-80c8-49cc637865d0","resolution":{"observed_at":"2026-08-06T13:02:28.430574Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.435742Z","title":"Instant-3d: Instant neural radiance field training towards on-device ar/vr 3d reconstruction,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.435742Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:273ae9e6dee4dacb7d2938c308f5bed4a918fb718d91ee0d017468ed758b5e4d","observation_id":"f1308af7-0091-4275-9884-107bc55a7e00","resolution":{"observed_at":"2026-08-06T13:02:28.435742Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.06886","last_updated":"2025-08-25T02:20:09Z","snapshot_observed_at":"2026-08-06T16:04:56.349386Z","submitted_at":"2024-07-09T14:14:47Z","title":"Aligning Cyber Space with Physical World: A Comprehensive Survey on Embodied AI","version":8},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.06886","snapshot_observed_at":"2026-08-06T13:02:28.441307Z","title":"Aligning cyber space with physical world: A comprehensive survey on embodied ai,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.441307Z"},"links":{"cited_paper":"/paper/2407.06886","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:5162950e1de47749152ab61c281467779fac2ae790d76dd0a37facaac8c41287","observation_id":"d11c3e30-22fb-4d8a-bc94-76b53e6e1d85","resolution":{"observed_at":"2026-08-06T13:02:28.441307Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.06665","last_updated":"2024-04-29T23:21:33Z","snapshot_observed_at":"2026-07-06T17:28:10.554625Z","submitted_at":"2024-02-06T17:15:33Z","title":"The Essential Role of Causality in Foundation World Models for Embodied AI","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.06665","snapshot_observed_at":"2026-08-06T13:02:28.447709Z","title":"The essential role of causality in foundation world models for embodied ai,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.447709Z"},"links":{"cited_paper":"/paper/2402.06665","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:2d794edb1925639930706c940be943830e1fc7b77c367f9ec7cf443e19979007","observation_id":"10b34215-92ce-4319-a3b9-465a753d4ec2","resolution":{"observed_at":"2026-08-06T13:02:28.447709Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12871","last_updated":"2024-05-09T17:35:44Z","snapshot_observed_at":"2026-08-05T10:01:43.401012Z","submitted_at":"2023-11-18T01:21:38Z","title":"An Embodied Generalist Agent in 3D World","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.12871","snapshot_observed_at":"2026-08-06T13:02:28.454159Z","title":"An embodied generalist agent in 3d world,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.454159Z"},"links":{"cited_paper":"/paper/2311.12871","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:bc439a517cdb907ae7c08d3d3577bd2d3491a9b80e9eeb35e39cb2e215763c64","observation_id":"08f6783d-50f3-49c1-b0aa-eaffd47c4d48","resolution":{"observed_at":"2026-08-06T13:02:28.454159Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.09631","last_updated":"2024-03-14T17:58:41Z","snapshot_observed_at":"2026-08-05T20:54:58.855446Z","submitted_at":"2024-03-14T17:58:41Z","title":"3D-VLA: A 3D Vision-Language-Action Generative World Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.09631","snapshot_observed_at":"2026-08-06T13:02:28.459549Z","title":"3d-vla: A 3d vision-language-action generative world model,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.459549Z"},"links":{"cited_paper":"/paper/2403.09631","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:f3d6e7698e6eb86c4e1e9110dc547e68a1161b14418df9261bdb52f75a8720dd","observation_id":"cbe25985-f876-4a60-ab34-47c1b18f445c","resolution":{"observed_at":"2026-08-06T13:02:28.459549Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.464591Z","title":"Recent advances in 3d gaussian splatting,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.464591Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:32987c3f230dbcc8d58c256e4cd1806191f58905a34c4ba8bc182e0fccdda2f7","observation_id":"12475e8c-247c-47c1-b67b-f839ece9e26e","resolution":{"observed_at":"2026-08-06T13:02:28.464591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.469592Z","title":"3d gaussian splatting as new era: A survey,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.469592Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:b49a6c3dc2a726497d88327089de8c9804833b5390a955507ee8856e9abde6b6","observation_id":"f5aae689-edcb-470c-8ff1-fc3d8c5d79ff","resolution":{"observed_at":"2026-08-06T13:02:28.469592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.474115Z","title":"Stereo matching algorithm based on deep learning: A survey,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.474115Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:95466ec3ed9c03b5bf5a813e4c7a742ac359d163b09c4f9dc709700097ee2c43","observation_id":"8f860274-898a-4c50-92db-2087b2bdd7e2","resolution":{"observed_at":"2026-08-06T13:02:28.474115Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.479707Z","title":"Review of stereo matching algorithms based on deep learning,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.479707Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:25e3137f8eefcdb6ba0b6a187738e06306e8bfdfd33c253eacfda19ecf899724","observation_id":"8a6f37d9-86e6-4149-ad3e-b98bd52ccb78","resolution":{"observed_at":"2026-08-06T13:02:28.479707Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.484726Z","title":"A sur- vey on deep learning techniques for stereo-based depth estima- tion,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.484726Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:a43bc6b80226fe15fecc4b8d0badaa7721a7ad78ed7065e27fadbd01eb44fbc5","observation_id":"e335c623-4028-4dd5-9e27-2a7a37828bf5","resolution":{"observed_at":"2026-08-06T13:02:28.484726Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.00379","last_updated":"2026-02-12T16:09:50Z","snapshot_observed_at":"2026-07-06T13:58:37.592838Z","submitted_at":"2022-10-01T21:35:11Z","title":"NeRF: Neural Radiance Field in 3D Vision: A Comprehensive Review (Updated Post-Gaussian Splatting)","version":8},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.00379","snapshot_observed_at":"2026-08-06T13:02:28.489403Z","title":"Nerf: Neural radiance field in 3d vision, a comprehensive review,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.489403Z"},"links":{"cited_paper":"/paper/2210.00379","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:8e3b9e7d6c54026b783719e82d5660cdfd899aa1d0f954ea883a05f4e1009919","observation_id":"8199cb7a-909d-4cf8-93ed-c88a93082c50","resolution":{"observed_at":"2026-08-06T13:02:28.489403Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.494181Z","title":"3d gaussian splatting: Survey, technologies, challenges, and opportunities,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.494181Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:d004247bec7cfa89caa57acf5325ed663619a06ed289b1ef4858c02ffea05e6a","observation_id":"b8434432-44e7-42c3-8fd5-39b8219a516a","resolution":{"observed_at":"2026-08-06T13:02:28.494181Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.01333","last_updated":"2025-07-02T01:46:31Z","snapshot_observed_at":"2026-08-07T03:19:09.642620Z","submitted_at":"2024-05-02T14:38:18Z","title":"NeRFs in Robotics: A Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.01333","snapshot_observed_at":"2026-08-06T13:02:28.498546Z","title":"Nerf in robotics: A survey,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.498546Z"},"links":{"cited_paper":"/paper/2405.01333","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6f5bea23071d3c268b2428f2485268913e7318e3376831ff58c1969a17dab990","observation_id":"9fd19c95-e2bf-4a49-a4dd-3d524cef8f84","resolution":{"observed_at":"2026-08-06T13:02:28.498546Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.503607Z","title":"Nerf: Representing scenes as neural radiance fields for view synthesis,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.503607Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:c58ae663bffe627eb2dbd61a76ef1d6562f53a9e8712e095b5c3df51e50f70fe","observation_id":"1a0fa5f9-580b-44b7-91e1-6dbd12b14ae0","resolution":{"observed_at":"2026-08-06T13:02:28.503607Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.507780Z","title":"Deep marching tetrahedra: a hybrid representation for high-resolution 3d shape synthesis,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.507780Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:35a38300c44ab0097cff5471270f169622cd0fcf45c979889aa8b7806b8c4e67","observation_id":"f3979560-811f-43b3-878a-647906ce60ce","resolution":{"observed_at":"2026-08-06T13:02:28.507780Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.512197Z","title":"3d gaussian splatting for real-time radiance field rendering,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.512197Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6eac0da2ff9b2027f9b81974bb0b469a49561461cd0df38f42c3099957223dd1","observation_id":"511d9e9f-576b-4770-82c7-cfa382134a63","resolution":{"observed_at":"2026-08-06T13:02:28.512197Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.517181Z","title":"Video diffusion models,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.517181Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:fba54b81fe47980551b74d59ba6a96b2efcb092ba0bda09ef0b0a4c59479ba2d","observation_id":"128e6d1b-aa1c-4e0b-908a-a210d71abd99","resolution":{"observed_at":"2026-08-06T13:02:28.517181Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.02303","last_updated":"2022-10-05T14:41:38Z","snapshot_observed_at":"2026-07-06T13:59:57.800591Z","submitted_at":"2022-10-05T14:41:38Z","title":"Imagen Video: High Definition Video Generation with Diffusion Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.02303","snapshot_observed_at":"2026-08-06T13:02:28.524542Z","title":"Imagen video: High definition video generation with diffusion models,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.524542Z"},"links":{"cited_paper":"/paper/2210.02303","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:18d6b1251e00e9b9683a547c37d484da4ae73d23fa9e064cfbf9104dfe04cdac","observation_id":"666c0483-83c7-428b-95bc-bb04896ddbff","resolution":{"observed_at":"2026-08-06T13:02:28.524542Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.15127","last_updated":"2023-11-25T22:28:38Z","snapshot_observed_at":"2026-08-07T21:47:08.589400Z","submitted_at":"2023-11-25T22:28:38Z","title":"Stable Video Diffusion: Scaling Latent Video Diffusion Models to Large Datasets","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.15127","snapshot_observed_at":"2026-08-06T13:02:28.529792Z","title":"Stable video diffusion: Scaling latent video diffusion models to large datasets,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.529792Z"},"links":{"cited_paper":"/paper/2311.15127","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6eece53f5f97a3e77af089f3c95d14307b9a6f3215e2c95b7b0a94a456070c6e","observation_id":"257614ee-a28b-41f6-999c-cb64003f62ae","resolution":{"observed_at":"2026-08-06T13:02:28.529792Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.537987Z","title":"Sift: Predicting amino acid changes that affect protein function,","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.537987Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:75415239254e8ad0096ade0644549d86216a23d2ca68b71e717692fc5d4296bf","observation_id":"0026f052-50fe-44d6-915c-ec1f4a5599d9","resolution":{"observed_at":"2026-08-06T13:02:28.537987Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.543574Z","title":"R2d2: Reliable and repeatable detector and descriptor,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.543574Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:96bc26ef277737d1aa4aab71fab11f7bd326f68514f05e85b8982d7f708119e3","observation_id":"2a0b0285-bac7-4c27-a7d8-22b012cef2a3","resolution":{"observed_at":"2026-08-06T13:02:28.543574Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.548175Z","title":"Superpoint: Self- supervised interest point detection and description,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.548175Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:0f27c4fb5d7cdb392f2a0c64257b2c1da474183095765bd58c96a93697f9d025","observation_id":"366d6eb0-b2f1-4fb9-a9b7-fbec2de8975b","resolution":{"observed_at":"2026-08-06T13:02:28.548175Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.553255Z","title":"Superglue: Learning feature matching with graph neural net- works,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.553255Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:4342376becb76851c44a19489a243a4088ea3ddca2f9a14869a54e2871ae9f63","observation_id":"88fd6f4b-98b3-4e1d-a9e7-0c90f73a8496","resolution":{"observed_at":"2026-08-06T13:02:28.553255Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.557555Z","title":"Loftr: Detector- free local feature matching with transformers,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.557555Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:31e582b8f93547e3f8985bd69ec5d11646808f9f4a288c33c46ff142c18fe6bc","observation_id":"5e6d1284-c534-45e4-bd56-f4cc777e0d14","resolution":{"observed_at":"2026-08-06T13:02:28.557555Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.562152Z","title":"Neural-guided ransac: Learn- ing where to sample model hypotheses,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.562152Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:a391e685ed32dccc697f049aec5a7ad4ec9363b676ec1b478c402f1f26b7c269","observation_id":"34d7f5a8-f705-49c1-961c-b4919e2e17ac","resolution":{"observed_at":"2026-08-06T13:02:28.562152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.566848Z","title":"Lightglue: Local feature matching at light speed,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.566848Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:e1c803ab329edb07a2b62611a29bc46af15cb56e432c05efd05e17d21e387427","observation_id":"ec9b5140-aafd-484a-9d94-94f5280c877b","resolution":{"observed_at":"2026-08-06T13:02:28.566848Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15381","last_updated":"2023-07-28T08:05:36Z","snapshot_observed_at":"2026-07-06T15:59:38.850337Z","submitted_at":"2023-07-28T08:05:36Z","title":"AffineGlue: Joint Matching and Robust Estimation","version":1},"cited_work":{"arxiv_id":"2307.15381","doi":null,"metadata_source":"pith","pith_arxiv_id":"2307.15381","snapshot_observed_at":"2026-08-06T13:02:32.717939Z","title":"AffineGlue: Joint Matching and Robust Estimation","venue":"cs.CV","work_id":"cb7170ff-993b-4fb8-88cb-e612f7610c01","year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.571255Z"},"links":{"cited_paper":"/paper/2307.15381","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:1b8da67fb5b92fdfe58ffd07bd9a72bc2aebd95d29926734e28123b4655ebc93","observation_id":"975cc149-dd51-4be9-9ffc-494037e71293","resolution":{"observed_at":"2026-08-06T13:02:32.777980Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.576091Z","title":"Structure-from-motion revis- ited,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.576091Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:0766d96bcfb715a669d14372021226954361717dee71d22374d22a1a09231724","observation_id":"41971552-fc4a-4857-96ce-5a32801b96a5","resolution":{"observed_at":"2026-08-06T13:02:28.576091Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.580811Z","title":"A survey of structure from motion*","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.580811Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:e093ace4e51fdf0e565138c11da2d73ea96ed38338f2362b472cd13d822893a5","observation_id":"4f06510c-a1b3-4d1c-ab3b-5167dfb6e869","resolution":{"observed_at":"2026-08-06T13:02:28.580811Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.585947Z","title":"Structure from motion photogrammetry in forestry: A review,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.585947Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:1bf688118c4e0b772c2fc5363b54d66d8a3611dfb8eb5e15ab47cc026ba1b113","observation_id":"40f9676d-8e61-4347-a1a2-d8b1ec8294ca","resolution":{"observed_at":"2026-08-06T13:02:28.585947Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.591059Z","title":"Pixel- perfect structure-from-motion with featuremetric refinement,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.591059Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:42769ef07f16700d0f5cbb58424dda3ec1d8c832b777c1fe46635707a450eadf","observation_id":"6dcf4a3e-a2fc-4040-bcf7-c900a80f725b","resolution":{"observed_at":"2026-08-06T13:02:28.591059Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.597662Z","title":"Bundle adjustment in the large,","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.597662Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:c5856df54d504e7edc006ec17dac7bff3a6b0b2c799173ea58c842975dc76690","observation_id":"38862f6b-859d-407d-803c-782905b38a5e","resolution":{"observed_at":"2026-08-06T13:02:28.597662Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.602366Z","title":"Bundle adjustment rules,","venue":null,"work_id":null,"year":2006},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.602366Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:c97fc87107ab4797e58b4448441b6fdd06b2e2e09cf9875927a2d9e5adab6211","observation_id":"8746f4f0-2a83-4795-ab42-403380c68b94","resolution":{"observed_at":"2026-08-06T13:02:28.602366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.608307Z","title":"Robust bundle adjustment revisited,","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.608307Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:b37308e5493d7b4d29545813b48dc2114104fbd8979b05d4de760a4f5a705fb1","observation_id":"b5996345-34e6-4457-8242-e100488ec87f","resolution":{"observed_at":"2026-08-06T13:02:28.608307Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.612979Z","title":"Bundle adjustment—a modern synthesis,","venue":null,"work_id":null,"year":1999},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.612979Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:c52c9194bb3d2bba6d2c8e2f0adc0f048e978861be14a9c90cf8492c9c1b080d","observation_id":"33e0e349-8f07-450c-9408-e2e420ca138f","resolution":{"observed_at":"2026-08-06T13:02:28.612979Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.617699Z","title":"Pixelwise view selection for unstructured multi-view stereo,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.617699Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:ed29aaf4d5f1d49d9ea98e05f09f165bb76247bc69286225344fb4ed4a76a412","observation_id":"e92b7b1a-159a-4835-8bbf-ec6e49f89f62","resolution":{"observed_at":"2026-08-06T13:02:28.617699Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.622713Z","title":"Visibility-aware multi-view stereo network,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.622713Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:752eacfd6858f1579b6b8c4556fde946cad1b272f26a795074e1db57ad513012","observation_id":"7f42f9e8-75a4-4f72-9793-c0f850207b64","resolution":{"observed_at":"2026-08-06T13:02:28.622713Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.628213Z","title":"Cost volume pyramid based depth inference for multi-view stereo,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.628213Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:d049a57259252ca0b670b32f450e1e78af74a0184018e0a049960a7f055b2c0f","observation_id":"10f2e351-bd53-4dfb-ad38-1942653ea95a","resolution":{"observed_at":"2026-08-06T13:02:28.628213Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.633394Z","title":"Patchmatchnet: Learned multi-view patchmatch stereo,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.633394Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:fa091072aa804beeed86a5025bdf82b952dd0b36e59cede664efb39771ad2228","observation_id":"aa00828e-46ef-4130-a524-759172ae9bd9","resolution":{"observed_at":"2026-08-06T13:02:28.633394Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.637870Z","title":"Cascade cost volume for high-resolution multi-view stereo and stereo matching,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.637870Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:485ef87b7df2f62adb9c68d93cdaafee9ab18e32498f1279b7295b7244a3a298","observation_id":"9f86db23-6669-415c-8b5c-c8aec544bfc9","resolution":{"observed_at":"2026-08-06T13:02:28.637870Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.642579Z","title":"Dust3r: Geometric 3d vision made easy,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.642579Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:2e6cd9457a7f54cec44daa5feb2fffc1600e913b098213622ca9d6bd10ade4f7","observation_id":"67445b25-e7cd-4187-b631-553b8cc354dc","resolution":{"observed_at":"2026-08-06T13:02:28.642579Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.12392","last_updated":"2025-06-02T13:44:10Z","snapshot_observed_at":"2026-08-09T04:40:44.217468Z","submitted_at":"2024-12-16T23:00:05Z","title":"MASt3R-SLAM: Real-Time Dense SLAM with 3D Reconstruction Priors","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.12392","snapshot_observed_at":"2026-08-06T13:02:28.647067Z","title":"Mast3r-slam: Real- time dense slam with 3d reconstruction priors,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.647067Z"},"links":{"cited_paper":"/paper/2412.12392","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:4dd0fff46b43b304ede502436d40c921a86d3f29d44e13af0c31905ad354ef28","observation_id":"40dc9659-ae0c-48ff-a452-0fe0c59f6063","resolution":{"observed_at":"2026-08-06T13:02:28.647067Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.03825","last_updated":"2025-05-08T08:32:16Z","snapshot_observed_at":"2026-07-06T19:28:08.374354Z","submitted_at":"2024-10-04T18:00:07Z","title":"MonST3R: A Simple Approach for Estimating Geometry in the Presence of Motion","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.03825","snapshot_observed_at":"2026-08-06T13:02:28.651730Z","title":"Monst3r: A simple approach for estimating geometry in the presence of motion,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.651730Z"},"links":{"cited_paper":"/paper/2410.03825","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:c6f47ce323ef29b9675fb7a41a7b6ed2ece8599c3f1dd73ad0ae3832dd713491","observation_id":"7cb70d8f-b439-43c4-b707-d76e823d6213","resolution":{"observed_at":"2026-08-06T13:02:28.651730Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.03079","last_updated":"2024-12-05T14:16:07Z","snapshot_observed_at":"2026-07-06T20:01:23.373458Z","submitted_at":"2024-12-04T07:09:59Z","title":"Align3R: Aligned Monocular Depth Estimation for Dynamic Videos","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.03079","snapshot_observed_at":"2026-08-06T13:02:28.656186Z","title":"Align3r: Aligned monocular depth estimation for dynamic videos,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.656186Z"},"links":{"cited_paper":"/paper/2412.03079","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:b468c4f0f408d9194a4759f7cdb04b3a5d765354274ae3f1fde742a920cd0e5f","observation_id":"16f5040d-184b-4b9c-a63c-8153ed52ecbb","resolution":{"observed_at":"2026-08-06T13:02:28.656186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13928","last_updated":"2025-03-19T19:35:52Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:55Z","title":"Fast3R: Towards 3D Reconstruction of 1000+ Images in One Forward Pass","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13928","snapshot_observed_at":"2026-08-06T13:02:28.661267Z","title":"Fast3r: Towards 3d recon- struction of 1000+ images in one forward pass,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.661267Z"},"links":{"cited_paper":"/paper/2501.13928","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:eb28292371d4fb5addc098dd69b937d309f2c90e5f1e281f51592a12ce36785a","observation_id":"a2e8551b-ae97-4353-9b5b-f15144821955","resolution":{"observed_at":"2026-08-06T13:02:28.661267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.666024Z","title":"Transformer in transformer,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.666024Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:b209618172e6a302cae889db3482f651aee94f95566c229511b10180579682ec","observation_id":"252328cb-b926-4f86-bd4c-91d26f82b493","resolution":{"observed_at":"2026-08-06T13:02:28.666024Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.670423Z","title":"A survey on vision transformer,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.670423Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:bc3c88802bd354af6e173a54eefd897b7bd90107a754121f27be8caa9623d2e2","observation_id":"92a1cd8b-87db-47c2-b35f-dfbd5da7ecae","resolution":{"observed_at":"2026-08-06T13:02:28.670423Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2001.04451","last_updated":"2020-02-18T16:01:18Z","snapshot_observed_at":"2026-07-06T08:50:12.690900Z","submitted_at":"2020-01-13T18:38:28Z","title":"Reformer: The Efficient Transformer","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2001.04451","snapshot_observed_at":"2026-08-06T13:02:28.674898Z","title":"Reformer: The efficient transformer,","venue":null,"work_id":null,"year":2001},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.674898Z"},"links":{"cited_paper":"/paper/2001.04451","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:e670e88f4c880bb7913f18bc30cedf662f9790af5f656e86e848e5ac7a9965c2","observation_id":"0acb0fbc-e236-4bd8-bd30-88373d6024f3","resolution":{"observed_at":"2026-08-06T13:02:28.674898Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.680102Z","title":"Point trans- former,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.680102Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:2c5858c1fe685a9129651a3259f12f98ffc2609348dbc186ebc55b20dd1a4d17","observation_id":"d1abb665-928f-4dcc-b32e-895830f01ed2","resolution":{"observed_at":"2026-08-06T13:02:28.680102Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.684907Z","title":"Image transformer,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.684907Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:4b1a8668f27035e243b111794948001e27913b27f5b56e1d6ccc60926db5722a","observation_id":"41334d01-311f-40ab-98a5-0cdb94ed2091","resolution":{"observed_at":"2026-08-06T13:02:28.684907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.689601Z","title":"Vggt: Visual geometry grounded transformer,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.689601Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:a5fded37d8edb2e6b80f0cb8a782ac938ad71b5e9da2dd939b86a684a5176ce4","observation_id":"4ba37fd8-008c-4704-ab3c-9c67cdb302ec","resolution":{"observed_at":"2026-08-06T13:02:28.689601Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.694509Z","title":"Nerf: Representing scenes as neural radiance fields for view synthesis,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.694509Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:be4cc9f7be8936447a1cfd8cf8bbb3d4a6a677417f721b299bf217d6b6e6100e","observation_id":"7bf2e189-bced-459c-8c02-910aecd87fd4","resolution":{"observed_at":"2026-08-06T13:02:28.694509Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.699278Z","title":"3d gaus- sian splatting for real-time radiance field rendering,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.699278Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6de7632c1a91801df42b47abd127ffdc4e85d0254f89449d1a2a13cf18351c09","observation_id":"75832935-0107-4abd-9507-d9d6ce408870","resolution":{"observed_at":"2026-08-06T13:02:28.699278Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.703915Z","title":"Flexible isosurface extraction for gradient-based mesh optimization,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.703915Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:d5d5ff1814e00a47d2171ac527b7fdd95e161854df5fc756862e43df2d06e2d7","observation_id":"f3ed81cb-827d-4f4d-8a80-e9afd7176447","resolution":{"observed_at":"2026-08-06T13:02:28.703915Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.708634Z","title":"Nerfies: Deformable neural radi- ance fields,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.708634Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:abb776084ca1cc60b1eaebcb08068640762ea95b0e5f16a8062e83f525a92f3f","observation_id":"58544010-9747-49eb-a24f-7aa9b4cc55ed","resolution":{"observed_at":"2026-08-06T13:02:28.708634Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.713660Z","title":"Non-rigid neural radiance fields: Reconstruction and novel view synthesis of a dynamic scene from monocular video,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.713660Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6f4e28a3dc66122e8fe07b4a3ac92a033e24b1e20cf6b432ac52bb607855385f","observation_id":"3ec67928-0e2a-4aa6-a42e-001e9f2c5e9e","resolution":{"observed_at":"2026-08-06T13:02:28.713660Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.13228","last_updated":"2021-09-10T06:34:46Z","snapshot_observed_at":"2026-08-09T08:50:52.811232Z","submitted_at":"2021-06-24T17:59:03Z","title":"HyperNeRF: A Higher-Dimensional Representation for Topologically Varying Neural Radiance Fields","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.13228","snapshot_observed_at":"2026-08-06T13:02:28.718303Z","title":"Hypernerf: A higher-dimensional representation for topologically varying neu- ral radiance fields,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.718303Z"},"links":{"cited_paper":"/paper/2106.13228","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:51c9dedd209a3c0c6407387840786c60755a88a79664e6c19a2eb9c6dfec6647","observation_id":"25647800-e458-40c8-83d9-bad5b9cb2faa","resolution":{"observed_at":"2026-08-06T13:02:28.718303Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.723527Z","title":"Dylin: Making light field networks dynamic,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.723527Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:234ea3e82e19ae148a7f0114b9f1e8a851d6b6cd5e88452051434b80ad129e55","observation_id":"6fab469c-8dc3-46c0-b145-0966e13ca41e","resolution":{"observed_at":"2026-08-06T13:02:28.723527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.17249","last_updated":"2025-03-25T16:10:41Z","snapshot_observed_at":"2026-07-06T19:38:00.663759Z","submitted_at":"2024-10-22T17:59:56Z","title":"SpectroMotion: Dynamic 3D Reconstruction of Specular Scenes","version":3},"cited_work":{"arxiv_id":"2410.17249","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.17249","snapshot_observed_at":"2026-08-06T13:02:32.385813Z","title":"SpectroMotion: Dynamic 3D Reconstruction of Specular Scenes","venue":"cs.CV","work_id":"9e109631-5766-43cd-81c7-90dcadf08155","year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.728580Z"},"links":{"cited_paper":"/paper/2410.17249","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:343eca7cc19c63fd9aabeb5f79aab58bdeaba61f393c1428b5d31ea679ca0d76","observation_id":"3e8d9ef5-3ef1-494f-91be-707486e89ee3","resolution":{"observed_at":"2026-08-06T13:02:32.458725Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.733564Z","title":"Neural scene flow fields for space-time view synthesis of dynamic scenes,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.733564Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:58174e187dd5be0456a14758a1909c19bf90bf1dc3ab5956b837068f8002bba7","observation_id":"526f613c-91df-41a4-8f49-b0487d07a375","resolution":{"observed_at":"2026-08-06T13:02:28.733564Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.738151Z","title":"Dynamic view synthesis from dynamic monocular video,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.738151Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:a15c761832513ff02fcf08d5d7b8a6c98ed3557d62fbeb69f0097316b9ca71c2","observation_id":"a71e1f89-d020-4bcf-b018-61b28b62c8ba","resolution":{"observed_at":"2026-08-06T13:02:28.738151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.742732Z","title":"Decoupling dynamic monocular videos for dynamic view synthesis,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.742732Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:25312ab882363e0473c4837eaf25b8cbc4b9529714e93c66d7c7da9341d451be","observation_id":"a371938b-1eff-4c0e-8579-f9c0477f8d8e","resolution":{"observed_at":"2026-08-06T13:02:28.742732Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2105.05994","last_updated":"2021-05-12T22:38:30Z","snapshot_observed_at":"2026-07-06T11:08:59.470370Z","submitted_at":"2021-05-12T22:38:30Z","title":"Neural Trajectory Fields for Dynamic Novel View Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2105.05994","snapshot_observed_at":"2026-08-06T13:02:28.747523Z","title":"Neural trajec- tory fields for dynamic novel view synthesis,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.747523Z"},"links":{"cited_paper":"/paper/2105.05994","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:42277cc1e22fbf648d7bca85718b369ec1d7717104fa01ec6858167a9c5b2041","observation_id":"d96fb4e2-fee4-4077-9bc0-3bed4844e37e","resolution":{"observed_at":"2026-08-06T13:02:28.747523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.752556Z","title":"Neural scene chronology,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.752556Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:2cae4e7351486e84434f5727cc564cd7fdc7d5f25856241b73d96453e00f1a86","observation_id":"7ff990da-bae3-42d4-9fc9-72024058345b","resolution":{"observed_at":"2026-08-06T13:02:28.752556Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.757472Z","title":"Darenerf: Direction-aware representation for dynamic scenes,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.757472Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6966376b4a101bdf53ebc9fb863cc73e0e6f9db196f6a5c92564222510782997","observation_id":"3f50f4f7-ed4f-49e8-910e-d7795452af69","resolution":{"observed_at":"2026-08-06T13:02:28.757472Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.02077","last_updated":"2023-11-03T17:59:55Z","snapshot_observed_at":"2026-08-08T00:21:57.482694Z","submitted_at":"2023-11-03T17:59:55Z","title":"EmerNeRF: Emergent Spatial-Temporal Scene Decomposition via Self-Supervision","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.02077","snapshot_observed_at":"2026-08-06T13:02:28.762758Z","title":"Emernerf: Emer- gent spatial-temporal scene decomposition via self-supervision,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.762758Z"},"links":{"cited_paper":"/paper/2311.02077","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:9ed6a199d0463d5c01b9e4089d8e393e4278f719664d92e9ab46cea0db178d83","observation_id":"a685bbb5-ed60-4461-86a5-cb65df421f0e","resolution":{"observed_at":"2026-08-06T13:02:28.762758Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.767767Z","title":"Gravity-aware monocular 3d human-object reconstruction,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.767767Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:ac866a22b743cbca8ea3b9106895188b2085289791d63c81175e2297ed098659","observation_id":"860b554d-8800-41a5-bb7d-103f99e7ed60","resolution":{"observed_at":"2026-08-06T13:02:28.767767Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2108.08420","last_updated":"2021-08-19T00:49:01Z","snapshot_observed_at":"2026-08-07T21:09:54.846761Z","submitted_at":"2021-08-19T00:49:01Z","title":"D3D-HOI: Dynamic 3D Human-Object Interactions from Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2108.08420","snapshot_observed_at":"2026-08-06T13:02:28.772694Z","title":"D3d-hoi: Dynamic 3d human-object interactions from videos,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.772694Z"},"links":{"cited_paper":"/paper/2108.08420","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:dbe7a44b9b8146371d9969e7039ca19c9bd64411f28fb4cfb2bb1f353c7e49aa","observation_id":"e28055c2-e59e-4a3c-9788-d9d469d14fab","resolution":{"observed_at":"2026-08-06T13:02:28.772694Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.777960Z","title":"Behave: Dataset and method for tracking human object interactions,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.777960Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:97faeecb75414272e39a4f6e25f86bcddd1a98f0c9a192f905f1bcd154485ef8","observation_id":"6cd655e2-5df6-46d1-a02b-f6585917cccd","resolution":{"observed_at":"2026-08-06T13:02:28.777960Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.782864Z","title":"Intercap: Joint markerless 3d tracking of humans and objects in interaction,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.782864Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:70bdd9c5a01baa005df6b41c1bf2dbf1b53e94f41b69d18202176d3917f3cee4","observation_id":"eab26ea9-9390-45fd-9aa1-59a539bea58f","resolution":{"observed_at":"2026-08-06T13:02:28.782864Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.787541Z","title":"Full-body articulated human-object inter- action,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.787541Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6dab1f1975d529cc9ea791696261081bd3c4c04c1d7fee31ffddb7b90581aa65","observation_id":"b1f9718c-3c71-424d-9eb2-10345d3c733b","resolution":{"observed_at":"2026-08-06T13:02:28.787541Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.20545","last_updated":"2024-07-30T04:57:21Z","snapshot_observed_at":"2026-08-07T11:19:02.323666Z","submitted_at":"2024-07-30T04:57:21Z","title":"StackFLOW: Monocular Human-Object Reconstruction by Stacked Normalizing Flow with Offset","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.20545","snapshot_observed_at":"2026-08-06T13:02:28.792493Z","title":"Stackflow: Monocular human-object reconstruction by stacked normalizing flow with offset,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.792493Z"},"links":{"cited_paper":"/paper/2407.20545","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:5cadb0fd689b690e368601cf69e552356f3d20791ee241948f4533ab5d801ff7","observation_id":"fefc755b-2754-417e-9266-ac0ce537708c","resolution":{"observed_at":"2026-08-06T13:02:28.792493Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.797346Z","title":"Monocular human-object recon- struction in the wild,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.797346Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:68dbd434fbeaee349f952093f6bb668a95d4be5bdac00d5465fbf492d882340e","observation_id":"9fcbd8cd-8965-44e7-94c7-6ddf8bd7ba76","resolution":{"observed_at":"2026-08-06T13:02:28.797346Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.801891Z","title":"I’m hoi: Inertia-aware monocular capture of 3d human- object interactions,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.801891Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:3c6b4f471ebdb6a3665f3333938cd82ca96418c516fec4ae189df1beed53f9e8","observation_id":"0adb414a-c4d3-4366-9166-fa9a10709aff","resolution":{"observed_at":"2026-08-06T13:02:28.801891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17470","last_updated":"2025-02-27T21:52:39Z","snapshot_observed_at":"2026-08-04T05:24:01.192706Z","submitted_at":"2024-07-24T17:59:43Z","title":"SV4D: Dynamic 3D Content Generation with Multi-Frame and Multi-View Consistency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17470","snapshot_observed_at":"2026-08-06T13:02:28.806630Z","title":"Sv4d: Dynamic 3d content generation with multi-frame and multi-view consistency,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.806630Z"},"links":{"cited_paper":"/paper/2407.17470","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:eb7a6e6b194b33c7dd1f66d5f34ec87e5765190d824c77bb28388f99a8abc2af","observation_id":"98d946b5-81dd-48e5-bc77-2497abca4b06","resolution":{"observed_at":"2026-08-06T13:02:28.806630Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.13953","last_updated":"2024-08-25T22:26:46Z","snapshot_observed_at":"2026-08-06T11:13:19.673489Z","submitted_at":"2024-08-25T22:26:46Z","title":"InterTrack: Tracking Human Object Interaction without Object Templates","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.13953","snapshot_observed_at":"2026-08-06T13:02:28.813129Z","title":"Intertrack: Tracking hu- man object interaction without object templates,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.813129Z"},"links":{"cited_paper":"/paper/2408.13953","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:5a4dab0be587317020574a9c843b43fd4a458931bbdf9b5daa848935d3728b5b","observation_id":"434af2c3-3f40-490a-ab10-e58855fe3860","resolution":{"observed_at":"2026-08-06T13:02:28.813129Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.07164","last_updated":"2024-10-09T17:58:56Z","snapshot_observed_at":"2026-08-08T08:59:31.271794Z","submitted_at":"2024-10-09T17:58:56Z","title":"AvatarGO: Zero-shot 4D Human-Object Interaction Generation and Animation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.07164","snapshot_observed_at":"2026-08-06T13:02:28.818237Z","title":"Avatargo: Zero-shot 4d human-object interaction generation and anima- tion,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.818237Z"},"links":{"cited_paper":"/paper/2410.07164","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:441bd502040f5cb375602f069dacb56a1365af5d662576448162378ac507eb08","observation_id":"9e8406f3-fcbb-4376-a3c5-7477622babb3","resolution":{"observed_at":"2026-08-06T13:02:28.818237Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.824114Z","title":"The one where they reconstructed 3d humans and environments in tv shows,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.824114Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6791a8c4467d40640ce2ef64fe001c5b3951b527576fc75c0c433d512218ee01","observation_id":"e1cb13cc-12ff-413c-843e-b5944adee604","resolution":{"observed_at":"2026-08-06T13:02:28.824114Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.02158","last_updated":"2026-06-29T19:38:44Z","snapshot_observed_at":"2026-08-06T14:43:41.891230Z","submitted_at":"2025-01-04T01:53:51Z","title":"Joint Optimization for 4D Human-Scene Reconstruction in the Wild","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.02158","snapshot_observed_at":"2026-08-06T13:02:28.829575Z","title":"Joint optimization for 4d human-scene reconstruction in the wild,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.829575Z"},"links":{"cited_paper":"/paper/2501.02158","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:50581a546f6dadfcb637ec8175cf1be8ec66a8d435bc4e2db40c433f978289ed","observation_id":"ce736998-eb89-4015-b243-c4707f6a395a","resolution":{"observed_at":"2026-08-06T13:02:28.829575Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.835330Z","title":"Odhsr: Online dense 3d reconstruction of humans and scenes from monocular videos,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.835330Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:ffb19498703a8fbba835b293e398639d19ee56d8fd3547811049543af09ea6a2","observation_id":"28acbc36-f43f-4a04-be28-79d615eec000","resolution":{"observed_at":"2026-08-06T13:02:28.835330Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.839788Z","title":"Hosnerf: Dynamic human-object-scene neu- ral radiance fields from a single video,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.839788Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:cc6774e9d0b0c7a999e1c10b7f79475eae43f555c5ad1d319669f8f8e351480b","observation_id":"09ada350-2779-4d27-b659-8bf87af3c972","resolution":{"observed_at":"2026-08-06T13:02:28.839788Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.844890Z","title":"Neuman: Neural human radiance field from a single video,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.844890Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:0475433e148e4499eb21d9f141ab2c0ba4a3f3f66eeb8f83243292f3016afb8e","observation_id":"b14afb07-2f2f-4915-9cb0-20741adf8ec6","resolution":{"observed_at":"2026-08-06T13:02:28.844890Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.04393","last_updated":"2023-12-07T16:06:31Z","snapshot_observed_at":"2026-08-06T12:06:41.391790Z","submitted_at":"2023-12-07T16:06:31Z","title":"PhysHOI: Physics-Based Imitation of Dynamic Human-Object Interaction","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.04393","snapshot_observed_at":"2026-08-06T13:02:28.849195Z","title":"Physhoi: Physics-based imitation of dynamic human-object in- teraction,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.849195Z"},"links":{"cited_paper":"/paper/2312.04393","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:0c77ffbee0cad6155581e61fb74fff7cfbd57c616578926a370127f9d81ddc8f","observation_id":"a2830648-99d7-4592-b864-3c7ed6ba11b7","resolution":{"observed_at":"2026-08-06T13:02:28.849195Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.853667Z","title":"Perpetual humanoid control for real-time simulated avatars,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.853667Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:d2e312e46fa97b0d54b2902a4ef3927eae8e5f49ecd29e7ae98fb03135c461fb","observation_id":"8876312f-7d6b-4f7e-af69-40e384290c6c","resolution":{"observed_at":"2026-08-06T13:02:28.853667Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.857916Z","title":"Universal humanoid motion representations for physics-based control,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.857916Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:46542712f23d2d14b8c5ab16620bcadc5805025bfa52f90c48094651918003fb","observation_id":"906b0fe4-9eb3-4c03-a83d-83dc6009d54c","resolution":{"observed_at":"2026-08-06T13:02:28.857916Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2108.10470","last_updated":"2021-08-25T23:42:59Z","snapshot_observed_at":"2026-07-06T11:40:56.544714Z","submitted_at":"2021-08-24T01:38:11Z","title":"Isaac Gym: High Performance GPU-Based Physics Simulation For Robot Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2108.10470","snapshot_observed_at":"2026-08-06T13:02:28.862444Z","title":"Isaac gym: High performance gpu-based physics simulation for robot learning,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.862444Z"},"links":{"cited_paper":"/paper/2108.10470","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:ade487da41b20268eaf16a2e3adb646a2465f3c2aca22446726a92e10b5a017e","observation_id":"35a47705-4269-437a-9452-97c9b5f3155a","resolution":{"observed_at":"2026-08-06T13:02:28.862444Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.866857Z","title":"Reinforcement learning,","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.866857Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:3af05334f8cff86ad1a1d584082c2ead6fa05b47e4732ee0765e78de75b73fda","observation_id":"841601de-7b5d-4e99-983b-d3a06ef59ab3","resolution":{"observed_at":"2026-08-06T13:02:28.866857Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.871324Z","title":"Reinforcement learning: A survey,","venue":null,"work_id":null,"year":1996},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.871324Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6d79cb204459288a04dfe938235ce2013f368aaf5b86ea63bd998a1556a7db14","observation_id":"aa8a111f-079b-4c5d-979b-0411510c96cb","resolution":{"observed_at":"2026-08-06T13:02:28.871324Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.875639Z","title":"Reinforcement learning,","venue":null,"work_id":null,"year":1999},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.875639Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:bb47e006c6ec50df74003c2b0a51c88478991dc7f648c180d81b951f6591e3a0","observation_id":"dcc8069f-92bd-4c28-a8fa-bd400dddb976","resolution":{"observed_at":"2026-08-06T13:02:28.875639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.880151Z","title":"Physicsnerf: Physics-guided 3d reconstruction from sparse views,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.880151Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6704579e2057a3521d7dac2a2653c4526c8a945f7c901bb724557317692fc5aa","observation_id":"fb5c6197-befd-428c-a4cd-cb85b66efa7e","resolution":{"observed_at":"2026-08-06T13:02:28.880151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.884509Z","title":"Pbr-nerf: Inverse rendering with physics-based neural fields,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":94,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.884509Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:bbfe6ab678a06b578832265502f028e01a5351efff5583c9a65a43d97a579f94","observation_id":"f99a04f3-0cee-4190-ac03-044238855d2d","resolution":{"observed_at":"2026-08-06T13:02:28.884509Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.889401Z","title":"Cast: Component-aligned 3d scene reconstruction from an rgb image,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":95,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.889401Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:02d5d80d57dc6bcb8ae06e0c605137791209a974f146cedd3645bb27fe4b83c9","observation_id":"e475ce1c-ff34-4c2f-b74c-7838a936f74f","resolution":{"observed_at":"2026-08-06T13:02:28.889401Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.894051Z","title":"Sv3d: Novel multi- view synthesis and 3d generation from a single image using latent video diffusion,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":96,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.894051Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:b33ab89f6051b1e2178ec598e1ecf51189c3e9de450b8ede9270b23a1dbf632e","observation_id":"26f7958c-e65f-45bc-bbea-8ecd53848a47","resolution":{"observed_at":"2026-08-06T13:02:28.894051Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.06738","last_updated":"2024-03-11T14:03:36Z","snapshot_observed_at":"2026-08-09T00:18:19.629481Z","submitted_at":"2024-03-11T14:03:36Z","title":"V3D: Video Diffusion Models are Effective 3D Generators","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.06738","snapshot_observed_at":"2026-08-06T13:02:28.899027Z","title":"V3d: Video diffusion models are effective 3d generators,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":97,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.899027Z"},"links":{"cited_paper":"/paper/2403.06738","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:7f41415cf08a08066e3d488f5ad593cdb355a6c7128b751bae20fe46128306e7","observation_id":"4d743a00-7884-4d48-ab0d-6c448e4b69f5","resolution":{"observed_at":"2026-08-06T13:02:28.899027Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.903721Z","title":"Drea- mavatar: Text-and-shape guided 3d human avatar generation via diffusion models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.903721Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:00cd3d9eacaa3338d91518738d41b1d14034127a6077f19aa9d87b3c7534b2e3","observation_id":"bf376e61-a411-4e0d-8c80-2d5ff58a91bc","resolution":{"observed_at":"2026-08-06T13:02:28.903721Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.908275Z","title":"4d-fy: Text-to-4d generation using hybrid score distillation sampling,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":99,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.908275Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:a6a5f3fb24b9dfc8b454aebcfdb819199b514ff6efee8e5b40b06f1333503463","observation_id":"318abee8-87ef-4918-8846-05ab9913f119","resolution":{"observed_at":"2026-08-06T13:02:28.908275Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:02:28.913510Z","title":"Tc4d: Trajectory- conditioned text-to-4d generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":100,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:28.913510Z"},"links":{"citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:6d88fb48b8f145dbc456b6ad2b43aa858749610961f4c43549eee306c0cee1ee","observation_id":"6be87aea-db7a-412f-b6d9-23d6d4503812","resolution":{"observed_at":"2026-08-06T13:02:28.913510Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":98,"verified_exact":2,"verified_fuzzy":0},"total_outbound_references":300},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 100 of 300 outbound references and 9 inbound Pith citation observations for arXiv:2507.21045."}