{"as_of":"2026-08-10T00:15:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:ac7cba81d56a8e2a7d1340d6a9708d564bb29a609af9b1a4e83f2e4fd49ea153","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":30,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":30,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":30,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":30,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T04:42:05.975670Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T20:00:07.459591Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2411.14295","last_updated":"2026-05-03T12:27:01Z","snapshot_observed_at":"2026-07-06T19:53:52.079055Z","submitted_at":"2024-11-21T16:41:55Z","title":"DissolveStereo: Coarse Depth Injection for Zero-Shot Stereo Video Generation","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-23T17:17:48.961672Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2411.14295"},"observation_digest":"sha256:7a68c11c416ad0737ce3ecb1b952e7c315770156952094e3388496ebf3ed1064","observation_id":"56e8e8b7-e70e-44c4-b103-8c88dfbc9e82","resolution":{"observed_at":"2026-05-23T17:18:13.976065Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-09T04:42:05.975670Z","title":"Video depth anything: Consistent depth estimation for super- long videos","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.03465","last_updated":"2025-03-17T06:29:41Z","snapshot_observed_at":"2026-08-09T04:36:06.051909Z","submitted_at":"2025-02-05T18:59:52Z","title":"Seeing World Dynamics in a Nutshell","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-09T04:42:05.975670Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2502.03465"},"observation_digest":"sha256:d80ce5040359979e9b08071c565bad9d8b4db4fbc999b7300a2af6d74ea10936","observation_id":"8a153380-aa41-472a-b42f-d93b78d6ce3f","resolution":{"observed_at":"2026-08-09T04:42:05.975670Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-07T12:26:57.495990Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24521","last_updated":"2025-05-30T12:31:59Z","snapshot_observed_at":"2026-08-09T04:17:29.326510Z","submitted_at":"2025-05-30T12:31:59Z","title":"UniGeo: Taming Video Diffusion for Unified Consistent Geometry Estimation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T12:26:57.495990Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2505.24521"},"observation_digest":"sha256:c2f7b39e525f6df3785753ceb1b8722499f466fd94a72a989903fcae9c58f141","observation_id":"677d1a65-46d6-4a29-8c90-a851b74eb0e9","resolution":{"observed_at":"2026-08-07T12:26:57.495990Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-07T11:36:06.197375Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.01933","last_updated":"2025-07-10T17:55:35Z","snapshot_observed_at":"2026-08-09T04:18:07.584928Z","submitted_at":"2025-06-02T17:53:09Z","title":"E3D-Bench: A Benchmark for End-to-End 3D Geometric Foundation Models","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T11:36:06.197375Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2506.01933"},"observation_digest":"sha256:c94728bfa727f5e96df8d8d6e5cc6d38242cf40d909ef82d7aa915d7c5e6dcc0","observation_id":"02a6ffb5-61ac-4584-91e1-58d98f7acf98","resolution":{"observed_at":"2026-08-07T11:36:06.197375Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-07T11:14:43.647637Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.03150","last_updated":"2025-06-03T17:59:52Z","snapshot_observed_at":"2026-08-09T04:47:45.249838Z","submitted_at":"2025-06-03T17:59:52Z","title":"IllumiCraft: Unified Geometry and Illumination Diffusion for Controllable Video Generation","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T11:14:43.647637Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2506.03150"},"observation_digest":"sha256:08a993e119baaef32d53e882d19b568fca0549062b7c89e599a6eacbe08acde0","observation_id":"53500216-ec5f-460f-b62c-3f9bbeec206c","resolution":{"observed_at":"2026-08-07T11:14:43.647637Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-06T23:58:41.807576Z","title":"Video depth anything: Consistent depth estimation for super-long videos,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.15560","last_updated":"2025-07-05T14:04:22Z","snapshot_observed_at":"2026-08-09T10:26:07.606919Z","submitted_at":"2025-06-18T15:35:16Z","title":"RaCalNet: Radar Calibration Network for Sparse-Supervised Metric Depth Estimation","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:41.807576Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2506.15560"},"observation_digest":"sha256:4b5b4b4f92662f0d4c2f058008bd9eaadc8ae022833ebfcf6af0ececbce7d743","observation_id":"118f8165-5b8a-471a-a0ce-a0ed023223e8","resolution":{"observed_at":"2026-08-06T23:58:41.807576Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-06T21:56:09.428813Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.23135","last_updated":"2025-06-29T08:19:45Z","snapshot_observed_at":"2026-08-07T02:31:43.164112Z","submitted_at":"2025-06-29T08:19:45Z","title":"RoboScape: Physics-informed Embodied World Model","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T21:56:09.428813Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2506.23135"},"observation_digest":"sha256:ce60ec278ff5108ce14327d24f8da4329ddb1de6481f95c01c8740f70a00dbc6","observation_id":"ddf523f2-5879-4fe0-9116-f32619e7454c","resolution":{"observed_at":"2026-08-06T21:56:09.428813Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-06T18:27:06.939745Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.08344","last_updated":"2025-08-05T08:39:55Z","snapshot_observed_at":"2026-08-09T11:31:18.200127Z","submitted_at":"2025-07-11T06:45:42Z","title":"MM-Gesture: Towards Precise Micro-Gesture Recognition through Multimodal Fusion","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T18:27:06.939745Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2507.08344"},"observation_digest":"sha256:58cdafd601bd8fb19f5e9936864363d26a478c7b5920e92aa0e912c5d2afa42a","observation_id":"ecf4f2db-7073-4213-beb6-daefed166cac","resolution":{"observed_at":"2026-08-06T18:27:06.939745Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-06T17:34:05.903479Z","title":"Video depth anything: Consistent depth estimation for super-long videos,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.10749","last_updated":"2025-07-14T19:16:13Z","snapshot_observed_at":"2026-08-09T09:40:07.554885Z","submitted_at":"2025-07-14T19:16:13Z","title":"RCG: Safety-Critical Scenario Generation for Robust Autonomous Driving via Real-World Crash Grounding","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T17:34:05.903479Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2507.10749"},"observation_digest":"sha256:62200c00e278d6741dc11fd13634dec2856c9fc4391cb0f51d33b3e97933019e","observation_id":"98217f03-2f15-440f-9424-261dcd1fb3f8","resolution":{"observed_at":"2026-08-06T17:34:05.903479Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-06T16:49:48.565573Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.12462","last_updated":"2025-07-19T02:07:12Z","snapshot_observed_at":"2026-08-08T08:02:10.460729Z","submitted_at":"2025-07-16T17:59:03Z","title":"SpatialTrackerV2: 3D Point Tracking Made Easy","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T16:49:48.565573Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2507.12462"},"observation_digest":"sha256:e36c9bae6fc1f526a570bb4ac7c783f1eba0140456a8245f0e7033a4ee5150f3","observation_id":"c3c5ea58-a021-4c05-b163-8065c0f45e55","resolution":{"observed_at":"2026-08-06T16:49:48.565573Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-06T13:02:29.086195Z","title":"Video depth anything: Consistent depth estimation for super- long videos,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21045","last_updated":"2025-08-03T14:18:19Z","snapshot_observed_at":"2026-08-07T20:14:50.999165Z","submitted_at":"2025-07-28T17:59:02Z","title":"Reconstructing 4D Spatial Intelligence: A Survey","version":2},"reference_index":135,"source":"pdf_text","source_observed_at":"2026-08-06T13:02:29.086195Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2507.21045"},"observation_digest":"sha256:a0c5a97a5fef9d298b8fee5191b04d9d45e5a3b316ff13fa9f69134fa4f9a029","observation_id":"96348a5e-db49-4c7b-86cc-e3b688b5be64","resolution":{"observed_at":"2026-08-06T13:02:29.086195Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2508.10934","last_updated":"2025-08-12T18:39:13Z","snapshot_observed_at":"2026-07-06T22:13:09.586800Z","submitted_at":"2025-08-12T18:39:13Z","title":"ViPE: Video Pose Engine for 3D Geometric Perception","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-16T16:41:08.620285Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2508.10934"},"observation_digest":"sha256:45e92282cee9b6cd51d828002389b4f40c5c3ec7f8bb1a347822c28bdff7c8de","observation_id":"3a193b00-fc7b-44d1-a0ab-53127283cb3e","resolution":{"observed_at":"2026-05-16T16:41:08.727344Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-05T13:46:44.930591Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.00361","last_updated":"2025-08-30T04:53:32Z","snapshot_observed_at":"2026-08-09T09:27:18.513708Z","submitted_at":"2025-08-30T04:53:32Z","title":"Generative Visual Foresight Meets Task-Agnostic Pose Estimation in Robotic Table-Top Manipulation","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-05T13:46:44.930591Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2509.00361"},"observation_digest":"sha256:ee8db8b3185ccb4037b650f03831d0fe30b494d8b1d53e26af804704160e169e","observation_id":"c11040f0-350d-4e24-a23f-8f6aa3498ec0","resolution":{"observed_at":"2026-08-05T13:46:44.930591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-04T11:35:58.419641Z","title":"Video depth anything: Consistent depth estimation for super-long videos,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.04074","last_updated":"2026-05-29T20:05:57Z","snapshot_observed_at":"2026-08-07T16:22:00.378390Z","submitted_at":"2025-10-05T07:26:40Z","title":"Feedback Matters: Augmenting Autonomous Dissection with Visual and Topological Feedback","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-04T11:35:58.419641Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2510.04074"},"observation_digest":"sha256:9607d18318fbce0810a82d9878af237b05d8df3da56b554b0e0fba2ad50a1bd9","observation_id":"96c1b55d-adfa-44e5-9ef0-2d72c5c1f332","resolution":{"observed_at":"2026-08-04T11:35:58.419641Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2512.11988","last_updated":"2026-04-19T00:55:15Z","snapshot_observed_at":"2026-07-06T22:38:54.914862Z","submitted_at":"2025-12-12T19:11:11Z","title":"CARI4D: Category Agnostic 4D Reconstruction of Human-Object Interaction","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-16T22:38:01.008280Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2512.11988"},"observation_digest":"sha256:63f40c329d4f587f4e6b077b1177f249b00fa1bb41161c12634e906a745873a8","observation_id":"6be966d9-301c-4180-901d-c6d8cd0236ce","resolution":{"observed_at":"2026-05-16T22:38:37.596051Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-13T22:13:53.383917Z","title":"arXiv:2501.12375 (2025) 5, 8, 9","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.19048","last_updated":"2026-07-05T07:55:03Z","snapshot_observed_at":"2026-08-07T14:01:57.520825Z","submitted_at":"2026-03-19T15:44:39Z","title":"Measuring 3D Spatial Geometric Consistency in Dynamic Video Generation","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-13T22:13:53.383917Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2603.19048"},"observation_digest":"sha256:43fda1296741991b47f3853c648f60a13ce982ef5310871bbf147a1dac4bcf22","observation_id":"e0f3cdf7-c86a-49ef-a33e-0492f7d05222","resolution":{"observed_at":"2026-07-13T22:13:53.383917Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2604.04016","last_updated":"2026-04-05T08:27:28Z","snapshot_observed_at":"2026-08-08T14:25:59.315820Z","submitted_at":"2026-04-05T08:27:28Z","title":"HOIGS: Human-Object Interaction Gaussian Splatting","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-13T17:24:50.043908Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2604.04016"},"observation_digest":"sha256:df0022ac20db5bc275dd9df441358fcbd0e44ea5247028c34bc80a5135370c86","observation_id":"f87dda88-df7c-41f9-ba6e-74ead34e61c1","resolution":{"observed_at":"2026-05-13T17:28:02.544902Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2604.04974","last_updated":"2026-06-03T17:05:28Z","snapshot_observed_at":"2026-08-05T03:49:14.637609Z","submitted_at":"2026-04-04T15:37:11Z","title":"From Video to Control: A Survey of Learning Manipulation Interfaces from Temporal Visual Data","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-05-13T17:02:18.358675Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2604.04974"},"observation_digest":"sha256:b0b75a1f700010457724d5f218c83e462dfa53f31c24ab822642466fe8202aed","observation_id":"2bb3e8db-2aea-477e-b6f7-934d607c3940","resolution":{"observed_at":"2026-05-13T17:03:01.206201Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2604.09352","last_updated":"2026-04-10T14:24:07Z","snapshot_observed_at":"2026-07-06T22:58:12.847181Z","submitted_at":"2026-04-10T14:24:07Z","title":"LuMon: A Comprehensive Benchmark and Development Suite with Novel Datasets for Lunar Monocular Depth Estimation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-10T17:18:51.726719Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2604.09352"},"observation_digest":"sha256:3aeb59e6726672ad16bdbf3f748777ab65015062b020b182ce4250e7250628ed","observation_id":"fbd445bf-0125-4f54-bbcb-db9b20b08ad8","resolution":{"observed_at":"2026-05-11T07:06:00.772158Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2604.14556","last_updated":"2026-04-16T02:39:15Z","snapshot_observed_at":"2026-07-06T23:02:18.425249Z","submitted_at":"2026-04-16T02:39:15Z","title":"Controllable Video Object Insertion via Multiview Priors","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T11:44:17.033051Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2604.14556"},"observation_digest":"sha256:8947716a8ed95ee67e45486d1fa0418a07748060032de43cf1432aaf6c847969","observation_id":"2c3a4796-b1a6-4246-94e3-9497a21452a6","resolution":{"observed_at":"2026-05-10T11:45:20.659728Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2604.22160","last_updated":"2026-06-26T16:26:41Z","snapshot_observed_at":"2026-08-04T04:08:04.701132Z","submitted_at":"2026-04-24T02:18:18Z","title":"GenMatter: Perceiving Physical Objects with Generative Matter Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-08T12:53:18.504365Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2604.22160"},"observation_digest":"sha256:556d9b9869cd97f5a6ba0b5f7aa6ca784924a27fcfb17e5d9cf11e416af5de71","observation_id":"3d744e5f-9ae0-4b6c-936c-3826d068ad69","resolution":{"observed_at":"2026-05-11T19:01:18.256786Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2604.22160","last_updated":"2026-06-26T16:26:41Z","snapshot_observed_at":"2026-08-04T04:08:04.701132Z","submitted_at":"2026-04-24T02:18:18Z","title":"GenMatter: Perceiving Physical Objects with Generative Matter Models","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-04T19:55:05.116930Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2604.22160"},"observation_digest":"sha256:319a820f23795091eada8d7c8ac6301487b80bd002191cb9f9edd8342053471f","observation_id":"a9a13e12-5f04-48e4-b699-7817f3f75bbd","resolution":{"observed_at":"2026-07-04T20:00:07.462600Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2604.24718","last_updated":"2026-04-27T17:29:22Z","snapshot_observed_at":"2026-07-06T23:10:42.926677Z","submitted_at":"2026-04-27T17:29:22Z","title":"WildLIFT: Lifting monocular drone video to 3D for species-agnostic wildlife monitoring","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-08T04:26:53.225501Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2604.24718"},"observation_digest":"sha256:23ed0938d600a609a85bd5592d45ab3f68831b54eb6f953fba8607b1acf576bf","observation_id":"86fb3299-e315-4310-b09a-74670ef9662b","resolution":{"observed_at":"2026-05-11T21:46:40.192466Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2605.00658","last_updated":"2026-05-01T13:40:56Z","snapshot_observed_at":"2026-07-06T23:14:01.852096Z","submitted_at":"2026-05-01T13:40:56Z","title":"UniVidX: A Unified Multimodal Framework for Versatile Video Generation via Diffusion Priors","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-05-09T20:05:21.723724Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2605.00658"},"observation_digest":"sha256:49722cf823e8c320dd474d6275201d72c5d42e64328e1568c494b35272c616b7","observation_id":"05596cf5-ac2a-47a5-9472-5b6e82108ab7","resolution":{"observed_at":"2026-05-11T15:26:07.902348Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2605.23098","last_updated":"2026-05-21T23:08:42Z","snapshot_observed_at":"2026-07-06T23:33:19.375527Z","submitted_at":"2026-05-21T23:08:42Z","title":"UfM*: Uncertainty from Motion* for DNN Depth Estimation Using Gaussians","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-25T05:16:28.339361Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2605.23098"},"observation_digest":"sha256:9b734d9e67f8bbf71560e17fdab8fe028e4cdc27a04027e6f07e8f3dd2215eaf","observation_id":"dd2f7c12-0bcf-428f-afbb-cffd18c90535","resolution":{"observed_at":"2026-05-25T05:16:39.146720Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2605.25308","last_updated":"2026-05-25T00:13:15Z","snapshot_observed_at":"2026-07-06T23:35:16.647496Z","submitted_at":"2026-05-25T00:13:15Z","title":"Stabilizing Streaming Video Geometry via Dynamic Feature Normalization","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-29T23:15:20.254402Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2605.25308"},"observation_digest":"sha256:2dcb5153f9179d59fe6adb7a8c433949d359ea493f680e9b6b341887e28a8869","observation_id":"faf378b5-508b-4f61-bdbc-ef01898e89df","resolution":{"observed_at":"2026-06-29T23:24:02.104002Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2606.26410","last_updated":"2026-06-24T22:06:35Z","snapshot_observed_at":"2026-08-07T03:54:49.094588Z","submitted_at":"2026-06-24T22:06:35Z","title":"Neural Voxel Dynamics: Learning Implicit 3D Physics via Volumetric Feature Advection","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-26T01:17:06.587807Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2606.26410"},"observation_digest":"sha256:52883b3a4f21936ef98f71b74725eb0ac66612fc76e282ff9b008e0f65cfee78","observation_id":"0bbdd005-672f-4bdc-a482-a51255ba26fe","resolution":{"observed_at":"2026-07-04T15:49:58.020111Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":"2501.12375","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-07-04T20:00:07.459591Z","title":"Video depth anything: Consistent depth estimation for super-long videos","venue":null,"work_id":"83e140fd-6eb9-4d1c-a155-d86ac7e58124","year":2025},"citing_paper":{"arxiv_id":"2606.28128","last_updated":"2026-06-26T14:30:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-06-26T14:30:18Z","title":"PhysisForcing: Physics Reinforced World Simulator for Robotic Manipulation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-29T04:34:38.286863Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2606.28128"},"observation_digest":"sha256:ea2d62bf7e034a870b2a41efcfc1e5c023b6482e3a345611fda3c74665db6f85","observation_id":"1c980236-bde6-4537-a8d9-05d5838c4609","resolution":{"observed_at":"2026-06-29T20:03:56.956882Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-02T06:16:27.239660Z","title":"Video depth anything: Consistent depth estimation for super-long videos.arXiv preprint arXiv:2501.12375, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.12993","last_updated":"2026-07-15T17:13:15Z","snapshot_observed_at":"2026-08-07T20:19:58.454005Z","submitted_at":"2026-07-14T17:45:50Z","title":"X-Lens: Real-Time Metric Depth Estimation with Heterogeneous Cameras","version":2},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-02T06:16:27.239660Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2607.12993"},"observation_digest":"sha256:e4bb4a663ab8df79054e39895a7f3a5308bcc4bd499df762b8997884f648a512","observation_id":"7b749165-ee2a-424f-81b7-5157dd18f758","resolution":{"observed_at":"2026-08-02T06:16:27.239660Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12375","snapshot_observed_at":"2026-08-06T18:39:26.592889Z","title":"Video depth anything: Consistent depth estimation for super-long videos.arXiv:2501.12375, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.04701","last_updated":"2026-08-05T11:10:52Z","snapshot_observed_at":"2026-08-08T23:12:19.042072Z","submitted_at":"2026-08-05T11:10:52Z","title":"UniWorld-View: Large-Baseline View Synthesis via Video Diffusion Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T18:39:26.592889Z"},"links":{"cited_paper":"/paper/2501.12375","citing_paper":"/paper/2608.04701"},"observation_digest":"sha256:c56716ce778d46deb2d81d83bd4bfc844affaa34c4e811bfa7114c8c0fcc016c","observation_id":"a1ac38ed-ec1a-4c70-8835-a89107fa843d","resolution":{"observed_at":"2026-08-06T18:39:26.592889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2501.12375/citation-record","integrity":"/paper/2501.12375/integrity","json":"/paper/2501.12375/citation-record.json","paper":"/paper/2501.12375"},"outbound":[],"paper":{"arxiv_id":"2501.12375","last_updated":"2025-06-15T03:50:49Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T20:24:03.105038Z","submitted_at":"2025-01-21T18:53:30Z","title":"Video Depth Anything: Consistent Depth Estimation for Super-Long Videos"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 30 inbound Pith citation observations for arXiv:2501.12375."}