{"as_of":"2026-08-17T19:33:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:14a5c3ac86d9c8e076a772de09f239ba63a3d772ced2be7e9801f0bcbb319b2b","coverage":[{"denominator":48,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":48,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T18:44:55.425136Z","state":"measured"},{"denominator":63,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":63,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":15,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":15,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T05:06:51.863744Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-03T17:58:47.750659Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-08-03T15:28:54.859365Z","title":"Vmem: Consistent interactive video scene gen- eration with surfel-indexed view memory.arXiv preprint arXiv:2506.18903, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2512.17796","last_updated":"2026-06-24T08:52:20Z","snapshot_observed_at":"2026-08-13T02:26:05.415227Z","submitted_at":"2025-12-18T18:59:18Z","title":"CustomX: Unified Character, Action, and Scene Customization in Video World Models","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-03T15:28:54.859365Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2512.17796"},"observation_digest":"sha256:3c72375973da3a850c7aa0f8ba6578c9e502b5bd8af46dd94c6e83b042eae557","observation_id":"e0e09699-5170-49d9-b20d-3498f807d1b4","resolution":{"observed_at":"2026-08-03T15:28:54.859365Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-08-02T20:36:11.425146Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.22960","last_updated":"2026-06-29T11:22:33Z","snapshot_observed_at":"2026-08-14T12:23:22.280746Z","submitted_at":"2026-02-26T12:54:46Z","title":"UCM: Unified Modeling of Camera Control and Memory with Time-aware Positional Encoding Warping for World Models","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-02T20:36:11.425146Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2602.22960"},"observation_digest":"sha256:26af4bc77a29a947d6fcb2563b033a251cac21bb817834cd2286b60e9385dd47","observation_id":"9e0238c0-1202-4efe-ab8c-69888bb3f6cf","resolution":{"observed_at":"2026-08-02T20:36:11.425146Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-08-02T17:54:05.655438Z","title":"arXiv preprint arXiv:2506.18903 (2025) 2, 3, 5, 8, 13, 25, 27, 28","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.19235","last_updated":"2026-07-17T06:42:00Z","snapshot_observed_at":"2026-08-14T02:16:35.171898Z","submitted_at":"2026-03-19T17:59:58Z","title":"Generation Models Know Space: Unleashing Implicit 3D Priors for Scene Understanding","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-02T17:54:05.655438Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2603.19235"},"observation_digest":"sha256:737705c91f22d26625705f597c19a67f0bd219bcb7ab485606325d3fcfcd5f20","observation_id":"62eb21f5-d460-4442-ada8-5f9fa0bdc3ae","resolution":{"observed_at":"2026-08-02T17:54:05.655438Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-08-03T02:30:19.482363Z","title":"arXiv preprint arXiv:2506.18903 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.23413","last_updated":"2026-07-31T11:13:43Z","snapshot_observed_at":"2026-08-10T10:39:59.339425Z","submitted_at":"2026-03-24T16:45:40Z","title":"I3DM: Implicit 3D-aware Memory Retrieval and Injection for Consistent Video Scene Generation","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-03T02:30:19.482363Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2603.23413"},"observation_digest":"sha256:7668a72601ceaff8163b5f492ec3af2da239ea504e07c931390b15ea42abf273","observation_id":"cf8a707d-f829-4682-9576-03af98789a77","resolution":{"observed_at":"2026-08-03T02:30:19.482363Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":"2506.18903","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-07-03T17:58:47.750659Z","title":"Vmem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":"8c6afc6c-7c32-41eb-87c1-e51bca941b63","year":2025},"citing_paper":{"arxiv_id":"2604.06339","last_updated":"2026-04-07T18:17:05Z","snapshot_observed_at":"2026-08-16T06:51:21.904839Z","submitted_at":"2026-04-07T18:17:05Z","title":"Evolution of Video Generative Foundations","version":1},"reference_index":287,"source":"pdf_text","source_observed_at":"2026-05-10T18:41:38.616611Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2604.06339"},"observation_digest":"sha256:b5fa214c6714d2f130263ca51dfb47dd60c0513ab38689a0d55b63264d86de34","observation_id":"d4919605-74b9-4e4c-abb6-af4493587453","resolution":{"observed_at":"2026-05-11T00:05:51.422882Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":"2506.18903","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-07-03T17:58:47.750659Z","title":"Vmem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":"8c6afc6c-7c32-41eb-87c1-e51bca941b63","year":2025},"citing_paper":{"arxiv_id":"2604.14268","last_updated":"2026-04-15T17:59:17Z","snapshot_observed_at":"2026-08-12T17:33:25.813833Z","submitted_at":"2026-04-15T17:59:17Z","title":"HY-World 2.0: A Multi-Modal World Model for Reconstructing, Generating, and Simulating 3D Worlds","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-10T13:45:24.961208Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2604.14268"},"observation_digest":"sha256:203bd73ed6e270ed7c9c323b91cc8b53d786926de40ef594088d8213f542808b","observation_id":"539fda55-bd04-4f4a-a8c2-d659a81ca089","resolution":{"observed_at":"2026-05-10T13:45:27.241056Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":"2506.18903","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-07-03T17:58:47.750659Z","title":"Vmem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":"8c6afc6c-7c32-41eb-87c1-e51bca941b63","year":2025},"citing_paper":{"arxiv_id":"2605.22718","last_updated":"2026-05-21T16:55:04Z","snapshot_observed_at":"2026-08-14T02:53:00.466176Z","submitted_at":"2026-05-21T16:55:04Z","title":"WorldKV: Efficient World Memory with World Retrieval and Compression","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-22T06:20:18.637998Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2605.22718"},"observation_digest":"sha256:167c8dd80eca7896647877da45a254c97029353993449704647d9f89134a65e1","observation_id":"48857495-eb75-4a88-803c-68b0c6c75231","resolution":{"observed_at":"2026-05-22T06:21:10.125264Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":"2506.18903","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-07-03T17:58:47.750659Z","title":"Vmem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":"8c6afc6c-7c32-41eb-87c1-e51bca941b63","year":2025},"citing_paper":{"arxiv_id":"2605.26316","last_updated":"2026-05-25T20:13:16Z","snapshot_observed_at":"2026-08-01T03:37:55.526828Z","submitted_at":"2026-05-25T20:13:16Z","title":"E$^3$C: Video Generation with 3D Environmental Memory and Ego-Exo Human Pose Control","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-29T22:24:38.389294Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2605.26316"},"observation_digest":"sha256:2c4c97791c3c4ce4c250f54ff7d9a6e0384bd77dbb9840455234ff5f1e386776","observation_id":"1f5e6537-9dde-4498-955a-9bad94af393a","resolution":{"observed_at":"2026-06-29T22:34:02.493682Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":"2506.18903","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-07-03T17:58:47.750659Z","title":"Vmem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":"8c6afc6c-7c32-41eb-87c1-e51bca941b63","year":2025},"citing_paper":{"arxiv_id":"2605.30855","last_updated":"2026-06-01T09:19:36Z","snapshot_observed_at":"2026-08-16T22:18:33.146183Z","submitted_at":"2026-05-29T05:21:33Z","title":"Robust Dreamer: Deviation-Aware Latent Gaussian Memory for Action-Controlled AR Video Generation","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-06-28T22:56:35.530309Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2605.30855"},"observation_digest":"sha256:740da0f2ee4e121b5680b5fb3c8981c1341ea24ffed4947e4d0d38440b908a50","observation_id":"d6e19b2a-9e74-4d4b-a3e8-ec79cef350f3","resolution":{"observed_at":"2026-06-28T23:02:46.964385Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":"2506.18903","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-07-03T17:58:47.750659Z","title":"Vmem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":"8c6afc6c-7c32-41eb-87c1-e51bca941b63","year":2025},"citing_paper":{"arxiv_id":"2605.31336","last_updated":"2026-05-29T14:17:59Z","snapshot_observed_at":"2026-08-08T11:08:50.003696Z","submitted_at":"2026-05-29T14:17:59Z","title":"DecMem: Towards Minute-Long Consistent World Generation with Decoupled Memory","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-28T22:38:10.473453Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2605.31336"},"observation_digest":"sha256:6b6c749bc9365689d2456228d54aae2977ef0eb612b6b1d17dd1aa172fe78f66","observation_id":"d6ad3910-ebaf-4067-bdfe-de77ebedf16d","resolution":{"observed_at":"2026-06-28T22:42:46.841506Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":"2506.18903","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-07-03T17:58:47.750659Z","title":"Vmem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":"8c6afc6c-7c32-41eb-87c1-e51bca941b63","year":2025},"citing_paper":{"arxiv_id":"2606.02753","last_updated":"2026-06-01T18:20:20Z","snapshot_observed_at":"2026-08-16T00:06:16.567037Z","submitted_at":"2026-06-01T18:20:20Z","title":"MetaWorld: Scaling Multi-Agent Video World Model from Single-view Video Data","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-28T14:52:30.406683Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2606.02753"},"observation_digest":"sha256:e01c10cdadf708b3f155e09c02c500669a74251f89ee221d6be7fedffe87672e","observation_id":"d23b27c4-f2da-4dae-8c53-36a2b1dd05e3","resolution":{"observed_at":"2026-07-01T22:56:20.197617Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":"2506.18903","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-07-03T17:58:47.750659Z","title":"Vmem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":"8c6afc6c-7c32-41eb-87c1-e51bca941b63","year":2025},"citing_paper":{"arxiv_id":"2606.09828","last_updated":"2026-06-08T17:59:54Z","snapshot_observed_at":"2026-08-14T17:26:29.934753Z","submitted_at":"2026-06-08T17:59:54Z","title":"Latent Spatial Memory for Video World Models","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-06-27T16:47:42.761342Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2606.09828"},"observation_digest":"sha256:10af2df11f73cfb946245e3518f6a78160caccca530c0218b804b835722091dd","observation_id":"9abdcd1c-0cdf-44fa-acaa-293c453449d1","resolution":{"observed_at":"2026-07-03T01:07:30.579093Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":"2506.18903","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-07-03T17:58:47.750659Z","title":"Vmem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":"8c6afc6c-7c32-41eb-87c1-e51bca941b63","year":2025},"citing_paper":{"arxiv_id":"2606.16449","last_updated":"2026-06-16T02:33:07Z","snapshot_observed_at":"2026-08-17T01:09:27.291570Z","submitted_at":"2026-06-15T09:20:32Z","title":"PermaVid: Consistent Video Generation Across Edits via Disentangled Context Memory","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-27T03:23:58.455778Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2606.16449"},"observation_digest":"sha256:68da47aa010e3da105f5aa0cda284592819f381db2779eafa826bf5b33f1a148","observation_id":"80c55eca-3c03-4ad1-932c-dcd2be0be63a","resolution":{"observed_at":"2026-07-03T17:58:47.752156Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":"2506.18903","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-07-03T17:58:47.750659Z","title":"Vmem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":"8c6afc6c-7c32-41eb-87c1-e51bca941b63","year":2025},"citing_paper":{"arxiv_id":"2606.31734","last_updated":"2026-06-30T14:31:32Z","snapshot_observed_at":"2026-08-07T09:44:34.339704Z","submitted_at":"2026-06-30T14:31:32Z","title":"MemLearner: Learning to Query Context memory for Video World Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-07-01T05:30:56.140465Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2606.31734"},"observation_digest":"sha256:b0db2f2715daebfe3c9016f5867f1152e93a9d3c4cdf49f03434814a1bedef6f","observation_id":"f7f4d308-7a1d-48e0-b60e-b026f81abe37","resolution":{"observed_at":"2026-07-01T10:25:41.888623Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.18903","snapshot_observed_at":"2026-08-10T05:06:51.863744Z","title":"VMem: Consistent interactive video scene generation with surfel-indexed view memory","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.07408","last_updated":"2026-08-07T16:55:57Z","snapshot_observed_at":"2026-08-16T13:12:28.441632Z","submitted_at":"2026-08-07T16:55:57Z","title":"Addressable Memory for Video World Models","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-10T05:06:51.863744Z"},"links":{"cited_paper":"/paper/2506.18903","citing_paper":"/paper/2608.07408"},"observation_digest":"sha256:caff46c0bd52fad17e6a44576017480e377e0153f952677ba567be053ead6401","observation_id":"de891703-bf83-4994-a4e0-ab76812214f2","resolution":{"observed_at":"2026-08-10T05:06:51.863744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2506.18903/citation-record","integrity":"/paper/2506.18903/integrity","json":"/paper/2506.18903/citation-record.json","paper":"/paper/2506.18903"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2408.00653","last_updated":"2024-08-01T15:41:57Z","snapshot_observed_at":"2026-08-17T05:50:17.640415Z","submitted_at":"2024-08-01T15:41:57Z","title":"SF3D: Stable Fast 3D Mesh Reconstruction with UV-unwrapping and Illumination Disentanglement","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.00653","snapshot_observed_at":"2026-08-15T18:44:55.208668Z","title":"SF3D: Stable fast 3D mesh reconstruction with uv-unwrapping and illumination disentanglement","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.208668Z"},"links":{"cited_paper":"/paper/2408.00653","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:09a9c0440ae853acbbc124af8113e5adf8f32cae2e627c4d2bc2512edc8a2a0f","observation_id":"fd6c16f4-997a-4808-b95f-a6bca4642e77","resolution":{"observed_at":"2026-08-15T18:44:55.208668Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.214283Z","title":"Ge- nie: Generative interactive environments","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.214283Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:b38a839b9aad7f1a460b18a03708b9079e9c8636624fedc6a908a8801770414e","observation_id":"add98839-8a72-4a22-9955-ec86ad76498c","resolution":{"observed_at":"2026-08-15T18:44:55.214283Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:56.078419Z","title":"Mvsnerf: Fast general- izable radiance field reconstruction from multi-view stereo","venue":null,"work_id":"1b2ecc01-4681-44ec-8701-2c2575ce6c57","year":2021},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.219075Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:e6ffdfac35538ead9718720f6885bbdc0bd03b5dec1b0e288c39821f77a97f7c","observation_id":"77b16276-def5-471b-b7a1-1c9508db77af","resolution":{"observed_at":"2026-08-15T18:44:56.083602Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:56.064364Z","title":"Single-Stage Diffusion NeRF: A Unified Approach to 3D Generation and Reconstruction,","venue":null,"work_id":"5acf96fe-bc7d-4c19-a172-02dbb302b3c3","year":null},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.225074Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:1edb3705f2e792694411aa9ec338f7b1c04b622b88acb31346a959d828656f2e","observation_id":"c6664651-6cd3-49d8-9f1f-a86eb188dc0c","resolution":{"observed_at":"2026-08-15T18:44:56.068808Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.14627","last_updated":"2024-07-18T13:10:22Z","snapshot_observed_at":"2026-08-16T14:07:34.573303Z","submitted_at":"2024-03-21T17:59:58Z","title":"MVSplat: Efficient 3D Gaussian Splatting from Sparse Multi-View Images","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.14627","snapshot_observed_at":"2026-08-15T18:44:55.234369Z","title":"Mvsplat: Efficient 3d gaussian splatting from sparse multi-view images","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.234369Z"},"links":{"cited_paper":"/paper/2403.14627","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:eb812b0da9995728feb061aaefb1c3593a66911e680557de28da387863fe0eda","observation_id":"5f623029-d7ba-4e32-83dc-5c66ddb746a6","resolution":{"observed_at":"2026-08-15T18:44:55.234369Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.239841Z","title":"Mvsplat360: Feed-forward 360 scene synthesis from sparse views","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.239841Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:d76514c9c55bf93633a1ede07a2477a0848812434a50860acb65da808fb0fa44","observation_id":"2443616f-1e82-4804-9408-9f3a9f1f7e00","resolution":{"observed_at":"2026-08-15T18:44:55.239841Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.01133","last_updated":"2023-05-30T08:52:45Z","snapshot_observed_at":"2026-08-17T16:02:00.474920Z","submitted_at":"2023-02-02T14:47:19Z","title":"SceneScape: Text-Driven Consistent Scene Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.01133","snapshot_observed_at":"2026-08-15T18:44:55.244278Z","title":"Scenescape: Text-driven consistent scene generation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.244278Z"},"links":{"cited_paper":"/paper/2302.01133","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:2a3beb7c3e4959f3a56583b975a7243863400c38261d6b21b88242dfb5a01d0a","observation_id":"7ce0cb6b-711b-4c85-a335-48872021ae75","resolution":{"observed_at":"2026-08-15T18:44:55.244278Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:56.040843Z","title":"Srinivasan, Jonathan T","venue":null,"work_id":"e10ea5d6-aecc-45d5-af81-10fa7afbfdba","year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.248946Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:3631843289b34e2b28e427b32e009f99db64f6bf494eeb55f9b09d58e7129a8a","observation_id":"29a4e433-366f-4979-9be6-4e0873d68e30","resolution":{"observed_at":"2026-08-15T18:44:56.045240Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.02101","last_updated":"2025-03-13T18:35:06Z","snapshot_observed_at":"2026-08-16T17:39:09.604516Z","submitted_at":"2024-04-02T16:52:41Z","title":"CameraCtrl: Enabling Camera Control for Text-to-Video Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.02101","snapshot_observed_at":"2026-08-15T18:44:55.253743Z","title":"Cameractrl: Enabling camera control for text-to-video generation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.253743Z"},"links":{"cited_paper":"/paper/2404.02101","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:5305c0b3c5febf77f912c0322148200d20900fdbd119afb893ef771531c6b448","observation_id":"b8b79e16-57b8-43f2-8fa6-63b001bdce6f","resolution":{"observed_at":"2026-08-15T18:44:55.253743Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.258500Z","title":"Gans trained by a two time-scale update rule converge to a local nash equilib- rium","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.258500Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:c723fd17032b848f9f6e1f1ff95bc23441f939bf814bce8b70081e49ce87a4c3","observation_id":"8b3104ac-bd76-4f68-a3cb-b9466745bf82","resolution":{"observed_at":"2026-08-15T18:44:55.258500Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:56.017896Z","title":"Text2room: Extracting textured 3d meshes from 2d text-to-image models","venue":null,"work_id":"62594e3b-aaa0-47b6-88b7-b2299aa4a8f5","year":2023},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.263360Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:9ff25a1962e5bad7dc73f829690b4e9dbb1e38f8ccb31659e5712ef3b221be1e","observation_id":"97a3b91b-bbef-4ea5-9440-dea261f7a1f6","resolution":{"observed_at":"2026-08-15T18:44:56.022683Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:56.003665Z","title":"LRM: Large reconstruction model for single image to 3D","venue":null,"work_id":"145539d6-937d-4fbc-a1e7-b32ce50ebf86","year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.267806Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:16e4948aeb7ba00d3fbaf6d90694fc3e839347f8d07bf3afaa25219e031f3c64","observation_id":"a96b9ace-a67e-419f-9a15-1e9e4497eaf8","resolution":{"observed_at":"2026-08-15T18:44:56.008118Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.989022Z","title":"Hu, Yelong Shen, Phillip Wallis, Zeyuan Allen- Zhu, Yuanzhi Li, Shean Wang, Lu Wang, and Weizhu Chen","venue":null,"work_id":"c6577461-f60d-4499-bf8a-d50542b26b71","year":2022},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.272032Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:907fabf33b822d565d41bffbbe0f0d468285470c6f8edbc26e0a6de40e3c30b3","observation_id":"12f4eec1-c7d0-4b51-98c8-b948f44b65a8","resolution":{"observed_at":"2026-08-15T18:44:55.993610Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.276226Z","title":"Tanks and temples: Benchmarking large-scale scene reconstruction","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.276226Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:d5724a40279347097d1c40a34b944f71ad1382e7a9b823785b4345880abfa6d5","observation_id":"3fa5edc2-a6f7-4b3e-82f1-f33edd827b13","resolution":{"observed_at":"2026-08-15T18:44:55.276226Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.965513Z","title":"Simple and effective synthesis of in- door 3d scenes","venue":null,"work_id":"7d1ed75a-519c-4862-97ad-92eb99a4d1ca","year":2023},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.280486Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:3d44e3d697e741f99f3ef77b8b39a12363da996679c02435c46ba285b997fb42","observation_id":"76bf3b31-54ff-4263-a706-96c6e567517a","resolution":{"observed_at":"2026-08-15T18:44:55.970046Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.951119Z","title":"Infinite na- ture: Perpetual view generation of natural scenes from a sin- gle image","venue":null,"work_id":"82540a2f-8def-4bcb-bd01-84d2701f0464","year":2021},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.284750Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:aab01bb8f33cb3b0bbe63fd57693260e3077ebea8e2417b8469d62d0c975845c","observation_id":"522ef07d-661d-470f-82bd-26074ea1104b","resolution":{"observed_at":"2026-08-15T18:44:55.955697Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.936382Z","title":"RealFusion: 360 reconstruction of any ob- ject from a single image","venue":null,"work_id":"cf60e37c-f875-4146-8bcd-5011f5ea4fbc","year":null},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.289240Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:d43dc91ba97bc3430790023875cb3e006c09d072921b32ac8601b2ef9170d329","observation_id":"3f2d1189-5ab2-4e4d-bebc-a563a1538555","resolution":{"observed_at":"2026-08-15T18:44:55.941018Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.920983Z","title":"Multidiff: Consistent novel view synthesis from a single image","venue":null,"work_id":"73410973-0e6f-41d1-a3ab-42cf9c88ca8a","year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.293669Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:fa503fd4f2d72f2666a99822c6d344fb78c9604fc7a11764307f62109a180093","observation_id":"cd710309-14a7-4aa2-ada5-63548ab43bc2","resolution":{"observed_at":"2026-08-15T18:44:55.926091Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.905687Z","title":"Reg- nerf: Regularizing neural radiance fields for view synthesis from sparse inputs","venue":null,"work_id":"f9622e51-0981-4606-acb1-3eee1e98f759","year":2022},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.298196Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:3476858c716e8faaadf3d4b7875f7d0b5413c4f235798166c4ffe670902f81a6","observation_id":"f579dda6-066d-473c-9442-eb29520baa37","resolution":{"observed_at":"2026-08-15T18:44:55.910463Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.890922Z","title":"Genie 2: A large-scale foundation world model, 2024","venue":null,"work_id":"ebf381d4-129f-46b5-b86d-9db23f93c0dc","year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.302459Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:4c5ef705750c228fc5ee098e52b1c2a42ee7a89c0ef8592e97e70a80f1b85634","observation_id":"8ba15dd5-6070-4b62-965b-f38a573eb15e","resolution":{"observed_at":"2026-08-15T18:44:55.895654Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.876695Z","title":"Freeman, and Michael Rubinstein","venue":null,"work_id":"92d1b245-a7a9-4396-ac59-3b7df915e411","year":2025},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.306738Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:faf7a7305229b518fe9a5b91bd6dc906801595b643e936483eb69992e0afd703","observation_id":"7081efb7-ea7e-45a3-b550-3f656ea89aee","resolution":{"observed_at":"2026-08-15T18:44:55.881241Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.861788Z","title":"Look outside the room: Synthesizing a consistent long-term 3d scene video from a single image","venue":null,"work_id":"d33b5a53-c8c0-441f-b7fe-e8f1bb7e427c","year":2022},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.310977Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:e0d1d7ce0e204ad5fb107e05e2dff156390f31a0b5616c7a1105cb9793b2ef7f","observation_id":"56c8bafd-2187-4a89-84d7-f839015c4abe","resolution":{"observed_at":"2026-08-15T18:44:55.866790Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.847120Z","title":"Gen3c: 3d-informed world-consistent video generation with precise camera con- 9 trol","venue":null,"work_id":"499d6799-f61d-4c5a-ba53-4e44b88c1d86","year":2025},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.315241Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:ffbd3bc0438fc3d940ba0987e82b7ae199069aa1aace9f37f4a98d020feaf6fc","observation_id":"85bdd673-92de-462f-8f13-a136417e4ef2","resolution":{"observed_at":"2026-08-15T18:44:55.851839Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.319542Z","title":"Pixel- synth: Generating a 3d-consistent experience from a single image","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.319542Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:3c332309197f6648bc4747e4ece1bf8787f100673e6c02c4f0cb3914b7caaa0a","observation_id":"26acf9de-a973-4e3f-add2-e1f96c4e581d","resolution":{"observed_at":"2026-08-15T18:44:55.319542Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.823270Z","title":"Geometry-free view synthesis: Transformers and no 3d pri- ors","venue":null,"work_id":"1f8bf6f6-4332-4fd4-863a-b29aedfe2d33","year":2021},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.323810Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:1e4b28f7051206655f55d100b41118fe8b7a49704d4747cd289b4738c4913a31","observation_id":"d9994c7b-426a-4ad3-abd8-8ce97e726d8a","resolution":{"observed_at":"2026-08-15T18:44:55.828098Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.17251","last_updated":"2024-09-26T08:22:52Z","snapshot_observed_at":"2026-08-16T13:49:00.677326Z","submitted_at":"2024-05-27T15:07:04Z","title":"GenWarp: Single Image to Novel Views with Semantic-Preserving Generative Warping","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.17251","snapshot_observed_at":"2026-08-15T18:44:55.328050Z","title":"Genwarp: Single image to novel views with semantic-preserving generative warping","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.328050Z"},"links":{"cited_paper":"/paper/2405.17251","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:683f2ca9c02cc409bf10260d95c888ac7e0085f6063274c601e2adf17830f979","observation_id":"7b19f49a-8e83-45c9-a6c7-92098039c04d","resolution":{"observed_at":"2026-08-15T18:44:55.328050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.15110","last_updated":"2023-10-23T17:18:59Z","snapshot_observed_at":"2026-08-15T06:11:26.743203Z","submitted_at":"2023-10-23T17:18:59Z","title":"Zero123++: a Single Image to Consistent Multi-view Diffusion Base Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.15110","snapshot_observed_at":"2026-08-15T18:44:55.332942Z","title":"Zero123++: a single image to consistent multi-view dif- fusion base model","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.332942Z"},"links":{"cited_paper":"/paper/2310.15110","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:3c3ff7760eeaf96d5c69749f7d2641cd993874d86e87c5756c975ae75a54ac51","observation_id":"b2b7f035-5276-48cb-903c-ab9a51de64be","resolution":{"observed_at":"2026-08-15T18:44:55.332942Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.809179Z","title":"Flash3d: Feed-forward gener- alisable 3d scene reconstruction from a single image","venue":null,"work_id":"d2110bfb-7724-45d5-9c1d-249ce111d413","year":null},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.337573Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:8e64bbf25b55da7b00b7cebc94175f1934f28b87e520962d3c0d3fb1c25ceae1","observation_id":"c12bb08e-27c3-4249-9c70-b57a01ec86bd","resolution":{"observed_at":"2026-08-15T18:44:55.813889Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.342114Z","title":"Splatter Image: Ultra-fast single-view 3D recon- struction","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.342114Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:4f607dac0273880c8ee9c180039b710820ad4fd4871d26345fa17db24751f841","observation_id":"37e85b76-892e-4633-8138-ca5008a893f8","resolution":{"observed_at":"2026-08-15T18:44:55.342114Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.346293Z","title":"Sparf: Neural radiance fields from sparse and noisy poses","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.346293Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:20e0c542b57af762188cf668c3e9ecb36572a6c50988f00047c1ddaa3e1c2eb2","observation_id":"97969906-7731-41f5-8542-8db6c3b70c0a","resolution":{"observed_at":"2026-08-15T18:44:55.346293Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.776974Z","title":"Consistent view synthe- sis with pose-guided diffusion models","venue":null,"work_id":"edbbb93e-bb4f-4b36-886d-9911a146c59b","year":2023},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.350562Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:758ddf29aa22f89ba47512c7a821a5be3be4d6a919a21b431c59e03475de98e0","observation_id":"d4e791ea-bea5-4ef8-a0de-2a256589acbc","resolution":{"observed_at":"2026-08-15T18:44:55.781560Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.762777Z","title":"Efros, and Angjoo Kanazawa","venue":null,"work_id":"0d7790e4-f227-42de-b71e-fcae0329d53a","year":2025},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.354695Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:eb1307e94e7c3b340e67e7095e73ff9ecc381eb04cbe4739f03da634559b4dfb","observation_id":"60d1c29b-b81b-4af0-833a-b650fa015388","resolution":{"observed_at":"2026-08-15T18:44:55.767513Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.748835Z","title":"Dust3r: Geometric 3d vi- sion made easy","venue":null,"work_id":"776fec05-9520-4631-8d55-09c2796b687c","year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.358921Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:6c59846e7cf18eeeff5037363661db9687470aaefd4652d70f6512e118d3c0cd","observation_id":"10407b3e-656a-4128-be13-5b7d8ef6f9d5","resolution":{"observed_at":"2026-08-15T18:44:55.753361Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.363135Z","title":"Image quality assessment: from error visibility to structural similarity","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.363135Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:2869db7c6576e70c12d8e10dcdb8863e29f946a3bf121bfee1551cea3aee3486","observation_id":"8ac958c1-8384-498a-ba75-d1c06aab8d01","resolution":{"observed_at":"2026-08-15T18:44:55.363135Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.724761Z","title":"Motionctrl: A unified and flexible motion controller for video generation","venue":null,"work_id":"7db39ef2-c5ed-498a-8fea-7f72051f88da","year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.367690Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:48cbf936f232cbd755a3fb7ed954f0a11190d3aceb6ac627811d2b43435d0e36","observation_id":"a85e073b-92cf-436e-afea-3db49c5284a1","resolution":{"observed_at":"2026-08-15T18:44:55.729382Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.710110Z","title":"SynSin: End-to-end view synthesis from a single image","venue":null,"work_id":"95d08edd-8880-48c1-99bf-9b0e700075db","year":2020},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.372120Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:66338fa3509af92c2fdba6224d5aa259474bf2280c6a5a5190377f637bbee35f","observation_id":"aa654a19-ab04-4d5d-a84f-f0ef55ba838a","resolution":{"observed_at":"2026-08-15T18:44:55.714501Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.696022Z","title":"Reconfusion: 3d reconstruction with diffusion priors","venue":null,"work_id":"4d6a7228-edec-418c-b872-272c047d9e0f","year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.376410Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:3e9ee0a47301c64dd1fa6762910374fda515a9df7f2d2015338253dd2deceaef","observation_id":"72ca79eb-9775-48ee-b752-03bf820cfcff","resolution":{"observed_at":"2026-08-15T18:44:55.700711Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.01506","last_updated":"2025-05-30T11:13:13Z","snapshot_observed_at":"2026-08-17T11:46:18.301015Z","submitted_at":"2024-12-02T13:58:38Z","title":"Structured 3D Latents for Scalable and Versatile 3D Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.01506","snapshot_observed_at":"2026-08-15T18:44:55.380754Z","title":"Structured 3d latents for scalable and versatile 3d gen- eration","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.380754Z"},"links":{"cited_paper":"/paper/2412.01506","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:e2fd18985c30fbdc3d653162190e913954e760987beefb0d658ecceeae85b349","observation_id":"f47c0f9c-ed6a-49cd-98cb-76a2fadaccbe","resolution":{"observed_at":"2026-08-15T18:44:55.380754Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.682128Z","title":"Worldmem: Long- term consistent world simulation with memory, 2025","venue":null,"work_id":"d95e6997-683d-4b04-8292-e9366c686422","year":2025},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.385553Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:06bab266a3bec50488cc05bb61eb492afc09c714213520ec6de4659a972bbf19","observation_id":"ae01d324-f964-412f-9043-cc83edc457b2","resolution":{"observed_at":"2026-08-15T18:44:55.686515Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.03884","last_updated":"2024-04-12T16:47:05Z","snapshot_observed_at":"2026-08-17T18:23:28.199217Z","submitted_at":"2023-12-06T20:22:32Z","title":"WonderJourney: Going from Anywhere to Everywhere","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.03884","snapshot_observed_at":"2026-08-15T18:44:55.390198Z","title":"Freeman, Forrester Cole, Deqing Sun, Noah Snavely, Jiajun Wu, and Charles Her- rmann","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.390198Z"},"links":{"cited_paper":"/paper/2312.03884","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:3cbc1aa6f48e8ae3fac4c713c1dff98196eda8c4ad36a800e16360cce3ab3143","observation_id":"2fdfa670-b2b8-4fab-a882-45a16422e152","resolution":{"observed_at":"2026-08-15T18:44:55.390198Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.09394","last_updated":"2025-03-25T01:00:11Z","snapshot_observed_at":"2026-08-16T13:43:13.507178Z","submitted_at":"2024-06-13T17:59:10Z","title":"WonderWorld: Interactive 3D Scene Generation from a Single Image","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.09394","snapshot_observed_at":"2026-08-15T18:44:55.395021Z","title":"Freeman, and Jiajun Wu","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.395021Z"},"links":{"cited_paper":"/paper/2406.09394","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:b953904ad3142dc5c4f8b4cfd1b896d1d61f135b9f9dbcd75f057366b9aa38f0","observation_id":"e52e608a-0d15-4768-994e-7b4ffa8a1af0","resolution":{"observed_at":"2026-08-15T18:44:55.395021Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.668162Z","title":"Long-term photometric consistent novel view synthesis with diffusion models","venue":null,"work_id":"01ecdfc1-0639-47ed-aee6-8b5e4f2c4521","year":2023},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.400815Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:45c12535455579f8ed7caa5119ef0416b70ea0780015c55acf1425ded32336f3","observation_id":"0b32d7fc-e62f-44a1-a168-d39ca32cb2d2","resolution":{"observed_at":"2026-08-15T18:44:55.672684Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.02048","last_updated":"2024-09-03T16:53:19Z","snapshot_observed_at":"2026-08-14T12:22:44.567374Z","submitted_at":"2024-09-03T16:53:19Z","title":"ViewCrafter: Taming Video Diffusion Models for High-fidelity Novel View Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.02048","snapshot_observed_at":"2026-08-15T18:44:55.405963Z","title":"Viewcrafter: Taming video diffusion models for high-fidelity novel view synthesis.arXiv preprint arXiv:2409.02048, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.405963Z"},"links":{"cited_paper":"/paper/2409.02048","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:6023074bcb49b5cb1224cf979d304b40576b4658a183ca677924b9dbe2cd497c","observation_id":"dcf4d948-3bc3-4372-8a02-2ecc49237467","resolution":{"observed_at":"2026-08-15T18:44:55.405963Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.651388Z","title":"Stargen: A spatiotemporal autoregression framework with video dif- fusion model for scalable and controllable scene generation,","venue":null,"work_id":"f66d04dd-8691-4ddf-9a56-6c5b5f19c81e","year":null},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.410547Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:1f350df04f77ef9519579dfe8c3203f5780250bf06544644f2f478e0e19d1a7e","observation_id":"8882c8b9-204b-46dc-b8de-0a52f82ec8a1","resolution":{"observed_at":"2026-08-15T18:44:55.658212Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:44:55.415264Z","title":"The unreasonable effectiveness of deep features as a perceptual metric","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.415264Z"},"links":{"citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:f26288b9dd14a94acab0ade5f93660fcbd0658fa753380e763019ff30c4bd1ca","observation_id":"28f4a810-5771-4100-9514-b2432af03caf","resolution":{"observed_at":"2026-08-15T18:44:55.415264Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.14489","last_updated":"2025-04-01T18:22:54Z","snapshot_observed_at":"2026-08-16T12:48:53.497558Z","submitted_at":"2025-03-18T17:57:22Z","title":"Stable Virtual Camera: Generative View Synthesis with Diffusion Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.14489","snapshot_observed_at":"2026-08-15T18:44:55.419913Z","title":"Stable virtual camera: Gen- erative view synthesis with diffusion models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.419913Z"},"links":{"cited_paper":"/paper/2503.14489","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:4805f95b2fb5130c0e6040d43012bfc86925a5173993b37479f5cb2b18591e13","observation_id":"b7987984-6ecf-468e-b6fd-77a64e697e88","resolution":{"observed_at":"2026-08-15T18:44:55.419913Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1805.09817","last_updated":"2018-05-24T17:58:02Z","snapshot_observed_at":"2026-08-16T00:52:41.730888Z","submitted_at":"2018-05-24T17:58:02Z","title":"Stereo Magnification: Learning View Synthesis using Multiplane Images","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.09817","snapshot_observed_at":"2026-08-15T18:44:55.425136Z","title":"Stereo magnification: Learning view synthesis using multiplane images","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.425136Z"},"links":{"cited_paper":"/paper/1805.09817","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:ae93ecd7e32c8f9a88bdb40aeb37eda4aab30b31a2b12fa2505db82b32bbf9f6","observation_id":"8eb2ca01-6884-4ff9-bbae-f56586a1fbcf","resolution":{"observed_at":"2026-08-15T18:44:55.425136Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.06714","last_updated":"2023-08-25T14:36:57Z","snapshot_observed_at":"2026-08-16T15:40:37.360341Z","submitted_at":"2023-04-13T17:59:01Z","title":"Single-Stage Diffusion NeRF: A Unified Approach to 3D Generation and Reconstruction","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.06714","snapshot_observed_at":"2026-08-15T18:44:55.229391Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory","version":3},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-15T18:44:55.229391Z"},"links":{"cited_paper":"/paper/2304.06714","citing_paper":"/paper/2506.18903"},"observation_digest":"sha256:2f51394da056512826eb4f71f8a281abe5b9e43005bc217334d054b02a8c722a","observation_id":"c584eebf-9972-40e9-bf18-66558bbd58a7","resolution":{"observed_at":"2026-08-15T18:44:55.229391Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.18903","last_updated":"2025-08-14T14:03:30Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-15T18:38:37.788719Z","submitted_at":"2025-06-23T17:59:56Z","title":"VMem: Consistent Interactive Video Scene Generation with Surfel-Indexed View Memory"},"reference_resolution":{"displayed":48,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":22,"verified_exact":0,"verified_fuzzy":26},"total_outbound_references":48},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 17 August 2026, this Paper Citation Record lists 48 of 48 outbound references and 15 inbound Pith citation observations for arXiv:2506.18903."}