{"as_of":"2026-08-23T20:50:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2d74e4a6de573cb482232988f3c5e1210efc73932f6e10dcda5e46a0f6175764","coverage":[{"denominator":51,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":51,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T05:45:57.633484Z","state":"measured"},{"denominator":54,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":54,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-23T06:30:58.430688+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:01:10.839117Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T17:07:12.437951Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.19834","snapshot_observed_at":"2026-08-07T14:01:10.839117Z","title":"Animateanywhere: Rouse the background in human image animation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.20255","last_updated":"2025-07-07T17:56:30Z","snapshot_observed_at":"2026-08-14T10:03:15.306691Z","submitted_at":"2025-05-26T17:32:10Z","title":"AniCrafter: Customizing Realistic Human-Centric Animation via Avatar-Background Conditioning in Video Diffusion Models","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T14:01:10.839117Z"},"links":{"cited_paper":"/paper/2504.19834","citing_paper":"/paper/2505.20255"},"observation_digest":"sha256:2ba4f481e580e7494078ba73868ce1b7a5ecd74bc385a5c5572e07a53ad6554a","observation_id":"018c23f9-6a99-45c9-9f2f-bf3a4323cb1d","resolution":{"observed_at":"2026-08-07T14:01:10.839117Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"cited_work":{"arxiv_id":"2504.19834","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.19834","snapshot_observed_at":"2026-07-02T17:07:12.437951Z","title":"Ani- mateanywhere: Rouse the background in human image ani- mation.arXiv preprint arXiv:2504.19834, 2025","venue":null,"work_id":"e13c16b3-4efd-415e-8187-be5bdd8f83ab","year":2025},"citing_paper":{"arxiv_id":"2605.15042","last_updated":"2026-05-14T16:36:34Z","snapshot_observed_at":"2026-08-15T12:46:54.250833Z","submitted_at":"2026-05-14T16:36:34Z","title":"EverAnimate: Minute-Scale Human Animation via Latent Flow Restoration","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-06-30T21:21:58.123630Z"},"links":{"cited_paper":"/paper/2504.19834","citing_paper":"/paper/2605.15042"},"observation_digest":"sha256:e00411eae4b36e0807cdc8c2ee23cb727bdaa0d7a0c227085637e02c359e1fb4","observation_id":"ff5c8f7a-78d7-436d-b227-3f57faf0b46c","resolution":{"observed_at":"2026-06-30T21:25:04.889157Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"cited_work":{"arxiv_id":"2504.19834","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.19834","snapshot_observed_at":"2026-07-02T17:07:12.437951Z","title":"Ani- mateanywhere: Rouse the background in human image ani- mation.arXiv preprint arXiv:2504.19834, 2025","venue":null,"work_id":"e13c16b3-4efd-415e-8187-be5bdd8f83ab","year":2025},"citing_paper":{"arxiv_id":"2606.07508","last_updated":"2026-06-05T17:57:32Z","snapshot_observed_at":"2026-08-03T10:42:24.485829Z","submitted_at":"2026-06-05T17:57:32Z","title":"Streaming Video Generation with Streaming Force Control","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-06-27T22:14:32.465663Z"},"links":{"cited_paper":"/paper/2504.19834","citing_paper":"/paper/2606.07508"},"observation_digest":"sha256:2025c8583c850da5f64889a6b7a4f7706797f568230cc1786507d541fba434a6","observation_id":"e60b58b7-cb6d-491e-adb8-e65bfbb3adfc","resolution":{"observed_at":"2026-07-02T17:07:12.439608Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2504.19834/citation-record","integrity":"/paper/2504.19834/integrity","json":"/paper/2504.19834/citation-record.json","paper":"/paper/2504.19834"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.331104Z","title":"Animate anyone: Consistent and controllable image-to-video synthesis for character animation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.331104Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:892a930e39a09bf5f4d2de539432f833c989487f32be07a0e8d75c635834f69b","observation_id":"4474c04f-196f-4a8c-ab6b-dbf62ac787b4","resolution":{"observed_at":"2026-08-16T05:45:57.331104Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.374488Z","title":"Magicanimate: Temporally consistent human image animation using diffusion model,","venue":null,"work_id":"0ca9f6a8-009d-4de0-af78-e2c3b25e8dde","year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.367251Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:667efa8ef3c56e889b6e43394d3fdcca3299a745293bbdebca27db03127953aa","observation_id":"7471a6f4-dd2b-49f6-a5fa-4bcd65e37d02","resolution":{"observed_at":"2026-08-16T05:45:58.379567Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.14781","last_updated":"2024-06-01T08:27:23Z","snapshot_observed_at":"2026-08-18T08:53:24.734837Z","submitted_at":"2024-03-21T18:52:58Z","title":"Champ: Controllable and Consistent Human Image Animation with 3D Parametric Guidance","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.14781","snapshot_observed_at":"2026-08-16T05:45:57.385540Z","title":"Champ: Controllable and consistent human image animation with 3d parametric guidance,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.385540Z"},"links":{"cited_paper":"/paper/2403.14781","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:94e074a5d83ac1cfaf577a644b90026c71ec3d35b6c60def57a31661153477f5","observation_id":"1a078438-16de-4579-b041-8bd9c09691d5","resolution":{"observed_at":"2026-08-16T05:45:57.385540Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17438","last_updated":"2024-11-21T03:26:54Z","snapshot_observed_at":"2026-08-20T20:23:55.806173Z","submitted_at":"2024-07-24T17:15:58Z","title":"HumanVid: Demystifying Training Data for Camera-controllable Human Image Animation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17438","snapshot_observed_at":"2026-08-16T05:45:57.414009Z","title":"Humanvid: Demystifying training data for camera- controllable human image animation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.414009Z"},"links":{"cited_paper":"/paper/2407.17438","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:c4feafb5713b410cf28c94dcf5191d8c47adc33a7bba79a95e0922751e777421","observation_id":"8928703b-6734-4f66-a769-3335fbbaf772","resolution":{"observed_at":"2026-08-16T05:45:57.414009Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16393","last_updated":"2024-05-28T05:25:00Z","snapshot_observed_at":"2026-08-19T00:01:58.136953Z","submitted_at":"2024-05-26T00:53:26Z","title":"Disentangling Foreground and Background Motion for Enhanced Realism in Human Video Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16393","snapshot_observed_at":"2026-08-16T05:45:57.419863Z","title":"Disentangling foreground and background motion for enhanced realism in human video genera- tion,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.419863Z"},"links":{"cited_paper":"/paper/2405.16393","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:c3d8615a6001425e27bc61770f1723228137680c1252a6c5ce2f570292547d4e","observation_id":"136d68e3-1574-40e9-af17-bb3f43c42c66","resolution":{"observed_at":"2026-08-16T05:45:57.419863Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.06202","last_updated":"2024-09-10T04:14:11Z","snapshot_observed_at":"2026-08-16T13:19:54.079878Z","submitted_at":"2024-09-10T04:14:11Z","title":"RealisDance: Equip controllable character animation with realistic hands","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.06202","snapshot_observed_at":"2026-08-16T05:45:57.425265Z","title":"Realisdance: Equip controllable character animation with realistic hands,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.425265Z"},"links":{"cited_paper":"/paper/2409.06202","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:77b371f853184456d4dfb012b0ad284d9e320046f6b7e64fea4b7828ca1cb36c","observation_id":"ca53fc43-4166-44b9-85b7-54355eea232d","resolution":{"observed_at":"2026-08-16T05:45:57.425265Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.03644","last_updated":"2024-11-13T01:45:31Z","snapshot_observed_at":"2026-08-16T13:53:16.096347Z","submitted_at":"2024-09-05T16:02:11Z","title":"RealisHuman: A Two-Stage Approach for Refining Malformed Human Parts in Generated Images","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.03644","snapshot_observed_at":"2026-08-16T05:45:57.430778Z","title":"Realishuman: A two-stage approach for refining malformed human parts in generated images,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.430778Z"},"links":{"cited_paper":"/paper/2409.03644","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:34dca99af364305fac2090d1dba3c996e185e75d7acde63c93de2a795e4d0d0e","observation_id":"345c54a4-ab4d-4efb-a8ff-e1cdd7cfb2d8","resolution":{"observed_at":"2026-08-16T05:45:57.430778Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01188","last_updated":"2024-06-03T10:51:10Z","snapshot_observed_at":"2026-08-18T22:36:33.661827Z","submitted_at":"2024-06-03T10:51:10Z","title":"UniAnimate: Taming Unified Video Diffusion Models for Consistent Human Image Animation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01188","snapshot_observed_at":"2026-08-16T05:45:57.435663Z","title":"Unianimate: Taming unified video diffusion models for consistent human image animation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.435663Z"},"links":{"cited_paper":"/paper/2406.01188","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:e37b63c47e561a49ebf582e67457bbddd091fdd49608318616fd80684d6239f9","observation_id":"2d8eda43-db87-429e-a90d-157d61edfc8d","resolution":{"observed_at":"2026-08-16T05:45:57.435663Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.06072","last_updated":"2025-03-26T08:33:10Z","snapshot_observed_at":"2026-08-15T19:52:28.153954Z","submitted_at":"2024-08-12T11:47:11Z","title":"CogVideoX: Text-to-Video Diffusion Models with An Expert Transformer","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.06072","snapshot_observed_at":"2026-08-16T05:45:57.440577Z","title":"Cogvideox: Text-to-video diffusion models with an expert transformer,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.440577Z"},"links":{"cited_paper":"/paper/2408.06072","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:52771bf6219672f5dff9a0822177b3c51504006d92064f65e00fa3880992fa43","observation_id":"b611f5b4-ad6f-4ed5-8cfb-816beb9220fe","resolution":{"observed_at":"2026-08-16T05:45:57.440577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.446191Z","title":"Lora: Low-rank adaptation of large language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.446191Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:c54d92373d0fd52442f46fb442f2e768811b7ff17b966f911e1dfee7c0bbede0","observation_id":"80a9339d-5f48-40c3-a231-1b33e69cf46a","resolution":{"observed_at":"2026-08-16T05:45:57.446191Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.350495Z","title":"First or- der motion model for image animation,","venue":null,"work_id":"5a27b60d-7f0f-4253-b896-982463a246fb","year":2019},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.451042Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:cf51f5a63b7d882b53887d7112967f7fa5dec209d265523747819485ebd71e41","observation_id":"1f9afd59-7626-4d89-90ea-f0e10df593ad","resolution":{"observed_at":"2026-08-16T05:45:58.355153Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.335596Z","title":"Gac-gan: A general method for appearance-controllable human video motion transfer,","venue":null,"work_id":"0aa1d77e-0dfc-48d9-abcd-1ce8a5af7326","year":2020},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.455527Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:198ec5d91b6eb3777dce2dba350adc1e71d936ef78364fcf5da8a5003c1b1958","observation_id":"6865e358-fb24-4ca1-b399-5989d10f4e9a","resolution":{"observed_at":"2026-08-16T05:45:58.340541Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.460001Z","title":"One-shot free-view neural talking-head synthesis for video conferencing,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.460001Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:c3f70d6dc75b309c24247ea27f6a78d4bb4b36cfd70fd7fd9c728e9c6657a0be","observation_id":"9227b25e-8e4d-492d-8019-3164e817eb61","resolution":{"observed_at":"2026-08-16T05:45:57.460001Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.311665Z","title":"Humangan: A generative model of human images,","venue":null,"work_id":"2ea9e8b9-00b9-4eaa-b1ac-7e182a5c27d5","year":2021},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.464491Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:41448536c7ee7309b959e4737cf09682ebd309e38e7dd9095f9d807028e1d461","observation_id":"9f153d76-a392-4f34-b59c-0ba4fb71e0f7","resolution":{"observed_at":"2026-08-16T05:45:58.316144Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2104.15069","last_updated":"2021-04-30T15:38:41Z","snapshot_observed_at":"2026-08-16T18:27:50.672545Z","submitted_at":"2021-04-30T15:38:41Z","title":"A Good Image Generator Is What You Need for High-Resolution Video Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.15069","snapshot_observed_at":"2026-08-16T05:45:57.468904Z","title":"A good image generator is what you need for high- resolution video synthesis,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.468904Z"},"links":{"cited_paper":"/paper/2104.15069","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:2fe0a62eb5182cefa54edf4502ff6261bdb8a3c4c750cd793464ff8520cfcac4","observation_id":"a4d7bd59-f9a2-4b39-a98e-3ce7fa40bdb3","resolution":{"observed_at":"2026-08-16T05:45:57.468904Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.296580Z","title":"Coherent image animation using spatial-temporal correspondence,","venue":null,"work_id":"526b54a8-608d-42db-9363-3e59d0b962ee","year":2022},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.474067Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:df954ead15affe11efa06849ed90cee923540eb8b0bcfbba80264ab550033d08","observation_id":"cd549cab-50fc-4217-9da1-8587b252a2e2","resolution":{"observed_at":"2026-08-16T05:45:58.301546Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.281915Z","title":"One-shot human motion transfer via occlusion-robust flow prediction and neural texturing,","venue":null,"work_id":"1ecf34df-08f9-4a6b-8e97-aa46c8b6dc96","year":2025},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.478503Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:d1677258aa7068d6c30d29bbcfa8580e367c4fd30fb99f06168045e0da11eeb2","observation_id":"bcbe9452-3b80-4fe6-a959-03c9baab9c10","resolution":{"observed_at":"2026-08-16T05:45:58.286678Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02004","last_updated":"2025-04-01T18:12:22Z","snapshot_observed_at":"2026-08-16T12:44:50.791789Z","submitted_at":"2025-04-01T18:12:22Z","title":"Beyond Static Scenes: Camera-controllable Background Generation for Human Motion","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.02004","snapshot_observed_at":"2026-08-16T05:45:57.482838Z","title":"Beyond static scenes: Camera-controllable back- ground generation for human motion,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.482838Z"},"links":{"cited_paper":"/paper/2504.02004","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:8e9a97637bab230aff7fa8a7dbaa3b3b52cfdaa93d7909f2f058f00b4127b9f1","observation_id":"0368c88d-3b39-4252-a411-9020e59e0a2b","resolution":{"observed_at":"2026-08-16T05:45:57.482838Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.487439Z","title":"Musepose: a pose-driven image-to-video framework for virtual human generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.487439Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:78b71d562effd157b2c299a643369a6c9cd09ab6fa59ae0baf77ec2b85023f97","observation_id":"23f2ad14-6bd0-4863-a588-9a8b2a3be997","resolution":{"observed_at":"2026-08-16T05:45:57.487439Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.10306","last_updated":"2024-12-11T02:55:31Z","snapshot_observed_at":"2026-08-17T02:05:24.887374Z","submitted_at":"2024-10-14T09:06:55Z","title":"Animate-X: Universal Character Image Animation with Enhanced Motion Representation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.10306","snapshot_observed_at":"2026-08-16T05:45:57.491426Z","title":"Animate-x: Universal character image animation with enhanced motion representation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.491426Z"},"links":{"cited_paper":"/paper/2410.10306","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:353a8a81b992e5c0876b6316482b361deae979780f0c6b59f55d8f0a9db9b2b6","observation_id":"09366ab6-c20d-41b8-908a-05852c618758","resolution":{"observed_at":"2026-08-16T05:45:57.491426Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.16160","last_updated":"2025-06-11T08:28:01Z","snapshot_observed_at":"2026-08-16T13:58:12.955222Z","submitted_at":"2024-09-24T15:00:07Z","title":"MIMO: Controllable Character Video Synthesis with Spatial Decomposed Modeling","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.16160","snapshot_observed_at":"2026-08-16T05:45:57.495997Z","title":"Mimo: Controllable character video synthesis with spatial decomposed modeling,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.495997Z"},"links":{"cited_paper":"/paper/2409.16160","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:58ef0a8c573d4b1f128c61831fdf58be5a0f2eb9c91afbc76e64d129a7de38fc","observation_id":"250c77e4-b1f2-4288-948f-abd4bb38fd43","resolution":{"observed_at":"2026-08-16T05:45:57.495997Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06145","last_updated":"2025-02-10T04:20:11Z","snapshot_observed_at":"2026-08-14T23:19:43.982858Z","submitted_at":"2025-02-10T04:20:11Z","title":"Animate Anyone 2: High-Fidelity Character Image Animation with Environment Affordance","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06145","snapshot_observed_at":"2026-08-16T05:45:57.500801Z","title":"Animate anyone 2: High-fidelity charac- ter image animation with environment affordance,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.500801Z"},"links":{"cited_paper":"/paper/2502.06145","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:5d48ed66db91dfb0e4ddcf4305b5549f55a8b4fff01aaa83ec0e7011a3d867d1","observation_id":"c48a0ead-fc83-4bce-98e3-1df6ed52966d","resolution":{"observed_at":"2026-08-16T05:45:57.500801Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.17346","last_updated":"2024-09-02T17:30:45Z","snapshot_observed_at":"2026-08-16T14:06:26.055809Z","submitted_at":"2024-03-26T03:10:45Z","title":"TRAM: Global Trajectory and Motion of 3D Humans from in-the-wild Videos","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.17346","snapshot_observed_at":"2026-08-16T05:45:57.506525Z","title":"Tram: Global trajectory and motion of 3d humans from in-the-wild videos,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.506525Z"},"links":{"cited_paper":"/paper/2403.17346","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:fce57cccfd2e027cf3f45d1c5ba555530420d75fd613ac458f471743155e4e1f","observation_id":"a5f5f4a2-a868-4c4d-9e9a-9e09a0f41d3e","resolution":{"observed_at":"2026-08-16T05:45:57.506525Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.511098Z","title":"Droid-slam: Deep visual slam for monocular, stereo, and rgb-d cameras,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.511098Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:8a2daf968e720f0d95be1fb49f55a66aa778466495b8ecc9402975a06782ea22","observation_id":"16ef635b-d8b1-4140-889f-6dd278e7ed38","resolution":{"observed_at":"2026-08-16T05:45:57.511098Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.02101","last_updated":"2025-03-13T18:35:06Z","snapshot_observed_at":"2026-08-16T17:39:09.604516Z","submitted_at":"2024-04-02T16:52:41Z","title":"CameraCtrl: Enabling Camera Control for Text-to-Video Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.02101","snapshot_observed_at":"2026-08-16T05:45:57.516122Z","title":"Cameractrl: Enabling camera control for text-to-video generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.516122Z"},"links":{"cited_paper":"/paper/2404.02101","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:964431f151e6d4fb8a8d219d02aea535c96111f6ffb5674d750bbc43c8aef072","observation_id":"f153f53e-3ccd-45fe-954d-690795b16a96","resolution":{"observed_at":"2026-08-16T05:45:57.516122Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.520882Z","title":"Video diffusion models,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.520882Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:4d0f8981d935fd526e964953f0b93960f4a322fb4ca9925656463cb6716d7216","observation_id":"4ad1b228-fdfe-4689-a97e-1a9dffdac01a","resolution":{"observed_at":"2026-08-16T05:45:57.520882Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.04725","last_updated":"2024-02-08T18:08:57Z","snapshot_observed_at":"2026-08-17T14:04:31.230742Z","submitted_at":"2023-07-10T17:34:16Z","title":"AnimateDiff: Animate Your Personalized Text-to-Image Diffusion Models without Specific Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.04725","snapshot_observed_at":"2026-08-16T05:45:57.525580Z","title":"Animatediff: Animate your personalized text- to-image diffusion models without specific tuning,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.525580Z"},"links":{"cited_paper":"/paper/2307.04725","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:94f58d16d3727c93cbb516b2bc546cf7bbcc3619dd3fdd03a51b131b948821a5","observation_id":"7d569291-8eb4-4373-81c4-59defbaff5a6","resolution":{"observed_at":"2026-08-16T05:45:57.525580Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.15127","last_updated":"2023-11-25T22:28:38Z","snapshot_observed_at":"2026-08-20T09:39:07.545813Z","submitted_at":"2023-11-25T22:28:38Z","title":"Stable Video Diffusion: Scaling Latent Video Diffusion Models to Large Datasets","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.15127","snapshot_observed_at":"2026-08-16T05:45:57.530256Z","title":"Stable video diffusion: Scaling latent video diffusion models to large datasets,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.530256Z"},"links":{"cited_paper":"/paper/2311.15127","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:8b62eab4a04b8b016f530d8996c10e67df7bf301ec800e073e001ee8a8e81d74","observation_id":"de2b8774-aa09-4698-a77a-680cce7077ca","resolution":{"observed_at":"2026-08-16T05:45:57.530256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.12945","last_updated":"2024-02-05T16:36:30Z","snapshot_observed_at":"2026-08-20T04:48:52.471048Z","submitted_at":"2024-01-23T18:05:25Z","title":"Lumiere: A Space-Time Diffusion Model for Video Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.12945","snapshot_observed_at":"2026-08-16T05:45:57.535194Z","title":"Lumiere: A space-time diffusion model for video generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.535194Z"},"links":{"cited_paper":"/paper/2401.12945","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:9e294977002295e37c46a6dcacd41d6ff2f5b0738aee6b64481e8f014bdf6955","observation_id":"ef6fc260-4a12-4b4e-aa0b-39d3f06e3983","resolution":{"observed_at":"2026-08-16T05:45:57.535194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.240565Z","title":"Photorealistic video generation with diffusion models,","venue":null,"work_id":"f7f632d6-0088-483d-9f36-6f43b59187d9","year":2025},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.539652Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:a56327b7e9aea01622d3d45c91eabb7a0fd4577b9fa30ca664a81ba88a2d8455","observation_id":"7822d379-218b-42a0-b7f7-ac584d55c76d","resolution":{"observed_at":"2026-08-16T05:45:58.245208Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13077","last_updated":"2023-05-22T14:48:53Z","snapshot_observed_at":"2026-08-18T20:48:53.851874Z","submitted_at":"2023-05-22T14:48:53Z","title":"ControlVideo: Training-free Controllable Text-to-Video Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.13077","snapshot_observed_at":"2026-08-16T05:45:57.543938Z","title":"Con- trolvideo: Training-free controllable text-to-video generation,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.543938Z"},"links":{"cited_paper":"/paper/2305.13077","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:db32ef12ca7dc3f0944ebc7715d2bbd26ae58be8eae3919d85673c19d01d123f","observation_id":"766cea25-4102-4fed-aff1-25f2f1463dea","resolution":{"observed_at":"2026-08-16T05:45:57.543938Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.223315Z","title":"Sparsectrl: Adding sparse controls to text-to-video diffusion models,","venue":null,"work_id":"55bece20-2c19-4533-b27b-a30076670a0b","year":2025},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.548432Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:87df41918d76fe81aca1f89d7375adca88705e58d1036b218cb9829f6a60a929","observation_id":"0deb2852-f0cc-48e0-b000-8e59ed4a6feb","resolution":{"observed_at":"2026-08-16T05:45:58.229356Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.207977Z","title":"Dreamvideo: Composing your dream videos with customized subject and motion,","venue":null,"work_id":"08e12474-6ef9-478b-b704-038922ab509e","year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.552689Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:79c52de99d2a3fb37d0da2c15f21f5de7d5280557fed408e2e6c6942f00625b7","observation_id":"7a90a08e-7467-443d-92e8-b27df4b0f07d","resolution":{"observed_at":"2026-08-16T05:45:58.213072Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.556904Z","title":"Videodreamer: Customized multi-subject text-to-video generation with disen-mix finetuning on language-video foundation models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.556904Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:fa329267ec36140211519343a2816e6e1b5ea95c2f47db7851cb0c9f260068f8","observation_id":"bc554827-c28e-41d7-bf0f-a0695d5cf359","resolution":{"observed_at":"2026-08-16T05:45:57.556904Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.09685","last_updated":"2021-10-16T18:40:34Z","snapshot_observed_at":"2026-08-20T11:47:17.477107Z","submitted_at":"2021-06-17T17:37:18Z","title":"LoRA: Low-Rank Adaptation of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.09685","snapshot_observed_at":"2026-08-16T05:45:57.561504Z","title":"Lora: Low-rank adaptation of large language models,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.561504Z"},"links":{"cited_paper":"/paper/2106.09685","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:91c29490f5c6d84e3bd6dfea409010c08a4b5a0c44a9088fb1874896038f95f5","observation_id":"0871f705-2e2e-49a8-9ea0-7a25ed0335a2","resolution":{"observed_at":"2026-08-16T05:45:57.561504Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.183857Z","title":"Direct-a-video: Customized video generation with user- directed camera movement and object motion,","venue":null,"work_id":"fb9bfaf2-40c1-4092-a781-bd83e668ceca","year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.566026Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:11ce76d724f517e5db63ede11c9e70c64f4a4c4f1775692f85edcfbc7f904237","observation_id":"a64811e3-6959-4a98-ab56-f2362254e5d2","resolution":{"observed_at":"2026-08-16T05:45:58.188429Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.168924Z","title":"Motionctrl: A unified and flexible motion controller for video generation,","venue":null,"work_id":"7b04638d-7adc-457b-abf8-6ab2be2c71ab","year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.570300Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:0d7f634b2032a7ad9e16223d5b903b2d580de846a08b925660796c6e19817a0f","observation_id":"f914a20b-d7ef-459e-b30b-15597d3f8baa","resolution":{"observed_at":"2026-08-16T05:45:58.173580Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.02509","last_updated":"2024-06-04T17:27:19Z","snapshot_observed_at":"2026-08-21T19:57:02.214817Z","submitted_at":"2024-06-04T17:27:19Z","title":"CamCo: Camera-Controllable 3D-Consistent Image-to-Video Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.02509","snapshot_observed_at":"2026-08-16T05:45:57.574337Z","title":"Camco: Camera-controllable 3d-consistent image-to-video generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.574337Z"},"links":{"cited_paper":"/paper/2406.02509","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:fba678a68b6121d7d1878a0bd6bd7ccee83731e967ea83d5915584f69529de6b","observation_id":"225298ff-7c22-4a97-90b1-38d9ab7a019e","resolution":{"observed_at":"2026-08-16T05:45:57.574337Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.15957","last_updated":"2024-12-04T12:54:44Z","snapshot_observed_at":"2026-08-20T05:26:36.217850Z","submitted_at":"2024-10-21T12:36:27Z","title":"CamI2V: Camera-Controlled Image-to-Video Diffusion Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.15957","snapshot_observed_at":"2026-08-16T05:45:57.579261Z","title":"Cami2v: Camera-controlled image-to-video diffusion model,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.579261Z"},"links":{"cited_paper":"/paper/2410.15957","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:8475b4830f87dacdc53d87d6e01c500a5a5d46b4f38ee4e105da36109dc3489b","observation_id":"1bd559aa-63d9-4feb-a61f-da5bcd352df9","resolution":{"observed_at":"2026-08-16T05:45:57.579261Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10126","last_updated":"2025-02-25T00:32:29Z","snapshot_observed_at":"2026-08-22T03:09:26.395691Z","submitted_at":"2024-06-14T15:33:00Z","title":"Training-free Camera Control for Video Generation","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10126","snapshot_observed_at":"2026-08-16T05:45:57.583808Z","title":"Training-free camera control for video generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.583808Z"},"links":{"cited_paper":"/paper/2406.10126","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:c726eb22d2350fa0a1912f9343d3100ce2f72717d414d988da0c9c9660280743","observation_id":"26d8bb1e-432e-45a3-aa8d-f66f14067417","resolution":{"observed_at":"2026-08-16T05:45:57.583808Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.588617Z","title":"Adding conditional control to text-to-image diffusion models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.588617Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:7a7f9fd2aee1097d8c56d5e512e70d9caf6ee6fe926f7e1a5c8ea229806ce2da","observation_id":"06f43448-8587-409c-911e-d9321d5a87d9","resolution":{"observed_at":"2026-08-16T05:45:57.588617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.593171Z","title":"T2i- adapter: Learning adapters to dig out more controllable ability for text- to-image diffusion models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.593171Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:7f15ffa330058b2f9ab785496d5318a405abaf39f309807181128ff773080122","observation_id":"5e3eefbd-fc02-4c7f-99ed-461a16bee889","resolution":{"observed_at":"2026-08-16T05:45:57.593171Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.597868Z","title":"Scope of validity of psnr in im- age/video quality assessment,","venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.597868Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:4f2dc2d10f63ca5ff0cca5d8cf2b8e8fb671c481cb04a1cc227078e4ba37756d","observation_id":"9332772a-d7a7-4cf7-8bd4-3bfa6bdd9ec4","resolution":{"observed_at":"2026-08-16T05:45:57.597868Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.602205Z","title":"Image quality assessment: from error visibility to structural similarity,","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.602205Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:a08c22816ff0e115680cddfcd93f7b1a15805793e1070154e392cdd2da81191e","observation_id":"4e7d91e8-a343-4620-8a0c-a7a1514a35ff","resolution":{"observed_at":"2026-08-16T05:45:57.602205Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.606548Z","title":"The unreasonable effectiveness of deep features as a perceptual metric,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.606548Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:b25280bb489d840c93cb42aed1bc1a8b2e6b5959ef8ffb2dedf65011133116eb","observation_id":"870ea23b-addf-420e-a2a6-960a633454bf","resolution":{"observed_at":"2026-08-16T05:45:57.606548Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.611086Z","title":"Gans trained by a two time-scale update rule converge to a local nash equilibrium,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.611086Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:f881d937034dfdb284abbdfa8756a87456253c79cb91a076e378b2731f91295e","observation_id":"7705450e-4887-444d-8f7c-3678be21b181","resolution":{"observed_at":"2026-08-16T05:45:57.611086Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.100245Z","title":"Fvd: A new metric for video generation,","venue":null,"work_id":"3603a5e1-2211-4c00-b3d6-b3aa20b6fbfd","year":2019},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.615461Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:18251af282f7f5dba5872118b33d7a3b12808d037a9f441da144a594ca2ee301","observation_id":"f8a4a7f9-4f95-4548-8bd3-f6779cc87641","resolution":{"observed_at":"2026-08-16T05:45:58.104770Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:57.619973Z","title":"Grounded sam: Assembling open-world models for diverse visual tasks,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.619973Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:20cad05f92e695c72983a62a716691fc612fbeb617fbe249b7080218dd8c5455","observation_id":"98199709-b4a9-4f7b-875e-04b06b02e58c","resolution":{"observed_at":"2026-08-16T05:45:57.619973Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T05:45:58.073694Z","title":"ProPainter: Improving propagation and transformer for video inpainting,","venue":null,"work_id":"305b45b9-f6d3-4e1b-b835-9077206f0c50","year":2023},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.624477Z"},"links":{"citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:0efb7fb1aa4fc3404766d456d36fca6a581e4f605fe555b261e8197774b158ef","observation_id":"7f9e8115-8186-40c3-9829-63ff3ec3b8d7","resolution":{"observed_at":"2026-08-16T05:45:58.080208Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1711.05101","last_updated":"2019-01-04T21:01:49Z","snapshot_observed_at":"2026-08-14T20:13:52.872565Z","submitted_at":"2017-11-14T14:24:06Z","title":"Decoupled Weight Decay Regularization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.05101","snapshot_observed_at":"2026-08-16T05:45:57.628847Z","title":"Decoupled weight decay regularization,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.628847Z"},"links":{"cited_paper":"/paper/1711.05101","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:4d9de8cceb493b83ab8a3e5ed2ac69102f1a5ad8fc346975198c22dd845aa2a8","observation_id":"8f28af0b-3a0a-433e-bdd9-5fd40c99c510","resolution":{"observed_at":"2026-08-16T05:45:57.628847Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.21246","last_updated":"2025-05-18T16:24:19Z","snapshot_observed_at":"2026-08-19T06:15:02.338902Z","submitted_at":"2025-03-27T08:07:45Z","title":"DynamiCtrl: Rethinking the Basic Structure and the Role of Text for High-quality Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.21246","snapshot_observed_at":"2026-08-16T05:45:57.633484Z","title":"Dynamictrl: Rethinking the basic structure and the role of text for high-quality human image animation,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-16T05:45:57.633484Z"},"links":{"cited_paper":"/paper/2503.21246","citing_paper":"/paper/2504.19834"},"observation_digest":"sha256:fe391b015636df9fbc1396bc1fc9aa9c09e3e25313c1b475751f0cf535a4479e","observation_id":"08c8a930-2454-499f-a950-a5be3b6d519c","resolution":{"observed_at":"2026-08-16T05:45:57.633484Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2504.19834","last_updated":"2025-08-23T09:05:05Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-18T18:28:23.774385Z","submitted_at":"2025-04-28T14:35:01Z","title":"AnimateAnywhere: Rouse the Background in Human Image Animation"},"reference_resolution":{"displayed":51,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":38,"verified_exact":0,"verified_fuzzy":13},"total_outbound_references":51},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"thesis":"As of 23 August 2026, this Paper Citation Record lists 51 of 51 outbound references and 3 inbound Pith citation observations for arXiv:2504.19834."}