{"as_of":"2026-08-22T01:41:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:38ca814f30efa2fe01702f2c0ad14ac6b7085ad1731b7426d79c6e58b0941182","coverage":[{"denominator":67,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":67,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T11:56:28.346922Z","state":"measured"},{"denominator":84,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":84,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":17,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":17,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T23:48:14.277107Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T16:27:09.335288Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-10T16:12:00.016561Z","title":"Stableanimator: High-quality identity-preserving human image animation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.13452","last_updated":"2025-02-27T06:55:41Z","snapshot_observed_at":"2026-08-18T21:45:51.574465Z","submitted_at":"2025-01-23T08:06:11Z","title":"EchoVideo: Identity-Preserving Human Video Generation by Multimodal Feature Fusion","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T16:12:00.016561Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2501.13452"},"observation_digest":"sha256:e51e364b717997f4a243f5b3e7a32eea260bb9bb1e9901d7bcd09c0249e52605","observation_id":"a2832985-90fc-4bdf-bddd-aaba45524c76","resolution":{"observed_at":"2026-08-10T16:12:00.016561Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-15T23:48:14.277107Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.03730","last_updated":"2025-05-06T17:58:02Z","snapshot_observed_at":"2026-08-19T14:05:25.246376Z","submitted_at":"2025-05-06T17:58:02Z","title":"FlexiAct: Towards Flexible Action Control in Heterogeneous Scenarios","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-15T23:48:14.277107Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2505.03730"},"observation_digest":"sha256:3eb21dd9e1b4804f6882527e4ad89d27f9f4625f680877d3676fb652c72cdbce","observation_id":"2a7a58fe-83ae-4c93-bc97-bf129bd7feba","resolution":{"observed_at":"2026-08-15T23:48:14.277107Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-07T14:39:37.877008Z","title":"Stablean- imator: High-quality identity-preserving human image animation.arXiv preprint arXiv:2411.17697, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.18078","last_updated":"2025-05-23T16:37:14Z","snapshot_observed_at":"2026-08-13T00:43:53.956087Z","submitted_at":"2025-05-23T16:37:14Z","title":"DanceTogether! Identity-Preserving Multi-Person Interactive Video Generation","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-07T14:39:37.877008Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2505.18078"},"observation_digest":"sha256:fd7addfed43651426b9d2720a9177f97cfd20c254a5f67663d9745d15b9450ce","observation_id":"24dcd086-b1d8-49c5-b27f-026e69609a5c","resolution":{"observed_at":"2026-08-07T14:39:37.877008Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-07T14:01:13.094891Z","title":"Stableanimator: High- quality identity-preserving human image animation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20255","last_updated":"2025-07-07T17:56:30Z","snapshot_observed_at":"2026-08-14T10:03:15.306691Z","submitted_at":"2025-05-26T17:32:10Z","title":"AniCrafter: Customizing Realistic Human-Centric Animation via Avatar-Background Conditioning in Video Diffusion Models","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T14:01:13.094891Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2505.20255"},"observation_digest":"sha256:51717af53269d76a1a508b28fcd7297b4eb88fc3be1d94330a8ad6e90616f8e3","observation_id":"b91aba48-1ac2-404e-aaaa-151e43532135","resolution":{"observed_at":"2026-08-07T14:01:13.094891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-07T05:07:15.128129Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08797","last_updated":"2026-07-24T09:02:20Z","snapshot_observed_at":"2026-08-20T23:43:22.344539Z","submitted_at":"2025-06-10T13:45:00Z","title":"HunyuanVideo-HOMA: Generic Human-Object Interaction in Multimodal Driven Human Animation","version":3},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T05:07:15.128129Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2506.08797"},"observation_digest":"sha256:f2b8e8cf1f24365cc3b23b8fbd7e6418268752d1fe81fcf9523aa666a7b27104","observation_id":"a04d0853-748d-49a4-bd63-e9d7f45dcbd0","resolution":{"observed_at":"2026-08-07T05:07:15.128129Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-07T04:27:39.669220Z","title":"Stableanimator: High-quality identity-preserving human image anima- tion","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.10568","last_updated":"2025-08-27T03:34:57Z","snapshot_observed_at":"2026-08-13T20:45:09.381829Z","submitted_at":"2025-06-12T10:58:23Z","title":"DreamActor-H1: High-Fidelity Human-Product Demonstration Video Generation via Motion-designed Diffusion Transformers","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T04:27:39.669220Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2506.10568"},"observation_digest":"sha256:574b5284e44617d189597e070dbe1c56c02b89fc9428c5a5a0098b353ecd68af","observation_id":"ebbf65d4-ddb6-4f1e-b73d-f9404f9751e8","resolution":{"observed_at":"2026-08-07T04:27:39.669220Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-06T18:05:53.912385Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.09230","last_updated":"2025-07-12T09:59:31Z","snapshot_observed_at":"2026-08-13T12:30:03.440010Z","submitted_at":"2025-07-12T09:59:31Z","title":"EgoAnimate: Generating Human Animations from Egocentric top-down Views","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-06T18:05:53.912385Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2507.09230"},"observation_digest":"sha256:4f0b2ff507063f22d07769761a59e6669addecb3c9b99689ba47b88014fb04a2","observation_id":"e1bf8f9c-fb4a-43de-b3ae-997be2a34e1a","resolution":{"observed_at":"2026-08-06T18:05:53.912385Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-06T15:47:24.294368Z","title":"Sta- bleanimator: High-quality identity-preserving human image animation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-18T01:18:47.420718Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.294368Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:faeb447eb0042a93178c167d57bc62a4245ad65a55e7790885e43e0f3450aa5e","observation_id":"d7acabfa-108d-42a4-9331-43fa61e70b37","resolution":{"observed_at":"2026-08-06T15:47:24.294368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-06T13:07:28.441776Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.20987","last_updated":"2025-07-29T04:13:25Z","snapshot_observed_at":"2026-08-20T03:01:27.596912Z","submitted_at":"2025-07-28T16:47:44Z","title":"JWB-DH-V1: Benchmark for Joint Whole-Body Talking Avatar and Speech Generation Version 1","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T13:07:28.441776Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2507.20987"},"observation_digest":"sha256:821801cda69313396946a4ba31d736b1007cbfe5fc4cda7bf04ee1eef93d821a","observation_id":"17262ae0-cc5e-4735-ba82-327d2cc7e50b","resolution":{"observed_at":"2026-08-06T13:07:28.441776Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-05T21:24:08.163993Z","title":"Stableanimator: High- quality identity-preserving human image animation.arXiv preprint arXiv:2411.17697, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.08891","last_updated":"2025-08-12T12:25:56Z","snapshot_observed_at":"2026-08-10T04:41:21.813549Z","submitted_at":"2025-08-12T12:25:56Z","title":"Preview WB-DH: Towards Whole Body Digital Human Bench for the Generation of Whole-body Talking Avatar Videos","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-05T21:24:08.163993Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2508.08891"},"observation_digest":"sha256:f8f4a53271a3f8467f85aec2a49c0135ff20202a57993e4e3ace1d99ea637a35","observation_id":"cf2d0577-8566-4b21-bb82-73ba931ad1ae","resolution":{"observed_at":"2026-08-05T21:24:08.163993Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-05T20:45:32.553066Z","title":"StableAnimator: High- quality identity-preserving human image animation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.09973","last_updated":"2025-08-13T17:40:48Z","snapshot_observed_at":"2026-08-17T11:20:26.447375Z","submitted_at":"2025-08-13T17:40:48Z","title":"PERSONA: Personalized Whole-Body 3D Avatar with Pose-Driven Deformations from a Single Image","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-05T20:45:32.553066Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2508.09973"},"observation_digest":"sha256:7522f393c228207daaabecdc1084b14658750adc0d48985f1a0b9f2c70a33ec5","observation_id":"4aa853fe-682d-4d24-8449-fbc57013bd58","resolution":{"observed_at":"2026-08-05T20:45:32.553066Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-05T10:36:57.629627Z","title":"Stableanimator: High-quality identity-preserving human image animation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03883","last_updated":"2025-09-04T04:39:21Z","snapshot_observed_at":"2026-08-15T07:22:48.616400Z","submitted_at":"2025-09-04T04:39:21Z","title":"Human Motion Video Generation: A Survey","version":1},"reference_index":208,"source":"pdf_text","source_observed_at":"2026-08-05T10:36:57.629627Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2509.03883"},"observation_digest":"sha256:8868d3a9c4cfa93159a1a3174cc8fb17860f891d9e9292750f4980d777d31c41","observation_id":"3b56b76c-161a-4d4c-a5ce-8c3990f16500","resolution":{"observed_at":"2026-08-05T10:36:57.629627Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":"2411.17697","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-07-02T16:27:09.335288Z","title":"Stableani- mator: High-quality identity-preserving human image animation.arXiv preprint arXiv:2411.17697","venue":null,"work_id":"69c88f17-9f05-4b56-9dbd-6fc0b942dfe6","year":2024},"citing_paper":{"arxiv_id":"2509.04434","last_updated":"2026-04-06T12:49:32Z","snapshot_observed_at":"2026-08-15T11:14:58.284840Z","submitted_at":"2025-09-04T17:53:03Z","title":"Durian: Dual Reference Image-Guided Portrait Animation with Attribute Transfer","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-18T18:36:03.018873Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2509.04434"},"observation_digest":"sha256:25d5af9c864e78221eb05faf6a4b397d6e39cff626c24566d2924e1797679954","observation_id":"1238999e-b2e9-412a-a55d-e078d7cb43a2","resolution":{"observed_at":"2026-05-18T18:36:44.058557Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":"2411.17697","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-07-02T16:27:09.335288Z","title":"Stableani- mator: High-quality identity-preserving human image animation.arXiv preprint arXiv:2411.17697","venue":null,"work_id":"69c88f17-9f05-4b56-9dbd-6fc0b942dfe6","year":2024},"citing_paper":{"arxiv_id":"2511.20657","last_updated":"2026-05-02T12:09:38Z","snapshot_observed_at":"2026-08-13T15:15:19.905394Z","submitted_at":"2025-10-11T07:40:36Z","title":"Intelligent Agents with Emotional Intelligence: Current Trends, Challenges, and Future Prospects","version":2},"reference_index":178,"source":"pdf_text","source_observed_at":"2026-05-18T08:11:31.181704Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2511.20657"},"observation_digest":"sha256:2d133acde2634def446fc0519e2b24e8d7fa3d0c6456b8ea0d0b39c74c292331","observation_id":"9caf126f-d203-4e54-a6e1-1465a97d94d0","resolution":{"observed_at":"2026-05-18T08:12:30.372691Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":"2411.17697","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-07-02T16:27:09.335288Z","title":"Stableani- mator: High-quality identity-preserving human image animation.arXiv preprint arXiv:2411.17697","venue":null,"work_id":"69c88f17-9f05-4b56-9dbd-6fc0b942dfe6","year":2024},"citing_paper":{"arxiv_id":"2606.06885","last_updated":"2026-06-05T04:06:42Z","snapshot_observed_at":"2026-08-16T21:13:14.435615Z","submitted_at":"2026-06-05T04:06:42Z","title":"FreeAnimate: Training-Free Human Image Animation with Preview-Guided Denoising","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-27T22:56:19.233669Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2606.06885"},"observation_digest":"sha256:fb068925ebfce89cef8a171f801b364acdcce10ad03668c2acd90cd8d4a6d62f","observation_id":"edea3663-453d-4e0b-beb2-2b97a84b3260","resolution":{"observed_at":"2026-07-02T16:07:09.347191Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":"2411.17697","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-07-02T16:27:09.335288Z","title":"Stableani- mator: High-quality identity-preserving human image animation.arXiv preprint arXiv:2411.17697","venue":null,"work_id":"69c88f17-9f05-4b56-9dbd-6fc0b942dfe6","year":2024},"citing_paper":{"arxiv_id":"2606.06903","last_updated":"2026-06-05T04:39:46Z","snapshot_observed_at":"2026-08-10T00:09:42.544772Z","submitted_at":"2026-06-05T04:39:46Z","title":"Beyond Skeletons: Learning Animation Directly from Driving Videos with Same2X Training Strategy","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-27T22:41:06.600948Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2606.06903"},"observation_digest":"sha256:1d23345d38a001b23aff8aaa70eefee36ca313a194d9f8faeadaef663f27c08e","observation_id":"75e337de-d87d-48f6-8957-3abce64af783","resolution":{"observed_at":"2026-07-02T16:27:09.336727Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-01T05:27:19.166790Z","title":"Stableanimator: High- 16 quality identity-preserving human image animation.arXiv preprint arXiv:2411.17697, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.22241","last_updated":"2026-07-30T03:34:42Z","snapshot_observed_at":"2026-08-15T05:30:04.018365Z","submitted_at":"2026-07-24T12:17:57Z","title":"AgentHOI: Multi-Agent Reasoning for Human-Object-Interaction Video Generation via Implicit Representation Alignment","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-01T05:27:19.166790Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2607.22241"},"observation_digest":"sha256:a1ae40cbe30c317a004d8919f8a127e27634b8e18a475a13f2500a839b6ef4c0","observation_id":"2ded1f96-4a1b-4f7d-b7d2-06fc4cc75434","resolution":{"observed_at":"2026-08-01T05:27:19.166790Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2411.17697/citation-record","integrity":"/paper/2411.17697/integrity","json":"/paper/2411.17697/citation-record.json","paper":"/paper/2411.17697"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2405.04233","last_updated":"2024-05-07T11:52:49Z","snapshot_observed_at":"2026-08-16T13:54:44.432320Z","submitted_at":"2024-05-07T11:52:49Z","title":"Vidu: a Highly Consistent, Dynamic and Skilled Text-to-Video Generator with Diffusion Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04233","snapshot_observed_at":"2026-08-12T11:56:28.104758Z","title":"Vidu: a highly consistent, dynamic and skilled text-to-video generator with diffusion models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.104758Z"},"links":{"cited_paper":"/paper/2405.04233","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:916a4b831b08719df69fc6834e2e12a24f87a47c866d5cb224a361695bb62b56","observation_id":"427fb57a-4732-4d77-be2e-eb4973b5f6d0","resolution":{"observed_at":"2026-08-12T11:56:28.104758Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:29.042648Z","title":"Optimal control and viscosity solutions of Hamilton-Jacobi-Bellman equa- tions","venue":null,"work_id":"3ce0199d-1f95-44f8-a0b9-05035b9ead4d","year":1997},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.109549Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:91cb3383b25e4fbed871464d5ba0e5355a8deb94c3546235b3d1f72cd8706ae5","observation_id":"7b5291f8-b470-42b4-9a40-399f42ace729","resolution":{"observed_at":"2026-08-12T11:56:29.046985Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.15127","last_updated":"2023-11-25T22:28:38Z","snapshot_observed_at":"2026-08-20T09:39:07.545813Z","submitted_at":"2023-11-25T22:28:38Z","title":"Stable Video Diffusion: Scaling Latent Video Diffusion Models to Large Datasets","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.15127","snapshot_observed_at":"2026-08-12T11:56:28.113601Z","title":"Stable video diffusion: Scaling latent video diffusion models to large datasets","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.113601Z"},"links":{"cited_paper":"/paper/2311.15127","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:4e64c46053afbc7656269df68d3a6958cb1860edf34ea7d3c82c42815c4bcdda","observation_id":"8314ebe8-b45e-4a36-8c70-8b77bac7a363","resolution":{"observed_at":"2026-08-12T11:56:28.113601Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.117483Z","title":"Video generation models as world simulators","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.117483Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:0931bc4b46f0b81b5e0ed689216e7a52ff7ddc7ba7d9caa8f4aa563ba3cbdc59","observation_id":"414e4579-cb96-43b5-b4c2-dedbff922d4d","resolution":{"observed_at":"2026-08-12T11:56:28.117483Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:29.024881Z","title":"Genera- tive modeling with phase stochastic bridges","venue":null,"work_id":"c3f1eb22-0abb-4e33-ad4f-ca9655198753","year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.121522Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:3f642997913a2c78b665b15e02356ecf0623cc1a87bf1a5422a469704c387f86","observation_id":"13d1ecc4-d302-4908-90b1-4853b61dde91","resolution":{"observed_at":"2026-08-12T11:56:29.028936Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:29.013132Z","title":"Transformers are SSMs: General- ized models and efficient algorithms through structured state space duality","venue":null,"work_id":"4610d537-fda8-4de0-9f71-88c9bc3446f0","year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.125219Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:66b72f8f908addba3a13b87d174faedf733d8ef543aab197b776a19f510990b5","observation_id":"8ae070d9-581e-4426-9d7e-b1233c6628f7","resolution":{"observed_at":"2026-08-12T11:56:29.017195Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:29.001081Z","title":"Arcface: Additive angular margin loss for deep face recognition","venue":null,"work_id":"4ecec746-9ee7-4465-be20-c6b26431260a","year":2019},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.129046Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:7bd62436bd331a632d04484f4acb0d0f33ee2e842220f096e961940c51914431","observation_id":"9123a2b7-fc2f-4027-9bb7-ce4a39a1ea48","resolution":{"observed_at":"2026-08-12T11:56:29.004964Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.990128Z","title":"Diffusion models beat gans on image synthesis","venue":null,"work_id":"95083c2c-b314-4f57-b160-b314e9a27e3e","year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.132748Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:13bcb275c47976219c8a86739e7684ae1b664dde7d00dd2282be0adc83cbd967","observation_id":"341cbff8-7012-40f5-a4b8-a558f26ef871","resolution":{"observed_at":"2026-08-12T11:56:28.993933Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.979649Z","title":"Tweedie’s formula and selection bias","venue":null,"work_id":"618e80c5-6e59-45c0-bcb9-78fc76715826","year":2011},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.136394Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:c8d6c9c78729ffa761f344dc43717a82dba55097d24ab54cec7b5e3782ebc64d","observation_id":"5c8e1169-d663-443f-bd70-7e2148ab69d3","resolution":{"observed_at":"2026-08-12T11:56:28.983326Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.968854Z","title":"Deterministic and stochastic optimal control","venue":null,"work_id":"eb9168c0-5648-46ce-ba81-6948ef07cf16","year":2012},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.140075Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:4476b39d63c434e176a75a20485c4962a87917e34e7f424a81f93814eed2b009","observation_id":"16db3910-c98a-484d-af36-25537cf6ebeb","resolution":{"observed_at":"2026-08-12T11:56:28.972608Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.957583Z","title":"Generative adversarial networks","venue":null,"work_id":"b0fdfbae-52a4-437c-b1ac-ac535b23d10e","year":2020},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.143985Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:1bb52e324036f68e535da25ffccb80f8e75eddbe2d8330ee119ab4fbbea4e908","observation_id":"d78561a5-87bb-4096-9126-becaa97dcbc6","resolution":{"observed_at":"2026-08-12T11:56:28.961481Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.03168","last_updated":"2025-02-28T14:39:17Z","snapshot_observed_at":"2026-08-18T14:48:47.548305Z","submitted_at":"2024-07-03T14:41:39Z","title":"LivePortrait: Efficient Portrait Animation with Stitching and Retargeting Control","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.03168","snapshot_observed_at":"2026-08-12T11:56:28.147493Z","title":"Livepor- trait: Efficient portrait animation with stitching and retarget- ing control","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.147493Z"},"links":{"cited_paper":"/paper/2407.03168","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:685827a934f247f82f3b6146dd7b1938a254d672e735bb7b233d0c43b6c1493d","observation_id":"56b48c8b-3209-47b8-9649-4637d1ff9e65","resolution":{"observed_at":"2026-08-12T11:56:28.147493Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.151527Z","title":"Animatediff: Animate your personalized text-to- image diffusion models without specific tuning","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.151527Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:4b3c6c1282c8367856f27c0b0af96ede63f6374af0c88534999c557edd645f38","observation_id":"1da5af5c-6fb2-4b11-8fe7-07f811f19604","resolution":{"observed_at":"2026-08-12T11:56:28.151527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.940396Z","title":"Pulid: Pure and lightning id customization via con- trastive alignment","venue":null,"work_id":"07a33742-73e8-448a-8a1e-535ee15a8b02","year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.155312Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:4dcaeed8cf2bc5bcdb636304ea0c6fc5ab5f532b3975ebec81bd937c0a809737","observation_id":"ee7bb491-6aa8-451c-82c5-a91156209187","resolution":{"observed_at":"2026-08-12T11:56:28.943951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.928731Z","title":"Facefusion","venue":null,"work_id":"9b2ba2f6-3089-467c-9123-a1130fd24b2f","year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.159061Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:0bc11f21be9dbe96c49fda827c335dadc9ca7763bdaf4d3d352f2026f14a44cb","observation_id":"6f5dd755-cca1-4408-ac21-7068086621f4","resolution":{"observed_at":"2026-08-12T11:56:28.932264Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2208.01626","last_updated":"2022-08-02T17:55:41Z","snapshot_observed_at":"2026-08-17T00:44:11.103098Z","submitted_at":"2022-08-02T17:55:41Z","title":"Prompt-to-Prompt Image Editing with Cross Attention Control","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2208.01626","snapshot_observed_at":"2026-08-12T11:56:28.162500Z","title":"Prompt-to-prompt im- age editing with cross attention control","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.162500Z"},"links":{"cited_paper":"/paper/2208.01626","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:bc2f99c9554d976788bd46d0c9ceb9414263dc7da0b553b51006cb631e5e86f8","observation_id":"eaaafb75-a76a-4077-891a-7fe398ed3bc7","resolution":{"observed_at":"2026-08-12T11:56:28.162500Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.166039Z","title":"Denoising diffu- sion probabilistic models","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.166039Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:4ef5314e4859be01ea44c4189882fb405418cda231784ee81d5fc402cb9abfb9","observation_id":"ff3a60e1-04f2-4fa6-b493-28efbdb2b634","resolution":{"observed_at":"2026-08-12T11:56:28.166039Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.909551Z","title":"Cascaded diffusion models for high fidelity image generation","venue":null,"work_id":"5c2daf7e-09e1-43a8-a72f-c316ef1f522d","year":2022},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.169273Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:3657036206d5b7376ba9818a78d6eae14ea7a61ddc4b6cc5acdfd321c4b3cd9c","observation_id":"88046198-97d1-4e1f-91a1-cd9d02914496","resolution":{"observed_at":"2026-08-12T11:56:28.913265Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2205.15868","last_updated":"2022-05-29T19:02:15Z","snapshot_observed_at":"2026-08-12T19:21:15.526321Z","submitted_at":"2022-05-29T19:02:15Z","title":"CogVideo: Large-scale Pretraining for Text-to-Video Generation via Transformers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.15868","snapshot_observed_at":"2026-08-12T11:56:28.172814Z","title":"Cogvideo: Large-scale pretraining for text-to-video generation via transformers","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.172814Z"},"links":{"cited_paper":"/paper/2205.15868","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:5a6fb6bb43ef407b8f9eb9530e43963b2a048d23cd099a02c81df6ce8935b854","observation_id":"7f851ecc-174d-4258-b105-54c068a123b7","resolution":{"observed_at":"2026-08-12T11:56:28.172814Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.898010Z","title":"Image quality metrics: Psnr vs","venue":null,"work_id":"ae48a0e9-1c15-42dc-b8e0-ea4738f6c056","year":2010},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.176214Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:68f68c2389ebee61840bc4b636d2075ed43f79a4d234f7253b4848fa44c8ac22","observation_id":"cc39b472-4744-4c8b-81a1-cdff5d8dba08","resolution":{"observed_at":"2026-08-12T11:56:28.901992Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.885829Z","title":"Lora: Low-rank adaptation of large language models","venue":null,"work_id":"78c241f1-5ca7-4bdc-a948-0622b5e5b5ea","year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.179617Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:f42485025a102f21985465197b30b1b0fcc4b77e62d5726d584a2a8bd1754fc1","observation_id":"1d6ede91-3336-4e77-a4da-fc22abf2f0fa","resolution":{"observed_at":"2026-08-12T11:56:28.890122Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.874990Z","title":"Animate anyone: Consistent and controllable image- to-video synthesis for character animation","venue":null,"work_id":"bd45090e-d33a-4a43-b135-6f5173260103","year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.183074Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:e8f99e0f87345dadb0414a413cdab561d3a527a0954adcfe6256406fd0970a4b","observation_id":"adefa506-085a-48f1-ac1a-4226974118d2","resolution":{"observed_at":"2026-08-12T11:56:28.878738Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.16771","last_updated":"2024-12-28T17:42:44Z","snapshot_observed_at":"2026-08-16T13:57:48.601133Z","submitted_at":"2024-04-25T17:23:43Z","title":"ConsistentID: Portrait Generation with Multimodal Fine-Grained Identity Preserving","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.16771","snapshot_observed_at":"2026-08-12T11:56:28.186508Z","title":"Consistentid: Portrait generation with multimodal fine-grained identity preserving","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.186508Z"},"links":{"cited_paper":"/paper/2404.16771","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:569dc9594dd00b96e5ec4655a0e1706fea4e5f1170b98edf8b44e25081b2bf5a","observation_id":"0e724309-a88e-424a-9aeb-d9d36f38a45f","resolution":{"observed_at":"2026-08-12T11:56:28.186508Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.863983Z","title":"Few- shot human motion transfer by personalized geometry and texture modeling","venue":null,"work_id":"f4bc9dee-dd27-498e-8803-392fad870c2e","year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.190090Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:89de7c658d95ef1a64d72e9c843511254a42f0eccaa16317c75b3c8cf9186ac4","observation_id":"c6384b4b-5476-4719-9b0b-ddbe46c3e2a5","resolution":{"observed_at":"2026-08-12T11:56:28.868205Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.852862Z","title":"Learning high fidelity depths of dressed humans by watching social media dance videos","venue":null,"work_id":"ad8d76f4-d3c1-4f19-908d-0324bc759511","year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.193292Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:e0a813c20575bee6f007ddfec473a74d7a871e7a9cbb7bc590309bd94716e63f","observation_id":"77bb79d5-fb50-4797-bc53-ce417e053481","resolution":{"observed_at":"2026-08-12T11:56:28.856832Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.842014Z","title":"Elucidating the design space of diffusion-based generative models","venue":null,"work_id":"b9bc1cc9-b50a-49e2-9731-bcd1481e0739","year":2022},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.196666Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:d419cfea90e34890bab99a917304696524b542f48578cff9584155011c9ced44","observation_id":"4284dab2-10d7-4e83-a0b6-bff07ee7e342","resolution":{"observed_at":"2026-08-12T11:56:28.845771Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.200371Z","title":"Auto-encoding variational bayes","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.200371Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:4a6b10213115e4388b96883653007d91e54a4d4f95cf6b5d65c93a97b13283b9","observation_id":"e3bd64c7-00ad-4322-b40f-30234ac6ca30","resolution":{"observed_at":"2026-08-12T11:56:28.200371Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.824981Z","title":"Optimal control theory: an introduction","venue":null,"work_id":"59a79548-6a14-4e85-990b-30cfccd28733","year":2004},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.203760Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:2e03933501b22852d79cfcabb3b86f289d24e32be3e5e53ed279d41853bcc563","observation_id":"312a839a-73fd-43ac-81cb-e09228e95764","resolution":{"observed_at":"2026-08-12T11:56:28.828801Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.207311Z","title":"Photomaker: Customizing re- alistic human photos via stacked id embedding","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.207311Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:328b76f1c1367e2447b34ad703a06be3c89871ecabdf747ace7bdcfecbe47906","observation_id":"e0e3a08d-a05c-4542-899b-3ab7966896c3","resolution":{"observed_at":"2026-08-12T11:56:28.207311Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03048","last_updated":"2025-05-01T09:40:21Z","snapshot_observed_at":"2026-08-21T21:36:12.450915Z","submitted_at":"2024-01-05T19:55:15Z","title":"Latte: Latent Diffusion Transformer for Video Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03048","snapshot_observed_at":"2026-08-12T11:56:28.211369Z","title":"Latte: Latent diffusion transformer for video generation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.211369Z"},"links":{"cited_paper":"/paper/2401.03048","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:a58dc3e44f103cf7cebcf58facea226efb4aadc15dab0aeb60ba60785d0e0285","observation_id":"d41cc219-d2bc-4941-94d0-955c35685e5e","resolution":{"observed_at":"2026-08-12T11:56:28.211369Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.806912Z","title":"Sdedit: Guided image synthesis and editing with stochastic differential equa- tions","venue":null,"work_id":"5716e15b-a6c1-4987-a3c4-84f28955bf55","year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.214786Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:1239bf36a9e7c6cca0e36d0df0fa49dfb0cdffbd0611fc42f412a9043b1e6381","observation_id":"f20b33bf-7f9a-4637-8a01-f7d1d1923da5","resolution":{"observed_at":"2026-08-12T11:56:28.810744Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.217969Z","title":"Improved denoising diffusion probabilistic models","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.217969Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:cfeff998a145945278e6f3d277df5a8eb5b430a3a809e166a71768a28f2d2e63","observation_id":"db01eccd-7a4e-431f-9d2b-5c7d125d36fe","resolution":{"observed_at":"2026-08-12T11:56:28.217969Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.221307Z","title":"Scalable diffusion models with transformers","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.221307Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:bdb86460cb1e1bcbd37a5ddc8b5628ec88c8b89db58948b2c6f529be4bb92d61","observation_id":"fac896fd-f51c-4362-baf6-04f3b1428822","resolution":{"observed_at":"2026-08-12T11:56:28.221307Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.06070","last_updated":"2025-03-08T08:43:03Z","snapshot_observed_at":"2026-08-19T14:04:38.615117Z","submitted_at":"2024-08-12T11:41:18Z","title":"ControlNeXt: Powerful and Efficient Control for Image and Video Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.06070","snapshot_observed_at":"2026-08-12T11:56:28.224930Z","title":"Controlnext: Powerful and effi- cient control for image and video generation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.224930Z"},"links":{"cited_paper":"/paper/2408.06070","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:b4f85b8d1feab736ca255882f04b54e310b29539b84551d113357443c031272b","observation_id":"ed4798e3-a369-4d1f-8b39-2f27f6f9561e","resolution":{"observed_at":"2026-08-12T11:56:28.224930Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.782893Z","title":"Stochastic hamilton–jacobi–bellman equations","venue":null,"work_id":"ab3e1c9f-7e21-4ae3-a683-6161c0563566","year":1992},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.229071Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:5f0037a0c47cfe8ccdb97c33f1965cbc6b4a0b3ec7946a7c214a64139af17f4d","observation_id":"7d7aec55-6016-4a19-9e84-732a788b06d9","resolution":{"observed_at":"2026-08-12T11:56:28.787001Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.232340Z","title":"Learn- ing transferable visual models from natural language super- vision","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.232340Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:1021a4d02504c4e698ebf6b07c20e94132996a3e053827177e5a000cfb0b8409","observation_id":"fac261f6-04d0-4760-9d43-8b86f8a07592","resolution":{"observed_at":"2026-08-12T11:56:28.232340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.764827Z","title":"High-resolution image syn- thesis with latent diffusion models","venue":null,"work_id":"1fce9710-b06e-4d33-a1fe-77c723c28f82","year":2022},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.235843Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:c2eac02d652cb9829d5c90b97038516c0ea9485729eabb6bbae8c3257e2a6583","observation_id":"696e0a09-b647-4b0d-957b-af0c7c6dba41","resolution":{"observed_at":"2026-08-12T11:56:28.768924Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.754170Z","title":"First order motion model for image animation","venue":null,"work_id":"437301ce-90b3-48ac-840b-01d183423e34","year":2019},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.239405Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:ec890ef0aa2bcc7f56ef239af7695f8ef3d93766fe4efe9b3d5e5d195765d2a5","observation_id":"f2688643-cfa8-4707-8633-37ca2a9e6ad5","resolution":{"observed_at":"2026-08-12T11:56:28.757861Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.743500Z","title":"Motion representations for ar- ticulated animation","venue":null,"work_id":"8133562c-b6f8-465b-8460-aaba8a8c38f3","year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.242806Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:4ee23af8bb01eec7fe6459ae040d9a63042a9a7844328290d71286eff64cd2df","observation_id":"d14127b1-70c5-4b27-aae8-25bdb8262057","resolution":{"observed_at":"2026-08-12T11:56:28.747417Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2209.14792","last_updated":"2022-09-29T13:59:46Z","snapshot_observed_at":"2026-07-06T13:57:47.051387Z","submitted_at":"2022-09-29T13:59:46Z","title":"Make-A-Video: Text-to-Video Generation without Text-Video Data","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.14792","snapshot_observed_at":"2026-08-12T11:56:28.246244Z","title":"Make-a-video: Text-to-video generation without text-video data","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.246244Z"},"links":{"cited_paper":"/paper/2209.14792","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:89f3b67fbbab6dd7d85ad03497801722eaf2326e298ae345f837c9d8a73e8011","observation_id":"09ecccc3-27f1-4ec2-9a76-ca7925f80aef","resolution":{"observed_at":"2026-08-12T11:56:28.246244Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.732637Z","title":"Denois- ing diffusion implicit models","venue":null,"work_id":"1807374e-9251-4730-813e-f90878999e41","year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.250326Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:46f15d572f1c464b639e62dd3bec154dacc9c423d8efbbff4336e33ae3b61493","observation_id":"5232307c-265f-4e97-b029-e64699f273af","resolution":{"observed_at":"2026-08-12T11:56:28.736453Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.721526Z","title":"Score-based generative modeling through stochastic differential equa- tions","venue":null,"work_id":"ca69fac0-6984-47d6-b0d9-3ccd94d5fdc2","year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.253641Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:f0b8431eae41abf60bf819079b08a9aa4db60e070b945365a17bfe3f00493273","observation_id":"806e6287-11b6-4512-91ae-295fdae741f1","resolution":{"observed_at":"2026-08-12T11:56:28.725331Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.257380Z","title":"Motioneditor: Editing video motion via content-aware diffusion","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.257380Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:b8af7dfbf7f80eba0b5566cd4dc75fa2df33d8def6e34ec8b88f27e69454a4ba","observation_id":"ade9842f-30db-4418-a303-e051e70050e8","resolution":{"observed_at":"2026-08-12T11:56:28.257380Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20325","last_updated":"2024-05-30T17:57:30Z","snapshot_observed_at":"2026-08-18T20:34:12.681533Z","submitted_at":"2024-05-30T17:57:30Z","title":"MotionFollower: Editing Video Motion via Lightweight Score-Guided Diffusion","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20325","snapshot_observed_at":"2026-08-12T11:56:28.260875Z","title":"Motionfollower: Editing video motion via lightweight score-guided diffusion","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.260875Z"},"links":{"cited_paper":"/paper/2405.20325","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:85a28d0d2af95217d4da4274366b385cbf11205209cbfcf8d4d6d45af0af3b4c","observation_id":"a1df5108-a637-419e-8a5b-e62c2296ee40","resolution":{"observed_at":"2026-08-12T11:56:28.260875Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.703400Z","title":"Plug-and-play diffusion features for text-driven image-to-image translation","venue":null,"work_id":"3f7afce5-2aa8-479d-850d-287d8aff4246","year":2023},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.264937Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:2658e2fcfdbe445534770338a91e01e53b2c2f7429767564dcb506c4eb0b7285","observation_id":"8a57e28b-dd39-4bcb-a86f-ec562a43bdd8","resolution":{"observed_at":"2026-08-12T11:56:28.707226Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.07519","last_updated":"2024-02-02T16:15:22Z","snapshot_observed_at":"2026-08-16T06:18:35.139211Z","submitted_at":"2024-01-15T07:50:18Z","title":"InstantID: Zero-shot Identity-Preserving Generation in Seconds","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.07519","snapshot_observed_at":"2026-08-12T11:56:28.268744Z","title":"Instantid: Zero-shot identity-preserving gener- ation in seconds","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.268744Z"},"links":{"cited_paper":"/paper/2401.07519","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:96863e094649c4d6ac5423e34c9b8c003d618775f52afb6ee9c37545a97f9518","observation_id":"69ea6ccc-1576-4619-a27e-cf8d2244f793","resolution":{"observed_at":"2026-08-12T11:56:28.268744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.692570Z","title":"Disco: Disentangled control for realistic human dance generation","venue":null,"work_id":"dd4fc4d8-9b1d-4124-8b0d-8cb00b8e50eb","year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.273104Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:5c53acdce290b7eaa600992b90c429935efd5bf48ff7724da01f6bae780d9e4c","observation_id":"36ab99cb-6928-421f-b633-a8f09d4cd35b","resolution":{"observed_at":"2026-08-12T11:56:28.696646Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04468","last_updated":"2024-01-09T10:12:52Z","snapshot_observed_at":"2026-08-19T04:52:16.274409Z","submitted_at":"2024-01-09T10:12:52Z","title":"MagicVideo-V2: Multi-Stage High-Aesthetic Video Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04468","snapshot_observed_at":"2026-08-12T11:56:28.276528Z","title":"Magicvideo-v2: Multi- stage high-aesthetic video generation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.276528Z"},"links":{"cited_paper":"/paper/2401.04468","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:f7cb616b88bc13c60d5e46a3f1c991ad6cb668876643e6b3aa4bff22acfd3da2","observation_id":"6dceb1b7-f330-41c2-ae5b-af93af4c1e70","resolution":{"observed_at":"2026-08-12T11:56:28.276528Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.682024Z","title":"To- wards real-world blind face restoration with generative fa- cial prior","venue":null,"work_id":"b5a13ebe-7e11-472c-abf1-df2777f517ec","year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.280202Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:1d4626b0502a8466d53f89891a2296fa866ae97ca6534488c455caeebb15a742","observation_id":"ae03e0dc-abd8-4c07-8228-18c1eb6efa3e","resolution":{"observed_at":"2026-08-12T11:56:28.685881Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01188","last_updated":"2024-06-03T10:51:10Z","snapshot_observed_at":"2026-08-18T22:36:33.661827Z","submitted_at":"2024-06-03T10:51:10Z","title":"UniAnimate: Taming Unified Video Diffusion Models for Consistent Human Image Animation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01188","snapshot_observed_at":"2026-08-12T11:56:28.283714Z","title":"Unianimate: Taming unified video diffusion mod- els for consistent human image animation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.283714Z"},"links":{"cited_paper":"/paper/2406.01188","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:6da367b2db04c39e0a4d65ac2ff98581ee2f35b1c9537cec8d2a42f495e3260f","observation_id":"d14b7f2d-e61d-4d87-a873-8bb96da56d12","resolution":{"observed_at":"2026-08-12T11:56:28.283714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.15241","last_updated":"2024-11-12T06:08:29Z","snapshot_observed_at":"2026-08-16T13:23:25.032548Z","submitted_at":"2024-08-27T17:59:41Z","title":"GenRec: Unifying Video Generation and Recognition with Diffusion Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.15241","snapshot_observed_at":"2026-08-12T11:56:28.287757Z","title":"Genrec: Unifying video generation and recognition with diffusion models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.287757Z"},"links":{"cited_paper":"/paper/2408.15241","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:f82076da1b820430be1b25da93124f42d9f725a73485c56ad27e84c86b802c84","observation_id":"b66f8379-8248-4ff4-8019-e17f6eddb9d6","resolution":{"observed_at":"2026-08-12T11:56:28.287757Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.671956Z","title":"Tune-a-video: One-shot tuning of image diffusion models for text-to-video generation","venue":null,"work_id":"41fdd873-d22e-4f6b-98da-a0ccb02a4011","year":2023},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.291537Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:a28ec19410f34cc2cdb9649d84dda66a535ae5dbd616cb04670967239b8d6217","observation_id":"c0c2c198-c4e7-4ffd-8860-b9b837629054","resolution":{"observed_at":"2026-08-12T11:56:28.675530Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.18837","last_updated":"2023-11-30T18:59:52Z","snapshot_observed_at":"2026-08-18T19:43:05.465046Z","submitted_at":"2023-11-30T18:59:52Z","title":"VIDiff: Translating Videos via Multi-Modal Instructions with Diffusion Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.18837","snapshot_observed_at":"2026-08-12T11:56:28.295063Z","title":"Vidiff: Translating videos via multi-modal instructions with diffusion models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.295063Z"},"links":{"cited_paper":"/paper/2311.18837","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:9254bceb1bc96877eac2d96dd8026d6a35d160d0256f35b9c1c2f48ba71fdb7c","observation_id":"c6a996c8-93be-47c2-ae2f-8c82c76dedb6","resolution":{"observed_at":"2026-08-12T11:56:28.295063Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.660653Z","title":"Simda: Simple diffusion adapter for efficient video generation","venue":null,"work_id":"85750e69-bc02-428e-ad12-2b3467c75711","year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.298639Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:d9d9fa1c720b124808e3bde65ea119b57d52b0ecec9de97d6ff59ad808d555a2","observation_id":"e777ad79-d5b5-405a-a1be-b4c4addc3af0","resolution":{"observed_at":"2026-08-12T11:56:28.665092Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06465","last_updated":"2024-06-10T17:02:08Z","snapshot_observed_at":"2026-08-16T23:20:31.876028Z","submitted_at":"2024-06-10T17:02:08Z","title":"AID: Adapting Image2Video Diffusion Models for Instruction-guided Video Prediction","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06465","snapshot_observed_at":"2026-08-12T11:56:28.302100Z","title":"Aid: Adapting image2video diffusion mod- els for instruction-guided video prediction","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.302100Z"},"links":{"cited_paper":"/paper/2406.06465","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:d652c6e23a1c78b45361804ec4d5f408f97f651aca698e2469e36cc98134553f","observation_id":"9e033344-2575-4369-b93d-c086c230d080","resolution":{"observed_at":"2026-08-12T11:56:28.302100Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.305857Z","title":"A survey on video dif- fusion models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.305857Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:6c126a9e7553dbfccddb465b0377cd357f7d932b1e5b0d6284e7ce8978da3d4c","observation_id":"e5491765-04a8-4a09-98a7-40fe6acbbcb2","resolution":{"observed_at":"2026-08-12T11:56:28.305857Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.641912Z","title":"Magicanimate: Temporally consistent human image animation using diffusion model","venue":null,"work_id":"ee1ba0fc-d028-48d2-b2e7-02e797fe77ab","year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.309198Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:c8b58fc0dccd8790adeb50d78728025e744f1965695a101bbde275f5e6e68463","observation_id":"f30776cd-fb9a-463e-a367-dead0302dd39","resolution":{"observed_at":"2026-08-12T11:56:28.645744Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2104.10157","last_updated":"2021-09-14T21:20:06Z","snapshot_observed_at":"2026-08-17T06:54:13.493934Z","submitted_at":"2021-04-20T17:58:03Z","title":"VideoGPT: Video Generation using VQ-VAE and Transformers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.10157","snapshot_observed_at":"2026-08-12T11:56:28.312777Z","title":"Videogpt: Video generation using vq-vae and trans- formers","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.312777Z"},"links":{"cited_paper":"/paper/2104.10157","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:00cfda362c85ed83d993667546eb6521098137b5d305447d195e7dd1da14ed01","observation_id":"4f4aa327-46cb-4661-9017-069f2f3526a7","resolution":{"observed_at":"2026-08-12T11:56:28.312777Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.02663","last_updated":"2023-12-06T12:23:36Z","snapshot_observed_at":"2026-08-16T14:37:32.614146Z","submitted_at":"2023-12-05T11:02:45Z","title":"FaceStudio: Put Your Face Everywhere in Seconds","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.02663","snapshot_observed_at":"2026-08-12T11:56:28.316458Z","title":"Facestudio: Put your face everywhere in seconds","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.316458Z"},"links":{"cited_paper":"/paper/2312.02663","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:34508fd9e7c418cd3595ef73572e6de8f30c7ce18026a1c1bf75fcc45047e0db","observation_id":"bdef9dc0-048b-4176-a1fe-69655934de89","resolution":{"observed_at":"2026-08-12T11:56:28.316458Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.630440Z","title":"Effec- tive whole-body pose estimation with two-stages distillation","venue":null,"work_id":"e93fe2f7-deee-46a0-af60-cf8cff4b0cd0","year":2023},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.320102Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:ecbea6123597110c723240f6f6eaff62668d8a91daa5163cc4b42e6ae0c80f74","observation_id":"53f182a8-9ea1-4594-95f5-11755ae82948","resolution":{"observed_at":"2026-08-12T11:56:28.634272Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.06721","last_updated":"2023-08-13T08:34:51Z","snapshot_observed_at":"2026-07-06T16:05:39.158819Z","submitted_at":"2023-08-13T08:34:51Z","title":"IP-Adapter: Text Compatible Image Prompt Adapter for Text-to-Image Diffusion Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.06721","snapshot_observed_at":"2026-08-12T11:56:28.323435Z","title":"Ip- adapter: Text compatible image prompt adapter for text-to- image diffusion models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.323435Z"},"links":{"cited_paper":"/paper/2308.06721","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:1d809ccaebd272391f3eec442356ffd1073f4886e47032af367e6a3c696078f3","observation_id":"1f14646a-69ed-4e58-965a-c92d5e6bc456","resolution":{"observed_at":"2026-08-12T11:56:28.323435Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.327917Z","title":"Magvit: Masked generative video transformer","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.327917Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:f415847f7627ebb3ec6f8d29b800de52d359ff0210ffc72207877bc770da98ed","observation_id":"e79c66fb-deb4-43d2-9660-d1aaede010ee","resolution":{"observed_at":"2026-08-12T11:56:28.327917Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.331327Z","title":"Adding conditional control to text-to-image diffusion models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.331327Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:b7f51c6afeb46b2346a68b408bdf544ce71de3c094cd645d0c1df202d1182fa3","observation_id":"1aee2736-b81e-4129-8a3d-edd05663c096","resolution":{"observed_at":"2026-08-12T11:56:28.331327Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.19680","last_updated":"2025-06-27T10:06:13Z","snapshot_observed_at":"2026-08-16T13:38:51.127960Z","submitted_at":"2024-06-28T06:40:53Z","title":"MimicMotion: High-Quality Human Motion Video Generation with Confidence-aware Pose Guidance","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.19680","snapshot_observed_at":"2026-08-12T11:56:28.334665Z","title":"Mim- icmotion: High-quality human motion video generation with confidence-aware pose guidance","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.334665Z"},"links":{"cited_paper":"/paper/2406.19680","citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:240b298c910bad7c48dcb30a7801c1045d7999d09e9b1d35bac7218e313cccc6","observation_id":"cff38848-3d5d-4678-a73c-ccbbe79da686","resolution":{"observed_at":"2026-08-12T11:56:28.334665Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.606569Z","title":"Chan, Chongyi Li, and Chen Change Loy","venue":null,"work_id":"9e89b54f-f0e3-411d-b1d8-deb6713a08d7","year":2022},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.338546Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:83b9d9353031f03a18ab04517821c72a0c3521daabf55e21f0e8810a10ebdf90","observation_id":"a5159717-fd08-436d-8981-ae8fb45dd7ca","resolution":{"observed_at":"2026-08-12T11:56:28.610229Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.595227Z","title":"Champ: Controllable and consistent human image animation with 3d parametric guidance","venue":null,"work_id":"38768306-fdb9-4c28-8742-82ca56638369","year":2024},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.342095Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:e4ea969341146e2d7f3e988051c765fcb80c3ce6e723d3eb0046b4556e406ae8","observation_id":"09381c08-3393-4e41-b574-f903487ebd45","resolution":{"observed_at":"2026-08-12T11:56:28.599270Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:56:28.581244Z","title":"We can see that our proposed components can sig- nificantly facilitate the performance of different backbone- based models, particularly in the facial regions","venue":null,"work_id":"c883bbdf-f43d-408c-b5ed-72d07f4ae93a","year":null},"citing_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-12T11:56:28.346922Z"},"links":{"citing_paper":"/paper/2411.17697"},"observation_digest":"sha256:0876d75f9fee37331e84ca68b0a97a1e46141fa6078f0e8b1c6ffcaee3552a4d","observation_id":"6ce20fc0-ba36-4b2c-80f3-df93a53af312","resolution":{"observed_at":"2026-08-12T11:56:28.587046Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-15T07:21:59.072249Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation"},"reference_resolution":{"displayed":67,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":32,"verified_exact":0,"verified_fuzzy":35},"total_outbound_references":67},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 22 August 2026, this Paper Citation Record lists 67 of 67 outbound references and 17 inbound Pith citation observations for arXiv:2411.17697."}