{"as_of":"2026-08-14T21:41:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d02e6755ef858588189ed0bfedfa439ad9cad36e78cca330e93ae29299180ef2","coverage":[{"denominator":73,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":73,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T04:44:08.392472Z","state":"measured"},{"denominator":73,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":73,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2412.01106/citation-record","integrity":"/paper/2412.01106/integrity","json":"/paper/2412.01106/citation-record.json","paper":"/paper/2412.01106"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.198251Z","title":"Gaussian shell maps for efficient 3d human generation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.198251Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:85f7f39ceb9c96cbf39c045ec7797f57fa200751e2185427ad180391c04f61d3","observation_id":"347e9a91-8225-480d-9d61-8cce15c038f8","resolution":{"observed_at":"2026-08-12T04:44:08.198251Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:09.026815Z","title":"Single-image 3d human digitization with shape-guided diffusion","venue":null,"work_id":"b614f1e5-b05d-45c4-bd32-bdd88b2a3e7c","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.201649Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:6720f0dbc63d9d7454d80fce4b4c239b91b2ae2eadb65203677765dd1002288a","observation_id":"f16a3f02-58e1-4881-8d65-b12be1ba9f1b","resolution":{"observed_at":"2026-08-12T04:44:09.029625Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.15127","last_updated":"2023-11-25T22:28:38Z","snapshot_observed_at":"2026-08-07T21:47:08.589400Z","submitted_at":"2023-11-25T22:28:38Z","title":"Stable Video Diffusion: Scaling Latent Video Diffusion Models to Large Datasets","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.15127","snapshot_observed_at":"2026-08-12T04:44:08.204504Z","title":"Stable video diffusion: Scaling latent video diffusion models to large datasets","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.204504Z"},"links":{"cited_paper":"/paper/2311.15127","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:a16fc405d58d93f522ff0f96e859a42a07fcfc47a52845bfd4ca492c0dac42e8","observation_id":"6e2bab96-5313-4cb4-83c8-32003370b88c","resolution":{"observed_at":"2026-08-12T04:44:08.204504Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:09.017657Z","title":"Structured 3d features for reconstructing control- lable avatars","venue":null,"work_id":"801f6f90-1d3d-4a19-8033-71dc6daae53e","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.207970Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:2fb15f84bb1c207bb265e8aad3cb4a64265917717c3c1184fe32b988b34242d1","observation_id":"06bfdf63-98ff-4229-9e34-d63d279aabcf","resolution":{"observed_at":"2026-08-12T04:44:09.021816Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:09.009062Z","title":"Portrait4d-v2: Pseudo multi-view data creates better 4d head synthesizer","venue":null,"work_id":"655736de-a4dc-45ab-a3af-5a89c1e4480d","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.210804Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:4b887ad209866641e1f7cbc282b7a898b6903df78edb4e74a5fb6705cc1794ef","observation_id":"c7efecd7-79ee-436c-b457-00f9bb7f7baf","resolution":{"observed_at":"2026-08-12T04:44:09.012537Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:09.000870Z","title":"Ag3d: Learning to gener- ate 3d avatars from 2d image collections","venue":null,"work_id":"02d0cb96-d7a1-4a80-8860-cda584a6953c","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.213754Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:b03d675c791556ff025a865616646aadb94f49fdeed3960234810ccfdb818ed2","observation_id":"eec32bbf-2a6e-4f5e-a16b-ca5438c3d9e1","resolution":{"observed_at":"2026-08-12T04:44:09.003891Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.992019Z","title":"Dreamoving: A human video generation framework based on diffusion models","venue":null,"work_id":"381e73c5-f077-4026-8b30-00c56b8af8f1","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.216827Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:b650358cb27b13d41b30bcc7383c82d3bc116edb3d2fcf7b7552d1c571f6ad2b","observation_id":"cea38d1c-b277-4fde-a570-0326fef0ff90","resolution":{"observed_at":"2026-08-12T04:44:08.995120Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.03168","last_updated":"2025-02-28T14:39:17Z","snapshot_observed_at":"2026-08-13T23:43:50.484861Z","submitted_at":"2024-07-03T14:41:39Z","title":"LivePortrait: Efficient Portrait Animation with Stitching and Retargeting Control","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.03168","snapshot_observed_at":"2026-08-12T04:44:08.219779Z","title":"Livepor- trait: Efficient portrait animation with stitching and retarget- ing control","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.219779Z"},"links":{"cited_paper":"/paper/2407.03168","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:70af6e02ca33a07fe8359b588f2841ff4c55f1cc998962bb9a0eb639eb7fc877","observation_id":"4ad883ff-53aa-43da-8858-3c94c3164203","resolution":{"observed_at":"2026-08-12T04:44:08.219779Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.984933Z","title":"The re- lightables: V olumetric performance capture of humans with realistic relighting","venue":null,"work_id":"604f92d9-1c42-4015-8b63-f254110375d9","year":2019},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.222979Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:4de156c2ce0031be46ad61a29b745ebdbd74f71e89fd80226c5eb0231ff9ae2b","observation_id":"6b2cfb98-f781-4a9a-bf54-d11edbf5b704","resolution":{"observed_at":"2026-08-12T04:44:08.987616Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.977130Z","title":"Animatediff: Animate your personalized text-to- image diffusion models without specific tuning, 2023","venue":null,"work_id":"618119ff-bc40-43d5-ae12-fd7407ca1fd3","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.225694Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:7bdd075e3c9135d2ca626a8e1a8b5c98bc83addbb3b11fea3bf1e07b3eb64bb5","observation_id":"307b81ca-bfa0-4001-bbaf-95a7d9a908a5","resolution":{"observed_at":"2026-08-12T04:44:08.980063Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.968545Z","title":"Real-time deep dynamic characters","venue":null,"work_id":"be3f0ebc-4e32-43a7-a5df-b9afa1f3bdb6","year":2021},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.228374Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:857974729372071e4e3ae07e8e57d1d88cfd1f0097ad8c7eca0f9dbcb8cdfeec","observation_id":"1df59436-3eeb-48cb-ab7a-857c903ef3ff","resolution":{"observed_at":"2026-08-12T04:44:08.971540Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.960595Z","title":"Towards measuring fairness in ai: the casual conversations dataset","venue":null,"work_id":"df8b0b0e-ea0f-4baf-8a31-66e6acb30a3c","year":2021},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.230940Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:6d960c7d2b3307ea347625664ba3d843190a9fe0be0201792187165259f1abd5","observation_id":"54d7414b-a521-43a0-b919-dac3f6fca3b3","resolution":{"observed_at":"2026-08-12T04:44:08.963698Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.951742Z","title":"Sith: Single-view tex- tured human reconstruction with image-conditioned diffu- sion","venue":null,"work_id":"f302a5bc-5e2e-40f6-8ccd-41d86cab1ac9","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.233661Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:85db181824e299d4e0fba9f6530a8d5a49c861b3f2b69288a2966aa3377446f6","observation_id":"6fb070c2-d882-4ff6-ac5f-29cbdc83158b","resolution":{"observed_at":"2026-08-12T04:44:08.955529Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.943345Z","title":"Expres- sive gaussian human avatars from monocular rgb video","venue":null,"work_id":"27cb8d5f-ab35-4a24-a612-eda847a08c4e","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.236222Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:b63569cb5b400aa8cab4025b7a515bec401c34b01db313926ebc20876529307f","observation_id":"0bd69fd3-6a9f-49b4-a033-73c2dc2277b3","resolution":{"observed_at":"2026-08-12T04:44:08.946435Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.936141Z","title":"Animate anyone: Consistent and controllable image- to-video synthesis for character animation","venue":null,"work_id":"1b4311a3-e292-4015-a855-e6f3b7a0d210","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.238765Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:fed58ac350d33462dae342934defe3f069480582a854950429d155463aee32ff","observation_id":"08d07b3e-bf75-46c3-9a71-edb137df6490","resolution":{"observed_at":"2026-08-12T04:44:08.938739Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.928393Z","title":"Gaussianavatar: Towards realistic human avatar model- ing from a single video via animatable 3d gaussians","venue":null,"work_id":"78bd194e-7820-42f6-b622-be7fd6bb2f0b","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.241382Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:9174ead4ca9aad38681e8ca1529cf4594592441df6d8d3fefe3d02bc924b57bb","observation_id":"2126f106-9aab-44b9-87ad-074e1be076c5","resolution":{"observed_at":"2026-08-12T04:44:08.931603Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.919749Z","title":"Gauhuman: Articu- lated gaussian splatting from monocular human videos","venue":null,"work_id":"495f4183-defb-44cb-8091-d4ade7ba534e","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.243912Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:df0740310f266df56e27cd0adb25b79d2c61c0afa98a19b97548feb9148a11be","observation_id":"d900e3e6-3bca-4748-a38c-16333b40a115","resolution":{"observed_at":"2026-08-12T04:44:08.923068Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.911005Z","title":"One-shot implicit animatable avatars with model- based priors","venue":null,"work_id":"a88a0d63-35d6-4ad5-8086-80a052acf552","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.246631Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:c531fe77f46bc4dc93ead474fcd21a34acaf7341b6378c94979c466feea288c2","observation_id":"2cd4efe6-56db-4ac1-8d28-195ade7d011d","resolution":{"observed_at":"2026-08-12T04:44:08.914112Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.902397Z","title":"Tech: Text-guided reconstruction of lifelike clothed humans","venue":null,"work_id":"59e57429-d652-4ec4-a4aa-903b5e5059dd","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.249453Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:4610a53a5ace0fcbd7481fdb5be9bd0d42d88f7e56146cbfd5e30e71633d4958","observation_id":"c3880289-0906-4a48-b821-19ce23d40625","resolution":{"observed_at":"2026-08-12T04:44:08.905355Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.893885Z","title":"Make-your-anchor: A diffusion-based 2d avatar generation framework","venue":null,"work_id":"919336ef-c320-4e83-ab04-b8e2365f1f08","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.252214Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:e73349fb7a38aa03f5ccce225f51902a2e1b336b50c6b2bd16e2ccd9d2e5c8ac","observation_id":"c8eec7f8-c7f5-40ba-835f-e8359c5bb45d","resolution":{"observed_at":"2026-08-12T04:44:08.897083Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.884973Z","title":"Humanrf: High-fidelity neural radiance fields for humans in motion","venue":null,"work_id":"b66d04ce-f6a0-46be-bea8-f1aecb881cc9","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.254934Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:f9430e254921a030388df30fe04cbc92f4359c47fa1b794cf82dd0f6e9fc6b5c","observation_id":"94f50c39-e181-4c25-9415-414a2190a82a","resolution":{"observed_at":"2026-08-12T04:44:08.888363Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.876734Z","title":"Humangen: Generating hu- man radiance fields with explicit priors","venue":null,"work_id":"34667b97-e102-426c-9394-2b94d000e35a","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.257104Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:eb0ea057ba98a68feaaecfe82b97e719ad6c3996c02fd0291e4602c1257bc460","observation_id":"a4de71cd-f836-4308-ac7d-22c436455ba4","resolution":{"observed_at":"2026-08-12T04:44:08.879289Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.259366Z","title":"Neuman: Neural human radiance field from a single video","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.259366Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:9578ea8c6c15c3c31280c38efee8d9092014d6791abc39a5238cfe7bc416f390","observation_id":"02c72955-4242-41a1-a847-fbf7cc66634b","resolution":{"observed_at":"2026-08-12T04:44:08.259366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.864815Z","title":"Hifi4g: 9 High-fidelity human performance rendering via compact gaussian splatting","venue":null,"work_id":"821f7543-727e-4983-9ffc-c6a81cb2da8a","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.261555Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:8b7186ccc283f60628fb2e605253dece983408f04c3f638a5436441f45a6592c","observation_id":"bbcdc317-7abb-4759-829b-c4307d1e47d7","resolution":{"observed_at":"2026-08-12T04:44:08.867569Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.263699Z","title":"3d gaussian splatting for real-time radiance field rendering","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.263699Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:6128434e1b3e3207a7c1d20561bd05ca23ae2790254f5d636f62763027532475","observation_id":"7d716f39-d72c-4347-9fd9-d19e459acc6a","resolution":{"observed_at":"2026-08-12T04:44:08.263699Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6980","last_updated":"2017-01-30T01:27:54Z","snapshot_observed_at":"2026-08-14T18:51:16.666127Z","submitted_at":"2014-12-22T13:54:29Z","title":"Adam: A Method for Stochastic Optimization","version":9},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.6980","snapshot_observed_at":"2026-08-12T04:44:08.265920Z","title":"Adam: A method for stochastic opti- mization","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.265920Z"},"links":{"cited_paper":"/paper/1412.6980","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:be25967e3853cab262cd8e435c6a7b545375203a5eff4f30e99e8a27e2610a65","observation_id":"f8078669-b4f7-40b2-a37b-e9423386da05","resolution":{"observed_at":"2026-08-12T04:44:08.265920Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.268634Z","title":"HUGS: Human gaussian splatting","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.268634Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:4f3402c98d019d81947cbfd3520b3391a72ef39e6ba9a7d1eb5525d6a4fd1b18","observation_id":"850395c2-81ff-4249-934f-bf1ddf93db73","resolution":{"observed_at":"2026-08-12T04:44:08.268634Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.847945Z","title":"Modular primitives for high-performance differentiable rendering","venue":null,"work_id":"88df8c1c-e073-404b-b964-7e1de64e74e1","year":2020},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.270770Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:42e74ccc78151b3fc33fd9ce544c20d8c8c2b21e9c25f3dd446eea0cfb230562","observation_id":"c558f39c-b550-4348-a5b8-5cd0e7ae9572","resolution":{"observed_at":"2026-08-12T04:44:08.851025Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.272881Z","title":"Gart: Gaussian articulated template mod- els","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.272881Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:3524ca7168faebcfd4b5e11102b5603e98299cc156d4b3f598ec2ea314c5a9c0","observation_id":"de6b3bce-7afd-4fea-9e00-469efe760d28","resolution":{"observed_at":"2026-08-12T04:44:08.272881Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.10141","last_updated":"2026-06-30T17:32:07Z","snapshot_observed_at":"2026-08-13T20:46:52.585162Z","submitted_at":"2024-09-16T10:13:06Z","title":"PSHuman: Photorealistic Single-image 3D Human Reconstruction using Cross-Scale Multiview Diffusion and Explicit Remeshing","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.10141","snapshot_observed_at":"2026-08-12T04:44:08.275571Z","title":"Pshuman: Photorealistic single-view human reconstruction using cross-scale diffusion","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.275571Z"},"links":{"cited_paper":"/paper/2409.10141","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:b036433a5361abd2aa2cc056cd7923a88e8a83192168da51e0bcc0b263f47540","observation_id":"42329b99-f1d6-4101-a89c-d60a66e15cd8","resolution":{"observed_at":"2026-08-12T04:44:08.275571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.835054Z","title":"Learning a model of facial shape and expression from 4d scans","venue":null,"work_id":"04c0ad31-caea-497e-94aa-b0bff9df48ee","year":2017},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.278592Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:2d03ca0048c57441ce2f573daddf94b1211895749b441258ebedc3dc1f1cc6fe","observation_id":"3dd46b3f-f4ba-4c52-8600-f0735bd77efa","resolution":{"observed_at":"2026-08-12T04:44:08.838174Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.281141Z","title":"Neural actor: Neural free-view synthesis of human actors with pose con- trol","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.281141Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:736c81952b96a015ef141768aeabcd6a92539475bd6d271e633e4093313a4034","observation_id":"73f340fa-40f4-4ba0-a6d0-fc6a9eca85fe","resolution":{"observed_at":"2026-08-12T04:44:08.281141Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.284046Z","title":"Human-vdm: Learning single-image 3d human gaussian splatting from video diffusion models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.284046Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:8b0662856e3d6310672d93e31495dabb00fafd072115c7e8d6462cddbc41f8a2","observation_id":"943fc592-a48a-4d75-a355-e1900fcbbc5c","resolution":{"observed_at":"2026-08-12T04:44:08.284046Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.821144Z","title":null,"venue":null,"work_id":"dcffd665-0a9c-4969-85d1-d52ad1c0f017","year":2015},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.286778Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:5bdb3b0ffad1775f097b94711df77af6bfdef193507cb2ba5513ab720b60d01f","observation_id":"73d71ed5-cf1b-4ec3-8734-ffd93bc41a16","resolution":{"observed_at":"2026-08-12T04:44:08.824216Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.811940Z","title":"The contextual loss for image transformation with non-aligned data","venue":null,"work_id":"95918843-80e7-445f-b358-9bb4f6345b77","year":2018},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.289924Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:6d1114dfdde9997ec534e54c3244ab04cde21eaae966acce512ada78db9d5254","observation_id":"9bd6e639-b275-4d51-b65d-969035a2c0bb","resolution":{"observed_at":"2026-08-12T04:44:08.815241Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.802434Z","title":"Expressive whole-body 3D gaussian avatar","venue":null,"work_id":"cc6fd457-36e9-4ccd-859b-108cb15e1fee","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.292663Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:7cc5cbaa28505004e4948452fb72af86fbc2603008e5b32acd9200bc739e84c9","observation_id":"34fd06a0-c580-459d-9680-4064e8ee0ec4","resolution":{"observed_at":"2026-08-12T04:44:08.805876Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.792861Z","title":"Human gaussian splatting: Real-time rendering of animatable avatars","venue":null,"work_id":"c0def66f-6d10-4d4f-b68d-c6a00aeac6a2","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.295296Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:e9a855b061687cf1a51121471222911d6e0450042bd43a0b9c61d591238feaaf","observation_id":"9dfdf4c0-aa0a-4d52-aabf-212949d4611c","resolution":{"observed_at":"2026-08-12T04:44:08.796214Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12459","last_updated":"2024-10-30T12:50:27Z","snapshot_observed_at":"2026-08-12T23:40:15.918898Z","submitted_at":"2024-06-18T10:05:33Z","title":"HumanSplat: Generalizable Single-Image Human Gaussian Splatting with Structure Priors","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12459","snapshot_observed_at":"2026-08-12T04:44:08.297923Z","title":"Humansplat: Generalizable single-image human gaussian splatting with structure priors","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.297923Z"},"links":{"cited_paper":"/paper/2406.12459","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:10cd8c1915f8327386abf9a4207c5339bdb4ea793eb1525e5b5817bcb47761fe","observation_id":"03599f8f-22bb-4e9e-b5e4-55f3d1a466ce","resolution":{"observed_at":"2026-08-12T04:44:08.297923Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.784593Z","title":"Ash: Animatable gaussian splats for efficient and photoreal human rendering","venue":null,"work_id":"f33c5428-66f7-4280-9b4a-a423cb7c56fc","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.300787Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:6c88e5ef024c8dcc6bf361dc3535304b1dbfc4ccc78655f0f5360148180368bb","observation_id":"09a26c2d-8c01-4451-af1d-cb43cce659c6","resolution":{"observed_at":"2026-08-12T04:44:08.787584Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.776409Z","title":null,"venue":null,"work_id":"7b780879-134e-4f48-a437-6342be49f1eb","year":2019},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.303528Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:7bc4332991c01b5b0444302252bef74f0a0684713218e7e0543e67f21711f5b4","observation_id":"cc69aafe-d185-452b-b4d3-e6cc62ebedc4","resolution":{"observed_at":"2026-08-12T04:44:08.779167Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.767411Z","title":"Neural body: Implicit neural representations with structured latent codes for novel view synthesis of dynamic humans","venue":null,"work_id":"2eb53c0c-290a-4344-8e75-c0cfce134555","year":2021},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.306097Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:83f63b3b6f39cb58a514d14612d89b6171b84e9aee75d93b9ad21ba0827abab7","observation_id":"3fc15d7f-f1d5-42e2-a156-3bc3104a56b7","resolution":{"observed_at":"2026-08-12T04:44:08.771130Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2209.14988","last_updated":"2022-09-29T17:50:40Z","snapshot_observed_at":"2026-07-06T13:57:54.539656Z","submitted_at":"2022-09-29T17:50:40Z","title":"DreamFusion: Text-to-3D using 2D Diffusion","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.14988","snapshot_observed_at":"2026-08-12T04:44:08.308734Z","title":"Dreamfusion: Text-to-3d using 2d diffusion","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.308734Z"},"links":{"cited_paper":"/paper/2209.14988","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:cfbc0194562797bf0b629dcc8e1f50b509de72c0d39ea547c59935615685021f","observation_id":"1f106321-713e-4929-8cd9-5a2491150362","resolution":{"observed_at":"2026-08-12T04:44:08.308734Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.757751Z","title":"3dgs-avatar: Animatable avatars via deformable 3d gaussian splatting","venue":null,"work_id":"4b18ecd9-d95f-4bef-be4a-9723cb911d61","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.311675Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:a637486f9c5a64d876e1a66fec86daf562ead29ac38a4801758be48122ec866b","observation_id":"d6b24401-4518-4325-b6e0-22f9fa523230","resolution":{"observed_at":"2026-08-12T04:44:08.761366Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.314302Z","title":"Learning transferable visual models from natural language supervi- sion","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.314302Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:4692cd6fec51bb5a25fe649d82e4f13b719ceb11f723857d5e5da3e5db8a3c68","observation_id":"08927b38-321d-449b-aa24-2337eea79b00","resolution":{"observed_at":"2026-08-12T04:44:08.314302Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.316897Z","title":"High-resolution image synthesis with latent diffusion models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.316897Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:69f9bf4a28e52e25e3a058aa8604d417f305e805694613a020876e7fb3b1f584","observation_id":"064bf7ab-9bb9-4660-8fc3-df2616132416","resolution":{"observed_at":"2026-08-12T04:44:08.316897Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.737983Z","title":null,"venue":null,"work_id":"c5271618-03ef-4e47-8bc1-199a096edc1f","year":null},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.319532Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:0167defa1c6d6845b77ba55c95e4d784d9d7de55e136a6991cb0c09c4b3a82aa","observation_id":"d81a3a90-9333-4a17-b302-870a8f9a0852","resolution":{"observed_at":"2026-08-12T04:44:08.741347Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.729510Z","title":"Pifu: Pixel-aligned implicit function for high-resolution clothed human digitiza- tion","venue":null,"work_id":"fa51115a-e8d1-4ea1-9ea3-d21dc8b0e632","year":2019},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.322376Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:38f48c6c00d00a13b76da9968e04bbe0c4ac7ed57f4624f830e9bac76446257d","observation_id":"0e673204-109c-4dbd-ac09-3896067a68d9","resolution":{"observed_at":"2026-08-12T04:44:08.732551Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.721666Z","title":"Pifuhd: Multi-level pixel-aligned implicit function for high-resolution 3d human digitization","venue":null,"work_id":"693b7d33-efed-4a59-902b-d1281ce70cab","year":2020},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.325130Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:c58287f613b2aedfc599d9037ca01f28b59995f50ae34ca49e6b2d0c9c7216d5","observation_id":"2877f847-7298-4bf7-a8e4-4fa8d69927b1","resolution":{"observed_at":"2026-08-12T04:44:08.724413Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.712292Z","title":"10 SplattingAvatar: Realistic Real-Time Human Avatars with Mesh-Embedded Gaussian Splatting","venue":null,"work_id":"1098035e-0ebb-43d4-9d2d-8d8fbab1bdfd","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.327863Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:71b2038b53a3a9932cdd66afa9137b353ccced206be3c4cffc5a32483b9ee126","observation_id":"d172ffa9-7b40-4e7e-b0b2-b8327cd85b34","resolution":{"observed_at":"2026-08-12T04:44:08.715526Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1409.1556","last_updated":"2015-04-10T16:25:04Z","snapshot_observed_at":"2026-08-14T04:13:35.056640Z","submitted_at":"2014-09-04T19:48:04Z","title":"Very Deep Convolutional Networks for Large-Scale Image Recognition","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1409.1556","snapshot_observed_at":"2026-08-12T04:44:08.330680Z","title":"Very deep convolutional networks for large-scale image recognition","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.330680Z"},"links":{"cited_paper":"/paper/1409.1556","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:50b8925d2ae615fbd3947e69d5940684ef1240fd2c7186c4782af0949fff8c96","observation_id":"e61c03ed-d521-492f-a48f-dd8f7e66a857","resolution":{"observed_at":"2026-08-12T04:44:08.330680Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.16653","last_updated":"2024-03-29T08:39:23Z","snapshot_observed_at":"2026-07-06T16:25:05.493379Z","submitted_at":"2023-09-28T17:55:05Z","title":"DreamGaussian: Generative Gaussian Splatting for Efficient 3D Content Creation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.16653","snapshot_observed_at":"2026-08-12T04:44:08.333739Z","title":"Dreamgaussian: Generative gaussian splatting for effi- cient 3d content creation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.333739Z"},"links":{"cited_paper":"/paper/2309.16653","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:b933cdb8ada90a87c2d228a1e381a2b1a0ea1589f83874028f8a36124f3d5056","observation_id":"1ef1f7c7-6d1b-4e8c-a964-28efe60df26a","resolution":{"observed_at":"2026-08-12T04:44:08.333739Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.336647Z","title":"Hu- mannerf: Free-viewpoint rendering of moving people from monocular video","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.336647Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:691b94cd1ab1a0e004495f5e0f932c6024ca1a33a33c50bd331d7469c9922999","observation_id":"112ae2a0-e4e8-4d19-a627-81ba6be578a5","resolution":{"observed_at":"2026-08-12T04:44:08.336647Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.699025Z","title":"Driv- able avatar clothing: Faithful full-body telepresence with dy- namic clothing driven by sparse rgb-d input","venue":null,"work_id":"c1b88c8f-f398-4468-ab55-a27185697c1c","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.339258Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:e96d6d46cf9485b2203cd763d015aec7e7f3c6ddb70f813411fa9b101e12df6a","observation_id":"ae0e4a63-4455-4c85-b446-5d4dd90548e5","resolution":{"observed_at":"2026-08-12T04:44:08.702048Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.690704Z","title":"Flashavatar: High-fidelity head avatar with efficient gaussian embedding","venue":null,"work_id":"ca3d2ccf-5dfb-4e0f-94ef-c92f6b4da85e","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.342055Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:5bbbe8440b21f205a6142a4369034ba4ae15d633434bb0fd291d8cff6ba89eb9","observation_id":"86b8334c-c176-436f-9a7e-ea6193e75cb9","resolution":{"observed_at":"2026-08-12T04:44:08.693622Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.682561Z","title":"Get3dhuman: Lifting stylegan-human into a 3d generative model using pixel-aligned reconstruction priors","venue":null,"work_id":"81e2bc6d-671d-4222-9b1b-d6cc50b36665","year":null},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.344316Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:7e8b5a93855da5e38229c1967a483dfa9e621dc0852ae4c4af39f66e6942f21a","observation_id":"c9b68d37-0820-4184-ba41-512413e4c892","resolution":{"observed_at":"2026-08-12T04:44:08.685636Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.675412Z","title":"Icon: Implicit clothed humans obtained from nor- mals","venue":null,"work_id":"0c36615d-7412-42d9-b9c7-99e341093532","year":2022},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.346952Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:3d85698b0f72cdf022b38a1776edb5628b1834b350dd33a51c6d93c055a97bd2","observation_id":"3bed377f-3ded-477e-a393-03aed434472f","resolution":{"observed_at":"2026-08-12T04:44:08.677929Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.667144Z","title":"Econ: Explicit clothed humans optimized via normal integration","venue":null,"work_id":"aeb2773c-f37c-4c0e-ad72-d532228858f3","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.349123Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:09c8573c159ae32fe2ffd3232ad475d6b4cfda1dfe60dbe56d8a05d07871ba8d","observation_id":"10791162-0d7d-4784-9fbc-8a4cc8bc2f73","resolution":{"observed_at":"2026-08-12T04:44:08.670609Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.658374Z","title":"Puzzleavatar: Assembling 3d avatars from personal albums","venue":null,"work_id":"b44f1295-82ee-44bc-aa45-97b43f443d92","year":null},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.351345Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:27270bcb63496827adabb6be155ee3515fefff37b17625e86feb616bb3f8a6b5","observation_id":"e0ec72c7-2351-44f9-b2b9-3061caf891b3","resolution":{"observed_at":"2026-08-12T04:44:08.661760Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.354033Z","title":"Magicanimate: Temporally consistent human im- age animation using diffusion model","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.354033Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:d140e6b3ace8b181bfbffb938ae232c8abd1cf453f2ef64a34c1b309be3b3812","observation_id":"2d682c8c-3c0d-4bb3-b281-9dfa9c6d36b3","resolution":{"observed_at":"2026-08-12T04:44:08.354033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.645380Z","title":"Have-fun: Human avatar reconstruction from few-shot unconstrained images","venue":null,"work_id":"0c8fa25f-1abc-4ced-bd39-7b632fe3687e","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.356297Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:075d184c2800c70165ab0a7825d38f83a683bf6aa58a40a6b221b210d013bfa3","observation_id":"40dd4dc5-99d7-44c8-bd01-10c78e682179","resolution":{"observed_at":"2026-08-12T04:44:08.648424Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.635192Z","title":"Effec- tive whole-body pose estimation with two-stages distillation","venue":null,"work_id":"3bf4e2f4-557d-4310-9797-269e9c4a75eb","year":2023},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.358652Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:f10497358cfc97c963d7f7a5f152e0ba10a3d8b7732c40b27263352e6ddc6e93","observation_id":"d961390f-98eb-4671-80b5-f15de63da86d","resolution":{"observed_at":"2026-08-12T04:44:08.639057Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.361535Z","title":"Robots learn social skills: End-to-end learning of co-speech gesture generation for humanoid robots","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.361535Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:ad9f253a3696c00a262185ef3cae18037abeee07897d4936d22451dcceb177c1","observation_id":"779a5f55-d6d1-4cef-a8a6-66c92a56c0c2","resolution":{"observed_at":"2026-08-12T04:44:08.361535Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.621521Z","title":"Humanref: Single image to 3d human gen- eration via reference-guided diffusion","venue":null,"work_id":"7a06cedc-4486-4551-91e5-ce16b0b92a0d","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.364334Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:101fd800ee28effe93a3d19469f0b15884c079305618dab383ae2a58b9a55f71","observation_id":"abd39386-ba03-4526-b4f3-801aa664d60b","resolution":{"observed_at":"2026-08-12T04:44:08.624806Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.613280Z","title":"The unreasonable effectiveness of deep features as a perceptual metric","venue":null,"work_id":"5aeb5782-4a30-4294-81ea-591a7509f07f","year":2018},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.367008Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:4dda16d1ea7bf9653f6ba3ba3575830883131a8c8e11be89c479111798b07e1e","observation_id":"5c9198cc-907d-4f6d-a55d-78a5fe90a615","resolution":{"observed_at":"2026-08-12T04:44:08.616198Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.19680","last_updated":"2025-06-27T10:06:13Z","snapshot_observed_at":"2026-08-13T04:40:24.036354Z","submitted_at":"2024-06-28T06:40:53Z","title":"MimicMotion: High-Quality Human Motion Video Generation with Confidence-aware Pose Guidance","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.19680","snapshot_observed_at":"2026-08-12T04:44:08.370282Z","title":"Mim- icmotion: High-quality human motion video generation with confidence-aware pose guidance","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.370282Z"},"links":{"cited_paper":"/paper/2406.19680","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:f1dd72314b1e31ee9b70192c87cdb5eeb19f049bb385add834ab0fe25e30dda3","observation_id":"eb3184c5-c142-4afc-b16f-d64a2232bda8","resolution":{"observed_at":"2026-08-12T04:44:08.370282Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.605851Z","title":"Sifu: Side-view conditioned implicit function for real-world us- able clothed human reconstruction","venue":null,"work_id":"5bf29f92-8f85-4a2f-a70c-517572c61115","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.373316Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:62159c4b5e50d51c980f46e7dff36835b2bb5f961498e62e44e494d5b2ef16dd","observation_id":"d40bc59e-c6f2-4f5e-ac91-7b0ec39b838a","resolution":{"observed_at":"2026-08-12T04:44:08.608430Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.376182Z","title":"Humannerf: Efficiently gen- erated human radiance field from sparse inputs","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.376182Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:5a5da5263f3f804c7efe4417a75d93b6cbbad2a3a51b9d38bb9618b0a94d2584","observation_id":"005133d1-af98-4122-b16e-f7a3021c041e","resolution":{"observed_at":"2026-08-12T04:44:08.376182Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.593827Z","title":"Bilateral refer- ence for high-resolution dichotomous image segmentation","venue":null,"work_id":"663add35-783d-4619-b755-723089137557","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.378730Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:424709a92da01ac605541f0a52197b1a2a7aac0f60dd9e858e1878773306ece4","observation_id":"31020ae0-c16d-4561-bf52-936018e1b737","resolution":{"observed_at":"2026-08-12T04:44:08.596369Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.585711Z","title":"Gps- gaussian: Generalizable pixel-wise 3d gaussian splatting for real-time human novel view synthesis","venue":null,"work_id":"c6641693-3900-4ce8-9d38-1aba5c960bbf","year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.381391Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:793e61cdf63e7384822e70219f62493582fc5075d110137288698ae04b3b6ee3","observation_id":"5c583db8-ed99-4c31-9618-88de302b23e7","resolution":{"observed_at":"2026-08-12T04:44:08.588664Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.575447Z","title":"Pamir: Parametric model-conditioned implicit representa- tion for image-based human reconstruction","venue":null,"work_id":"40126015-e46f-465a-9c10-4fc7293720db","year":2021},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.384106Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:9ab5861733a51d993c0c46cb2c792bf7652d1d6dae02b7a172e3d68dddc16d3a","observation_id":"b0b27d47-0c46-4bc0-94ab-cc6336c93874","resolution":{"observed_at":"2026-08-12T04:44:08.579765Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.14781","last_updated":"2024-06-01T08:27:23Z","snapshot_observed_at":"2026-08-13T00:48:24.019003Z","submitted_at":"2024-03-21T18:52:58Z","title":"Champ: Controllable and Consistent Human Image Animation with 3D Parametric Guidance","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.14781","snapshot_observed_at":"2026-08-12T04:44:08.386669Z","title":"Champ: Controllable and consistent human image animation with 3d parametric guidance","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.386669Z"},"links":{"cited_paper":"/paper/2403.14781","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:5efe2f6e50d74937c74126552bd0ce4e62e05841341ea7a5593be2e132a163cf","observation_id":"193deaee-3610-48e2-bfd6-1ed699fb0235","resolution":{"observed_at":"2026-08-12T04:44:08.386669Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.08581","last_updated":"2025-02-10T20:17:32Z","snapshot_observed_at":"2026-08-13T05:25:13.326225Z","submitted_at":"2023-11-14T22:54:29Z","title":"Drivable 3D Gaussian Avatars","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.08581","snapshot_observed_at":"2026-08-12T04:44:08.389556Z","title":"Driv- able 3d gaussian avatars","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.389556Z"},"links":{"cited_paper":"/paper/2311.08581","citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:9590e81bdab0f8cad51955009f0cfb39ec1a4d428387c7746b1db0ff8d1b06e4","observation_id":"1b3328d4-3bd4-4197-85c5-c0897ce76f84","resolution":{"observed_at":"2026-08-12T04:44:08.389556Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T04:44:08.566014Z","title":"A detailed illustration of our pipeline","venue":null,"work_id":"f10b096a-c81b-4027-bf2a-2eb26fe8811b","year":null},"citing_paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-12T04:44:08.392472Z"},"links":{"citing_paper":"/paper/2412.01106"},"observation_digest":"sha256:86b073a40897b9e1fe67956123bcb058ed565244e285ad3e323aeb2caeec10ec","observation_id":"e10b32d7-0d86-4db1-a5e9-dedf2868569f","resolution":{"observed_at":"2026-08-12T04:44:08.569849Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2412.01106","last_updated":"2024-12-02T04:27:41Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-13T14:13:13.998883Z","submitted_at":"2024-12-02T04:27:41Z","title":"One Shot, One Talk: Whole-body Talking Avatar from a Single Image"},"reference_resolution":{"displayed":73,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":27,"verified_exact":0,"verified_fuzzy":46},"total_outbound_references":73},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 73 of 73 outbound references and 0 inbound Pith citation observations for arXiv:2412.01106."}