{"as_of":"2026-08-18T21:54:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:586029076af8bc5be310750aed9321a61513c2167cd888133a631fe03f932e97","coverage":[{"denominator":81,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":81,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T19:23:52.025941Z","state":"measured"},{"denominator":83,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":83,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T17:58:10.675454Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-07T15:22:34.125627Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"cited_work":{"arxiv_id":"2412.06786","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.06786","snapshot_observed_at":"2026-08-07T15:22:34.125627Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","venue":"cs.CV","work_id":"b02f4f5e-3200-4a26-999b-904aca1c3f00","year":2024},"citing_paper":{"arxiv_id":"2505.15385","last_updated":"2025-05-21T11:22:52Z","snapshot_observed_at":"2026-08-18T20:22:51.305922Z","submitted_at":"2025-05-21T11:22:52Z","title":"EVA: Expressive Virtual Avatars from Multi-view Videos","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-07T15:22:33.620089Z"},"links":{"cited_paper":"/paper/2412.06786","citing_paper":"/paper/2505.15385"},"observation_digest":"sha256:04b31d3f398a027a8dadbc6ed9092812b3999d233a71e3bd4b82b617e3cdd898","observation_id":"6ff922be-7bb7-42c5-8837-04ae2de1f921","resolution":{"observed_at":"2026-08-07T15:22:34.310454Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.06786","snapshot_observed_at":"2026-08-15T17:58:10.675454Z","title":"Retrieving se- mantics from the deep: an rag solution for gesture syn- thesis","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.19359","last_updated":"2025-07-25T15:10:15Z","snapshot_observed_at":"2026-08-18T18:31:43.022276Z","submitted_at":"2025-07-25T15:10:15Z","title":"SemGes: Semantics-aware Co-Speech Gesture Generation using Semantic Coherence and Relevance Learning","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-15T17:58:10.675454Z"},"links":{"cited_paper":"/paper/2412.06786","citing_paper":"/paper/2507.19359"},"observation_digest":"sha256:b402bf6f0b51c3e0e0105c40b542a17ed40e9f4b278c06abfd388a4fb2941373","observation_id":"b411826f-470f-4d9f-8202-386eced30e1d","resolution":{"observed_at":"2026-08-15T17:58:10.675454Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2412.06786/citation-record","integrity":"/paper/2412.06786/integrity","json":"/paper/2412.06786/citation-record.json","paper":"/paper/2412.06786"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:51.653144Z","title":"Nakano, and Louis-Philippe Morency","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.653144Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:c9f7bf0e9c23cfc0dd7680ef4ba3b628185825efb4e98453335149e5fcd46387","observation_id":"2449ae32-fe3a-410f-b805-3fae674a097c","resolution":{"observed_at":"2026-08-11T19:23:51.653144Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:53.118303Z","title":"Style-controllable speech-driven gesture synthesis using normalising flows","venue":null,"work_id":"2092c8c2-01ee-475b-adb6-35bdf8913195","year":2020},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.659282Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:207cffdc9835498bd0f6c7290878a789cb85fbdb19306a3f49abce1043797dc7","observation_id":"a65eb9fc-a926-4df2-8248-497bf3c9892a","resolution":{"observed_at":"2026-08-11T19:23:53.122870Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:53.103267Z","title":"Gesturediffuclip: Gesture diffusion model with clip latents","venue":null,"work_id":"3e21ba57-41a7-4f2a-baae-b85ea64a7595","year":2023},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.666236Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:38adbe2fcfe2ce6d9cb80dbe11767685a77f4691b46fbfd9b7a383fb8aa6020f","observation_id":"61dbc38a-68ca-420a-8a20-39cf0c957dc5","resolution":{"observed_at":"2026-08-11T19:23:53.107477Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:53.089684Z","title":"wav2vec 2.0: A framework for self-supervised learning of speech representations","venue":null,"work_id":"249b546c-d1ca-41ae-a27f-6af7d47392d9","year":2020},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.671065Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:7c51a0cd2a85bb06955252ad482de41ebf530c75a5a2a89f6fca83cb527002c8","observation_id":"3086194c-1df4-4ef8-a75f-b65b60194c05","resolution":{"observed_at":"2026-08-11T19:23:53.094387Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:53.076522Z","title":"Gnetic–using bayesian decision networks for iconic gesture generation","venue":null,"work_id":"c148b091-e3f4-4922-a4c0-ff082b7bf3db","year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.675352Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:74302aa07cd526b1981feb168df1aa9d0f4c42b91f9c37bee2d9fd179da00764","observation_id":"1c30fe27-dd94-4492-bf53-1db98f9a65d6","resolution":{"observed_at":"2026-08-11T19:23:53.080766Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:53.063349Z","title":"Bernard, Zachary B","venue":null,"work_id":"296a498e-4c0f-4c70-a19c-c1140838314a","year":2015},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.680651Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:5c36212539920a024552f2626db93b89759b854f79fe0b9e817770d382b38d89","observation_id":"0ecc64d8-fead-400c-957a-740055a6d36f","resolution":{"observed_at":"2026-08-11T19:23:53.067652Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:53.049137Z","title":"A repertoire of Ger- man recurrent gestures with pragmatic functions","venue":null,"work_id":"61b37171-b6b5-4353-83b8-a5c5e35551db","year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.685678Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:117beb07d68b36241be93e2154e94bebc274a311ea29c1c88b67e3c2769c8130","observation_id":"881571cc-e15e-4ccb-81d3-2e3ce5ed8ef5","resolution":{"observed_at":"2026-08-11T19:23:53.053622Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:53.021800Z","title":"Motion matching-the road to next gen animation","venue":null,"work_id":"5a336669-ad8e-47af-b2de-afb2cd3f0b79","year":2015},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.694564Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:1ae809def146cb4861b9a8a93b69f5ea2517305c657bcf3515c2b7fb60da98be","observation_id":"13830e35-bdfa-4d4e-99b8-1f1cc5a8181e","resolution":{"observed_at":"2026-08-11T19:23:53.026570Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:53.007728Z","title":null,"venue":null,"work_id":"9939e609-9300-4ff5-abde-bfc6fb1ddff8","year":2011},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.698965Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:e4a562f277c0aa6b288a5bdd2041bd6ae0fb4b30d72259a5a00909c8529b5a9a","observation_id":"d4f40b8c-09c9-43db-a07a-fd7a7ee2d603","resolution":{"observed_at":"2026-08-11T19:23:53.012167Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.991646Z","title":"Beat: The behavior expression animation toolkit","venue":null,"work_id":"82f6187d-884a-481f-a023-0a29ceebfe30","year":2001},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.703387Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:20276c2c3d431c0f135159974f0ba3e0c03ba03e2628c46d501d59475af146ec","observation_id":"dd347c1a-d827-4c80-9da8-85a4debe9221","resolution":{"observed_at":"2026-08-11T19:23:52.995822Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.977311Z","title":"Motion matching and the road to next- gen animation","venue":null,"work_id":"05621834-2544-4cf1-9423-0712fd0db2bc","year":2016},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.707593Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:8eb42f19d711c7c30bf51feb0c34d21d7b8a1c214b67b6edf83cdfac5dfccfd0","observation_id":"e6400655-bec4-4c26-ba0d-701722287662","resolution":{"observed_at":"2026-08-11T19:23:52.982424Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:51.711938Z","title":"Mofusion: A framework for denoising-diffusion-based motion synthesis","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.711938Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:815f507e66eb95cafda7d8463924e44812360f6d2099cb3e8634bfba1dbae43b","observation_id":"98ce2f11-92d0-4deb-8b8b-1ccb843fb229","resolution":{"observed_at":"2026-08-11T19:23:51.711938Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.953795Z","title":"Bert: Pre-training of deep bidirectional trans- formers for language understanding","venue":null,"work_id":"6ac87cac-89b9-4a86-b661-fe1ac5087c12","year":2019},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.716018Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:4dbd58896bb3be9c29454765e8eb32364daf8d07ff75190920e2dc04ea8c3c5f","observation_id":"bdfda7a9-7977-4da8-9c67-0efd15b1526d","resolution":{"observed_at":"2026-08-11T19:23:52.959281Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.940332Z","title":"Diffusion models beat gans on image synthesis","venue":null,"work_id":"282e4686-34c1-425b-ab2e-1e3452b9e901","year":2021},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.720277Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:1ba70baaf601cd572b359ebc6772567909d9915119032573b7dfa2513b68f034","observation_id":"35280a80-d952-4d8e-9b87-db13478ff697","resolution":{"observed_at":"2026-08-11T19:23:52.944804Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.927595Z","title":"Adam: A method for stochastic opti- mization","venue":null,"work_id":"f687e35b-a4d4-4b46-8b80-0cdf381ba948","year":2015},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.724566Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:90577fc465e03fef34255a56a054319619282ca2aa968554996da87323e74c38","observation_id":"21d1b94e-f45f-4c2f-8c97-3ee885d1e2b7","resolution":{"observed_at":"2026-08-11T19:23:52.931791Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.916199Z","title":"Gesture synthesis adapted to speech emphasis","venue":null,"work_id":"f1da53ef-98fa-4ade-b308-d8e90cea6df9","year":2014},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.728834Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:1e012ce7e718fcaf5fe2ccf261e061b343cdb52e582ffe69c424b9bb6ff15a38","observation_id":"87c169fa-b181-4b90-a5ab-9c2638aeac48","resolution":{"observed_at":"2026-08-11T19:23:52.920229Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.902413Z","title":"Investigating the use of recurrent motion modelling for speech gesture generation","venue":null,"work_id":"b2babde8-835a-46c5-9218-48604a602cb5","year":2018},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.733112Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:a71df66ce126100ca2014e0a150ffe96d512af8a89ec27013a271394d8050d89","observation_id":"f4e721fc-4dd4-471c-8100-cc1f3ac04f94","resolution":{"observed_at":"2026-08-11T19:23:52.907415Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.888353Z","title":"Adver- sarial gesture generation with realistic gesture phasing.Com- puters & Graphics, 89:117–130, 2020","venue":null,"work_id":"2b774149-4bc9-4148-8eb1-d354c1ba0172","year":2020},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.737859Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:25f05af8807ee03f9ce3c4207ab8bc08a91b1ef012357815bdb905aa3fcefeb8","observation_id":"7eb2cbe6-bcf9-4852-a9d0-a18ed1356941","resolution":{"observed_at":"2026-08-11T19:23:52.892702Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.872644Z","title":"Express- gesture: Expressive gesture generation from speech through database matching","venue":null,"work_id":"1fb45928-86ba-4446-b6b2-ae58d214da32","year":2021},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.741901Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:b7cbcedf1cded058375e3e407da5e0b5d4e80128d6d4570a5601733064ce78b2","observation_id":"ffa00226-c8f1-4144-9202-c7e0bf88ef3b","resolution":{"observed_at":"2026-08-11T19:23:52.877264Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.860765Z","title":"Troje, and Marc-Andr ´e Carbonneau","venue":null,"work_id":"d7e3a01f-b81c-4f11-8ebc-8f1843893d35","year":2023},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.747769Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:94fce30e1f2a7853383b98e8d50707674f6c9a6d7f690791c41bc16ae21f4560","observation_id":"3100fff6-a231-4f03-8a38-2e745f1e32f9","resolution":{"observed_at":"2026-08-11T19:23:52.864783Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.846189Z","title":"Imos: Intent-driven full-body motion synthesis for human-object interactions","venue":null,"work_id":"837d54bb-5da4-4a98-a952-d7f4507c629f","year":2023},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.752603Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:92f79a8400959a43ebf8d414589a7cb9390976c83d377dff7f6f7f35c8bfbe14","observation_id":"624fb8f2-e483-4875-a237-46ca0aab5a2f","resolution":{"observed_at":"2026-08-11T19:23:52.850954Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.832746Z","title":"Remos: 3d motion- conditioned reaction synthesis for two-person interactions","venue":null,"work_id":"30cb00fd-f9f0-4c83-aa4f-de9186a4f7d3","year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.756849Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:cd31105dce599d4e2272f1ca5b66a17bc7bebb6d7f35da1a9b5f5e98bc1feedc","observation_id":"aba544f6-62fb-49db-b88b-fc5be1b50974","resolution":{"observed_at":"2026-08-11T19:23:52.837141Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.819546Z","title":"Iterative motion editing with natural language","venue":null,"work_id":"13a3ce21-49d2-4bce-bda3-9455419b3011","year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.761881Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:2ba4dfc92d8cf05b23604361c26bbb397c6e42d1bb24e14d49d3a4bf03c73804","observation_id":"9c423cff-7b82-4eb1-9d36-a07d80575cc1","resolution":{"observed_at":"2026-08-11T19:23:52.823887Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.806014Z","title":"Learning speech-driven 3d conversational gestures from video","venue":null,"work_id":"e0db708e-dcf2-4490-966b-d983a96d9503","year":2021},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.766674Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:49854cfdf72954ac3f912ef94f6e19974c9d1fd4f8d0fe3d7d0809a32c3da1f6","observation_id":"a60b0272-e5df-42af-aed8-912e5e5ee4b4","resolution":{"observed_at":"2026-08-11T19:23:52.810757Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.791717Z","title":"A motion matching-based framework for controllable gesture synthesis from speech","venue":null,"work_id":"0d6c40ff-c5e8-491e-a5cb-f5589eb2b540","year":2022},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.771225Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:e630e370e49fc35f18d3f3050bc55ac81cd1be12d3c7511ee31ef61ec8e6f1aa","observation_id":"616cfef0-dab8-4e42-8515-8b119177db9d","resolution":{"observed_at":"2026-08-11T19:23:52.796442Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.775940Z","title":"The verbal-kinesic enactment of contrast in north american english","venue":null,"work_id":"5a9dec85-ef51-48a2-bc28-0434a17b30b6","year":2019},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.776129Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:2735ab3f64fdce5b610d798b0c9eb097181afb740abf962489fed9c83e676415","observation_id":"de698d6f-97b6-4a57-8d24-65c94e5f2ad6","resolution":{"observed_at":"2026-08-11T19:23:52.782241Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:51.781455Z","title":"Denoising diffu- sion probabilistic models","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.781455Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:48debcf1976e734604e932c493b85fb7fa0e2ce1fc281c810bd67e6c0f2f0efc","observation_id":"d0cfb653-e545-4e32-8acb-99cae99e18bd","resolution":{"observed_at":"2026-08-11T19:23:51.781455Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.753254Z","title":"Character control with neural networks and machine learning","venue":null,"work_id":"134e34d8-9aff-437d-b968-10fc5bb94159","year":2018},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.785731Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:6d3110fa87253bb13ed2016890ec8ac8070c0c472229afc4b9cddc600a3607d8","observation_id":"73e495f4-3450-48c9-8d7c-96f361299fb0","resolution":{"observed_at":"2026-08-11T19:23:52.757879Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.740729Z","title":"Learned motion matching","venue":null,"work_id":"09ebedb6-9b91-4aac-8dfa-0ecd4e8412e5","year":2020},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.790303Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:f7ccfea0233a0efd7b476c33cb3cb971f208442240b6260c30bd7d21033aa618","observation_id":"c06dd7d6-fe20-43af-a887-e8eef4f120c3","resolution":{"observed_at":"2026-08-11T19:23:52.744521Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.727848Z","title":"Como: Controllable motion generation through language guided pose code edit- ing","venue":null,"work_id":"487edbba-2d35-42d4-9a55-068bcc5a747c","year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.794688Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:5fdb951db7f5d6901840d3cfa224f18e6d482363a66408abef4d13c27f08babf","observation_id":"d8c12727-740a-4fdf-943c-bd57f09f0e11","resolution":{"observed_at":"2026-08-11T19:23:52.732249Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.715718Z","title":"The raised index finger gesture in Hebrew mul- timodal interaction","venue":null,"work_id":"0efe9260-f3a4-4985-b0f2-a07ad12b1198","year":2022},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.798796Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:1e15fe14a2af5bdd72e308a7ab9184f1c8a1571e92f8821a2b96d922d7aab654","observation_id":"57507b9b-11ca-44eb-9f8f-6fa3bdce9e60","resolution":{"observed_at":"2026-08-11T19:23:52.720033Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.702692Z","title":"Gesture: Visible action as utterance","venue":null,"work_id":"c2154cc3-dcf6-4726-a8c8-3b20e4f21910","year":2004},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.802810Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:682bffdb73dcfec522fcca50fa1f95655a0ebcc13faf9f5453f88ab8cd862b1a","observation_id":"1b9a1718-0323-4fb0-911a-7b44695c4668","resolution":{"observed_at":"2026-08-11T19:23:52.706944Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.689267Z","title":"Gesture generation by imitation: From hu- man behavior to computer character animation","venue":null,"work_id":"748e76de-cdbb-4746-8293-9656a429181b","year":2005},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.807735Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:0ac5e94e7bbe7a7107827e2576fe601cf6efaaf0f7e8dbf8caee5da0a4963f56","observation_id":"f71cb779-d0f0-4a6a-ab18-0996d69cadad","resolution":{"observed_at":"2026-08-11T19:23:52.693507Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.675763Z","title":"discopy: A neural system for shallow dis- course parsing","venue":null,"work_id":"8226e892-abd9-4e0e-99c3-bbc2a24ec0e3","year":2021},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.812111Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:1b9b8e0ddd88e012b0d2663d8d22427626b15b423fd8cbd87ec0c6a939d5be3f","observation_id":"26ce5775-e48f-4258-aa73-7b4dec2e2f05","resolution":{"observed_at":"2026-08-11T19:23:52.679974Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.661302Z","title":"Analyzing input and output representations for speech-driven gesture genera- tion","venue":null,"work_id":"59b05546-3e38-4242-92af-257dce6ac1c1","year":2019},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.816022Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:27bf33b7e308d3ee6d88f889bc7d688f6378738885647dd351064dce536ef92c","observation_id":"214820b3-8d73-444f-beb2-5ec1c218af8a","resolution":{"observed_at":"2026-08-11T19:23:52.666712Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.643669Z","title":"Speech2Properties2Gestures: Gesture-property predic- tion as a tool for generating representational gestures from speech","venue":null,"work_id":"2c7a463d-3b07-46b7-b0aa-0e15f9f4be47","year":2021},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.819902Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:58d7ca1fc3af3377ec89bbd1e21acad2da7144a1b41e956a6094f597e46ea85b","observation_id":"89a2d9aa-6bae-4e47-b4e3-862ee65cf42f","resolution":{"observed_at":"2026-08-11T19:23:52.648299Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.630092Z","title":"Tracking discourse topics in co-speech gesture","venue":null,"work_id":"c47abd7e-5065-49c8-b78e-73f5b7572aca","year":2021},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.824391Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:f186b78a5dba22ef0fb8b29e9e276e7d6b9887842b696cd3486dcb12ba686611","observation_id":"7e324f01-d5c5-4ce7-a60b-6bbabf8dff9f","resolution":{"observed_at":"2026-08-11T19:23:52.634591Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.615988Z","title":"Ross, and Angjoo Kanazawa","venue":null,"work_id":"fa3c5585-bfc3-467f-afce-031b8fb8c351","year":2021},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.830075Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:11b22a4707259e764cb041c0dc4982040356d18f9fad8e408d4bd85bc1fea4ac","observation_id":"c0f6b0c9-bd61-4a66-bdd8-2983529fc13a","resolution":{"observed_at":"2026-08-11T19:23:52.621075Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.595402Z","title":"Beat: A large-scale semantic and emotional multi-modal dataset for conversational gestures synthesis","venue":null,"work_id":"5a89f38a-263b-4a11-8a24-79853dddfea2","year":2022},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.834037Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:7fcbc6370f14ae65093822fe4ea428c65a6c8ff972c2a3ff2c20dd7f8b5d1be8","observation_id":"20025e7e-1417-4512-8f05-06b146bc90f8","resolution":{"observed_at":"2026-08-11T19:23:52.600436Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.581726Z","title":"Emage: Towards unified holis- tic co-speech gesture generation via expressive masked audio gesture modeling","venue":null,"work_id":"8ef132fc-b6bb-4b47-805d-0466291ab6ff","year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.838841Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:15414652e546e466c0e65aae6e593f8bd0fb4b59f1cf6a72cb97ed6b2597b058","observation_id":"e06f1bb9-f939-44e2-b760-0966764f2b90","resolution":{"observed_at":"2026-08-11T19:23:52.586346Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:51.843921Z","title":"Repaint: Inpainting using denoising diffusion probabilistic models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.843921Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:6fe8863a5df6bc0a7b453690467ba514f460fa39d0405c322c27f14885c1170a","observation_id":"e6d5e1af-d904-49fb-8f18-285981359513","resolution":{"observed_at":"2026-08-11T19:23:51.843921Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.560655Z","title":"Rhetorical Struc- ture Theory: Toward a functional theory of text organization","venue":null,"work_id":"a47e0eeb-c5dd-4bed-a0e5-78ff3c074c3f","year":1988},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.849904Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:ceca7a6cba99b443d39b6ce343c8c8529324f84ada8914d754c9c10fd8e14fa8","observation_id":"89dac410-a0fa-4a1a-ba4d-02b658955c6c","resolution":{"observed_at":"2026-08-11T19:23:52.565108Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.548941Z","title":"Hand and Mind: What Gestures Reveal About Thought","venue":null,"work_id":"804be6aa-a4d7-4dc0-9eee-eb1983be58e2","year":1992},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.854279Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:ec8bc3dc8b3ad00aa7fd20c80143b56a1cc2db1940b8c82b040c9eddbf3e796d","observation_id":"5f3704e5-2257-4311-803f-9d743f9e7755","resolution":{"observed_at":"2026-08-11T19:23:52.552804Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.536243Z","title":"Gesture and Thought","venue":null,"work_id":"ace6cc4d-55e7-4dc2-ab4e-aa6111d320fc","year":2005},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.858394Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:699b9072bbfedad2693126b2d9e0b2cf83ab2256309f8313e976fa390d9a2f46","observation_id":"abfaf61d-8982-42b1-ae6b-48134e4e7ef7","resolution":{"observed_at":"2026-08-11T19:23:52.540554Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.519750Z","title":"Gesture in discourse","venue":null,"work_id":"92669461-d759-4198-a0d8-d09476c4a5e9","year":2014},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.863180Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:90c4f4ca6c1ed56b10cd676779910f2116a59ba48e35f652795a41b415a8e4a1","observation_id":"6ea67c23-da58-453b-8bae-e2941e3d5a0a","resolution":{"observed_at":"2026-08-11T19:23:52.523781Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.506777Z","title":"Null-text inversion for editing real images using guided diffusion models","venue":null,"work_id":"d9b0cb1f-cb27-4c8c-9a53-d48556caafb6","year":2023},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.867225Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:2d6e205ca787e52755ac54952f8dc7350c21fd4b2b1e055784901f55ac61e267","observation_id":"5fa30d6c-370c-4d0f-98d8-2d4a46216df2","resolution":{"observed_at":"2026-08-11T19:23:52.510923Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.493221Z","title":"Convofusion: Multi-modal conversational diffu- sion for co-speech gesture synthesis","venue":null,"work_id":"178f2bcf-d253-4dcb-a1bc-e0929b0a477d","year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.871514Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:4e72078535fdff80b2dd932dff497de63d67f3c73e1172a346eb0ab5e963cd08","observation_id":"5eb645d4-ad03-41ce-993a-f22092326256","resolution":{"observed_at":"2026-08-11T19:23:52.498087Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.06327","last_updated":"2024-10-08T20:05:08Z","snapshot_observed_at":"2026-08-16T13:11:21.924791Z","submitted_at":"2024-10-08T20:05:08Z","title":"Towards a GENEA Leaderboard -- an Extended, Living Benchmark for Evaluating and Advancing Conversational Motion Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.06327","snapshot_observed_at":"2026-08-11T19:23:51.875514Z","title":"Towards a genea leaderboard–an extended, living benchmark for evaluating and advancing conversational mo- tion synthesis","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.875514Z"},"links":{"cited_paper":"/paper/2410.06327","citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:ff92f68d756147e97ab981b49309dd54dd814c631295a479f6cad0472d6dc3ef","observation_id":"54508a8b-1814-4340-9d57-3f38bb210f57","resolution":{"observed_at":"2026-08-11T19:23:51.875514Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.480673Z","title":"Hand gesture synthesis for conversational characters","venue":null,"work_id":"652b2dc9-ce92-4bb3-869a-5b5eb6d7e7ec","year":2016},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.879939Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:c401f6e92b81af18466c80e3b6f10769cd8f905516f387c4749d433652a22e10","observation_id":"75c2fab8-f1a1-4ab9-9124-703250c94f5c","resolution":{"observed_at":"2026-08-11T19:23:52.485606Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.467557Z","title":"Gesture modeling and animation based on a proba- bilistic re-creation of speaker style","venue":null,"work_id":"f363ba22-fdfd-416b-83fb-8535f4e9bbf5","year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.884839Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:cdb2bd42bca8b28069a19daaea2c6d25c09fa24c080e22a711e3ea53276278ee","observation_id":"ab34a8a9-07f2-4300-a269-5a288c8fc4c5","resolution":{"observed_at":"2026-08-11T19:23:52.471626Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.454337Z","title":"From audio to photoreal embodiment: Synthesizing humans in conversations","venue":null,"work_id":"cb5561c5-85fd-4ede-8215-da5eaf31243c","year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.889177Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:d2525d521553c2f0a8ee4d7af9a9296b166f65040ee08fca6458f1ee1d0d35b5","observation_id":"1773cbb9-536d-472b-8c31-b2ec9ecd981a","resolution":{"observed_at":"2026-08-11T19:23:52.458395Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.443157Z","title":"A comprehensive re- view of data-driven co-speech gesture generation","venue":null,"work_id":"b05a50a9-378d-442d-a074-68f195e0e51d","year":2023},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.893054Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:dcb85e5a7de1da89c51aaa66cedf59ffac392c04d97648ae585f616268a089d4","observation_id":"8e6cac11-b500-44ce-ac85-107ba7981b23","resolution":{"observed_at":"2026-08-11T19:23:52.447101Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-17T09:58:46.058102Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-11T19:23:51.897125Z","title":"Gpt-4 technical report, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.897125Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:de5b1314ed39e54fdab4d3bb31e997ebe091f5a5ebea424c9798b8f115f6baff","observation_id":"df10fd17-9c5d-43ef-8cb1-e163771df815","resolution":{"observed_at":"2026-08-11T19:23:51.897125Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.431328Z","title":"Visualizing prosodic structure: Manual gestures as highlighters of prosodic heads and edges in english academic discourses","venue":null,"work_id":"9dbf036a-91f8-4dee-924a-328a33328a28","year":2023},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.901591Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:0ed814f1818c7d693ea8c57ffe66b240dd7421228eb8721bbdc938112569bc5b","observation_id":"99970d84-8ddb-4563-afa7-51ff9cfd71ba","resolution":{"observed_at":"2026-08-11T19:23:52.435437Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.419257Z","title":"High-resolution image syn- thesis with latent diffusion models","venue":null,"work_id":"9d2b4fa1-e4b6-4219-af00-b86e519567d0","year":2021},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.906340Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:11c8edbddf7aa9162ebbc92ab87ed258ce10bd68c173ffbe0867dfd5043ffa2f","observation_id":"95ebb77d-b5d1-4a87-b24b-e40167f5e03d","resolution":{"observed_at":"2026-08-11T19:23:52.423147Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.406479Z","title":"Toward a taxonomy of coherence relations","venue":null,"work_id":"1a0cafb2-4c1e-487a-ba4a-a71ebeb8b319","year":1992},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.910843Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:c9bdc3a7ebd54cd8840850052f7f83a2967b9800ffa89686da21b0d53821f1a7","observation_id":"beb57ae2-a707-49d3-9c7b-dd4ac88fb407","resolution":{"observed_at":"2026-08-11T19:23:52.410590Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.393002Z","title":"Denois- ing diffusion implicit models","venue":null,"work_id":"6cf049b6-ddd7-4650-a354-60f271046dd0","year":2021},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.915563Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:3d5b6a1b974ce8254f042ff2ddd42ee841d4e79dda51122d5c2265b9a28e5e68","observation_id":"b2e5544c-3c18-4933-9cbc-d90637b7a472","resolution":{"observed_at":"2026-08-11T19:23:52.397268Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.379891Z","title":"Categorical codebook matching for embodied character controllers","venue":null,"work_id":"27ebfe4f-3fcd-42fa-9e71-ac597ab33e9b","year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.920770Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:07905cd2489622f58376503b7813f7f30b1a7465671b4204b18fe6d8c361e6a0","observation_id":"69d6083c-fefe-41ac-bd17-c042fd13184f","resolution":{"observed_at":"2026-08-11T19:23:52.384047Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.367654Z","title":"Hierarchical representation and estimation of prosody using continuous wavelet transform","venue":null,"work_id":"e491a30c-c248-4640-b54c-587aec3ac375","year":2017},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.925236Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:302868133e06645a58178b2ef393dd7645c601c48d4e0051aa41e93347438616","observation_id":"3827e948-c410-4898-9051-1cc239e5bc94","resolution":{"observed_at":"2026-08-11T19:23:52.371830Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.354700Z","title":"The perception of prosodic prominence","venue":null,"work_id":"937da10b-2c0a-4b96-8020-8caa8dddff65","year":2000},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.930692Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:4379e75c4cff519801749aaaa18913aeb2574bb5a5295db70e369e4bfbe76166","observation_id":"9893f28b-47a9-480c-a324-38cf5412f7fb","resolution":{"observed_at":"2026-08-11T19:23:52.359326Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.343335Z","title":"Smartbody: Behavior realization for em- bodied conversational agents","venue":null,"work_id":"ec2d3cce-ea05-47df-b177-bfc6fd9e7cae","year":2008},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.934883Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:5bb35abd2d14a25df340a1e30ee7973447b27cd4201852be1faf8f2f3f265dbc","observation_id":"688c8211-3955-4c23-a9ad-3743f7bf1b32","resolution":{"observed_at":"2026-08-11T19:23:52.347324Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.331019Z","title":"Experimental and theoretical advances in prosody: A review","venue":null,"work_id":"adffa6aa-1c5d-4d7e-a41b-9c7214bb4d5b","year":2010},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.939655Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:bb61f173908e790b121ef1a7e26ba1dfe1d34f879ad4b2c22a05aa4832a94cea","observation_id":"4ffff9e3-60fb-4bfa-837a-c30f402d8f2b","resolution":{"observed_at":"2026-08-11T19:23:52.335666Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.318562Z","title":"A re- view of evaluation practices of gesture generation in embod- ied conversational agents","venue":null,"work_id":"40b13415-80ea-4c14-962b-d5bf231f906a","year":2022},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.943795Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:b2ec53dd71afea938ee599ef5f5fafbbbe32f84bd18a260390231604c9866120","observation_id":"55819386-9f76-4dda-b384-d75e1ac29dc9","resolution":{"observed_at":"2026-08-11T19:23:52.322637Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.305635Z","title":"Generating holistic 3D human motion from speech","venue":null,"work_id":"26c98731-7275-4c6f-bc09-c9cc7f3c85e4","year":2023},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.947997Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:34132be74c1cc4b8772444a8f05df89bdad4fd78034ea18d000b659540bdb9d5","observation_id":"3e2dd088-c573-4aef-9214-1e4ed17eaab6","resolution":{"observed_at":"2026-08-11T19:23:52.310666Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.291449Z","title":"Robots learn social skills: End-to-end learning of co-speech gesture generation for humanoid robots","venue":null,"work_id":"45b7b6a8-a92d-45d6-941c-1d600f5500e0","year":2019},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.952661Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:369a09431bacb742397fac4b6b3bb67131455ca51e88c8d9cd9d1110c14d3191","observation_id":"89e05fbb-1a99-474c-b6d7-c4dd807a816c","resolution":{"observed_at":"2026-08-11T19:23:52.296150Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.278659Z","title":"Speech ges- ture generation from the trimodal context of text, audio, and speaker identity","venue":null,"work_id":"2af18f52-e7a3-43f6-95cb-0bfd8e62043f","year":2020},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.958374Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:b96ec2be911fe94a37dcb8980854b1b19c4cf538d6d38671f57c1995d9d729d5","observation_id":"449f99c4-b792-4bf4-b7f9-1a9198f0bd15","resolution":{"observed_at":"2026-08-11T19:23:52.282879Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.265798Z","title":"The genea challenge 2022: A large evaluation of data- driven co-speech gesture generation","venue":null,"work_id":"c231b221-30e8-4461-9ce2-e3155233730e","year":2022},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.963006Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:cc73c825dcfb7a166ad4fb28d3a1ad2c5ace066209f2df0eca5fb34be282cbe0","observation_id":"ece68872-dc2c-4d88-a9a6-2af54ae4f15f","resolution":{"observed_at":"2026-08-11T19:23:52.270143Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.252098Z","title":"Re- modiffuse: Retrieval-augmented motion diffusion model","venue":null,"work_id":"712ae947-9212-439f-9c74-152850f21c05","year":2023},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.968772Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:42ce5e319839de4fd207aac09f86079929d8785a1fc920f34f97bba014ded527","observation_id":"4f7527c8-c244-4be8-8e08-1aa78030c131","resolution":{"observed_at":"2026-08-11T19:23:52.256906Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.239607Z","title":"Motiondif- fuse: Text-driven human motion generation with diffusion model","venue":null,"work_id":"45fc078c-3b5e-48e2-ae22-b3ad69821200","year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.973947Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:e12b87af2365bcb3b95d849c5085f532fb66f93beaa680b57738fc0403b6722e","observation_id":"c61cb747-c7b5-4705-9cd3-f7d41472517b","resolution":{"observed_at":"2026-08-11T19:23:52.243406Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.223734Z","title":"Roam: Robust and object-aware motion genera- tion using neural pose descriptors","venue":null,"work_id":"9468d5df-7b13-47bf-b908-7b35b60c9571","year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.979307Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:4f90ca5aa90d0624f015b2c68bbefc4297ede9f7d1f9f76c098bfbd5dd2efcef","observation_id":"aeb575bc-3b55-4d0c-acde-9c498589f49c","resolution":{"observed_at":"2026-08-11T19:23:52.231186Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.210769Z","title":"Semantic gestic- ulator: Semantics-aware co-speech gesture synthesis","venue":null,"work_id":"f4123c5b-b8c1-4c0a-a1d3-4a915baf16a7","year":2024},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.986469Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:c3d516f5a29865efab6340c66a6550ee3529a1b10a56dceea9bfbd2e8b3f9923","observation_id":"cc77a17a-c39e-43a7-ae3c-dab0b49905c1","resolution":{"observed_at":"2026-08-11T19:23:52.215709Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.198495Z","title":"Diffugesture: Generating human gesture from two-person dialogue with diffusion models","venue":null,"work_id":"9e1a3b83-4e21-4d3b-8beb-43b80ab02971","year":2023},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.990984Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:5b5be90ae2454df87239135302f31a0b29733881cb91e78cd5f652539e8efa62","observation_id":"23a16ec0-4457-4261-820b-1a84a7c5a73a","resolution":{"observed_at":"2026-08-11T19:23:52.202920Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.18223","last_updated":"2026-03-18T05:34:39Z","snapshot_observed_at":"2026-08-14T10:40:26.323157Z","submitted_at":"2023-03-31T17:28:46Z","title":"A Survey of Large Language Models","version":19},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.18223","snapshot_observed_at":"2026-08-11T19:23:51.994974Z","title":"A survey of large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.994974Z"},"links":{"cited_paper":"/paper/2303.18223","citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:31ffad6f0d3f690cdeb7899d515d47416fc27f3e5fa93cc24ba5655d807b6246","observation_id":"8f0d70f0-db90-4a6b-bae7-4aaefa8181e7","resolution":{"observed_at":"2026-08-11T19:23:51.994974Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.183885Z","title":"Livelyspeaker: Towards semantic-aware co-speech gesture generation","venue":null,"work_id":"5053ef7f-cafa-4474-ae8c-3ea0b3eaef07","year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.999646Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:c7fe7ac41fe128e0e4c06b0e3686e154a30767aaeade96b63f314ce4417ac48e","observation_id":"64fa6c12-912a-48f3-8612-9de833f931f6","resolution":{"observed_at":"2026-08-11T19:23:52.189697Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.170148Z","title":"On the continuity of rotation representations in neural networks","venue":null,"work_id":"664d1ac1-06f4-4a3d-a49d-af9a254d1dee","year":2019},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:52.003556Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:19d133e44d4540cd3f21378051c2025ba747c7e706a60696888e6b6ff79227c7","observation_id":"d24bd7a5-e174-4db2-9359-5419493e21fa","resolution":{"observed_at":"2026-08-11T19:23:52.174626Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.154995Z","title":"This creates a scarcity of good exemplars which can be used during the database matching steps","venue":null,"work_id":"9bc0313d-b974-4a8f-833c-e9298d0aa2b1","year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:52.007423Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:1a534a1543a1cc9ec39b8a2ec018a59ce50878f45e0b8efb4c5e077bf2f2e8e3","observation_id":"ce66ac91-e075-4f19-97fa-eb9a7c449383","resolution":{"observed_at":"2026-08-11T19:23:52.160283Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.138720Z","title":"Which of the two gestures look natural?","venue":null,"work_id":"777e8738-f217-4335-901c-3401b7c51099","year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:52.013218Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:00cd4eaa2df308bb405c7f4bc708670e7809a9347453662c65c4fdc7c3847f1f","observation_id":"0f127f02-2c9d-4492-9c2b-84ecf75c7641","resolution":{"observed_at":"2026-08-11T19:23:52.144760Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.123730Z","title":"We employ the Frechet Inception Distance (FID) metric inspired by Yoon et al","venue":null,"work_id":"5321eb64-636c-4812-a841-f9a156a9d772","year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:52.017531Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:9fee771ce66fc38b916bf63c1b4c534773f7d5cda6d14b530c47b0fe0992572e","observation_id":"a9ab2d3e-475d-49c6-a5ec-d8b5de6c46c4","resolution":{"observed_at":"2026-08-11T19:23:52.129073Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.110735Z","title":null,"venue":null,"work_id":"e6808fdf-5180-4122-af2c-0fe09ea3bfc0","year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:52.021488Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:d4e4cdf13e39d9e78a66175b10a626140a12c791c8b58abad987d7ce04f44e2a","observation_id":"33a4b9f7-ee97-40f9-ac40-c18b42334e2f","resolution":{"observed_at":"2026-08-11T19:23:52.115177Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:52.096168Z","title":"scaled linear","venue":null,"work_id":"37d6b47f-11b1-42e3-8e43-93dde34015ce","year":null},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:52.025941Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:f567b3d8bb5a382e4176a59c95129689b81bc5e697bfbbbde5c9f0c6ee947069","observation_id":"551d37dd-111f-4901-a67e-ba760b59613b","resolution":{"observed_at":"2026-08-11T19:23:52.101504Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T19:23:53.035527Z","title":null,"venue":null,"work_id":"c5ecfbcf-193e-4067-ac7d-d89b0c2d1b31","year":2014},"citing_paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis","version":3},"reference_index":1591,"source":"pdf_text","source_observed_at":"2026-08-11T19:23:51.690085Z"},"links":{"citing_paper":"/paper/2412.06786"},"observation_digest":"sha256:7560425100c777e898f6025e7be98aac9729446adf91bdea5b4402e7985f4b3e","observation_id":"3cb99638-b88c-4069-b928-bb9b06c8d627","resolution":{"observed_at":"2026-08-11T19:23:53.039780Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2412.06786","last_updated":"2025-04-04T07:48:19Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-18T20:23:21.229920Z","submitted_at":"2024-12-09T18:59:46Z","title":"Retrieving Semantics from the Deep: an RAG Solution for Gesture Synthesis"},"reference_resolution":{"displayed":81,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":10,"verified_exact":0,"verified_fuzzy":70},"total_outbound_references":81},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 81 of 81 outbound references and 2 inbound Pith citation observations for arXiv:2412.06786."}