{"as_of":"2026-08-18T04:18:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b611a94e47bb1312d6d203ea29b7b0a788052c6b5fac58cb39a4b06896ffd4a6","coverage":[{"denominator":63,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":63,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T23:27:45.288936Z","state":"measured"},{"denominator":64,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":64,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-13T20:16:35.411267Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-13T20:18:13.202646Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"cited_work":{"arxiv_id":"2506.17912","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.17912","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Planmogpt: Flow-enhanced progressive planning for text to motion synthesis","venue":null,"work_id":"2c784af1-8cb6-4ea0-bd1a-9e8c8df89242","year":2025},"citing_paper":{"arxiv_id":"2604.02908","last_updated":"2026-04-19T15:13:54Z","snapshot_observed_at":"2026-08-14T09:38:21.018074Z","submitted_at":"2026-04-03T09:26:28Z","title":"SentiAvatar: Towards Expressive and Interactive Digital Humans","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-13T20:16:35.411267Z"},"links":{"cited_paper":"/paper/2506.17912","citing_paper":"/paper/2604.02908"},"observation_digest":"sha256:d42809d96fb6d37f8fbde46104a9fbc25ea9aca52ab97864bed40120d28bb9c0","observation_id":"bbbe2e75-3a77-4fa9-bc92-ab3620eeb88b","resolution":{"observed_at":"2026-05-13T20:18:13.204278Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.17912/citation-record","integrity":"/paper/2506.17912/integrity","json":"/paper/2506.17912/citation-record.json","paper":"/paper/2506.17912"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2312.08895","last_updated":"2023-12-14T12:57:35Z","snapshot_observed_at":"2026-08-18T00:23:16.847280Z","submitted_at":"2023-12-14T12:57:35Z","title":"Motion Flow Matching for Human Motion Synthesis and Editing","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.08895","snapshot_observed_at":"2026-08-06T23:27:40.790842Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:40.790842Z"},"links":{"cited_paper":"/paper/2312.08895","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:2a8fcd05bb46e11b77d6405f23d2d543a64de8213440e1167cb5a6a6cc4dbc90","observation_id":"90a763a1-4952-46ba-bd87-8a8c9bbd7aad","resolution":{"observed_at":"2026-08-06T23:27:40.790842Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-17T09:58:46.058102Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-06T23:27:40.884564Z","title":"Gpt-4 technical report.arXiv preprint arXiv:2303.08774, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:40.884564Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:3d505cf3bfd39bd7600d59dd1fc40a1fee3479484cd0cdfb12e8b27aa594a203","observation_id":"57669eca-3d83-4974-9a06-9ffab7a0cc2d","resolution":{"observed_at":"2026-08-06T23:27:40.884564Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:40.953922Z","title":"Flamingo: a visual language model for few-shot learning.Advances in neural information processing systems, 35:23716–23736, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:40.953922Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:ed3eba82430d7f10d0734a4ad2c6d0a375efcbb4b6d1e4a4914308c9c37c1167","observation_id":"1f6e173e-f83a-4c82-8e5d-4701e13d8fb2","resolution":{"observed_at":"2026-08-06T23:27:40.953922Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:46.002361Z","title":"Gesturediffuclip: Gesture diffusion model with clip latents.ACM Transactions on Graphics (TOG), 42(4):1–18, 2023","venue":null,"work_id":"13a8e676-bc89-466b-8101-f5b768c22324","year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.030977Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:ac6da694f1cf9adaab041905c54b7c55ddc189c6e98339d81eb18f3591c6707e","observation_id":"816bef85-a25e-4377-abc5-bd7efa5c1b7a","resolution":{"observed_at":"2026-08-06T23:27:46.006530Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:41.111558Z","title":"Language models are few-shot learners.Advances in neural information processing systems, 33:1877–1901, 2020","venue":null,"work_id":null,"year":1901},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.111558Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:3ada739e8a3b687f4ca9f3b43f79e414733fbec1fae722aee3aa93c550f2cfb1","observation_id":"7a526be9-24ed-4d37-92df-ea0fc9ce24b0","resolution":{"observed_at":"2026-08-06T23:27:41.111558Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.981285Z","title":"Long- term human motion prediction with scene context","venue":null,"work_id":"068d46f4-8a99-46c6-88d0-0a8f76dab9d4","year":2020},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.177602Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:2e3534daea789c090e413cd6302ff9f2867ac0ef052fb1e2fc82a823b8f34174","observation_id":"78b91fcb-316a-40c2-a690-9e5016e15408","resolution":{"observed_at":"2026-08-06T23:27:45.985661Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.968390Z","title":"Cmu graphics lab motion capture database, 2003","venue":null,"work_id":"5016f1f5-c4ff-4c61-afb4-33bd9b83102e","year":2003},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.269465Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:93d4181cad85c78b872c0c504a6f8b273d4f55c93da4555ee6a2f35aa8e93abe","observation_id":"3ce38eb4-dcee-4427-8e8c-c372f81d34f5","resolution":{"observed_at":"2026-08-06T23:27:45.973174Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1709.06158","last_updated":"2017-09-18T20:34:48Z","snapshot_observed_at":"2026-08-14T20:32:01.406964Z","submitted_at":"2017-09-18T20:34:48Z","title":"Matterport3D: Learning from RGB-D Data in Indoor Environments","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1709.06158","snapshot_observed_at":"2026-08-06T23:27:41.329444Z","title":"Matterport3d: Learning from rgb-d data in indoor environments.arXiv preprint arXiv:1709.06158, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.329444Z"},"links":{"cited_paper":"/paper/1709.06158","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:8c7a34879d4abaf6d6d8242b345fc3b9d72e42d4f7a05eb88f8db268f1373a29","observation_id":"ecd3cc6c-2636-43ba-8afd-fe32c3b7ee89","resolution":{"observed_at":"2026-08-06T23:27:41.329444Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:41.368660Z","title":"Executing your commands via motion diffusion in latent space","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.368660Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:70fd5cb4c969a4cc9a9cbff4f314984c10a5b41d2091f4f3ebdd30be32444a20","observation_id":"7696776e-6263-4799-a737-657b688c2ee8","resolution":{"observed_at":"2026-08-06T23:27:41.368660Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.947719Z","title":"Generative adversarial graph convolutional networks for human action synthesis","venue":null,"work_id":"7fe4808b-ea90-4069-962a-f35043d529de","year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.450727Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:c8534fef956d15e304d00e0af21957ae6a15bd7eef3a074f2186ce4a1f34a366","observation_id":"5de944e7-8c27-4f3c-a07f-2b2f8248296f","resolution":{"observed_at":"2026-08-06T23:27:45.951577Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.934926Z","title":"Learning individual styles of conversational gesture","venue":null,"work_id":"6500d3c3-2626-4886-a1d7-3b04e29c9334","year":2019},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.496393Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:aa168168020555a81ec6d4e55521369b3ff12746e680a0e8b00ee63819d8f0de","observation_id":"b013ad27-0cd9-4466-88a4-132ced0e390c","resolution":{"observed_at":"2026-08-06T23:27:45.939695Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:41.563826Z","title":"Generative adversarial nets.Advances in neural information processing systems, 27, 2014","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.563826Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:a7bb43e819eb0d6e375c5d42036e5e13bdfae2d4cff5dde77c0004e23998cc54","observation_id":"41102238-5526-4254-9e40-284c4da4ad7d","resolution":{"observed_at":"2026-08-06T23:27:41.563826Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:41.655390Z","title":"Momask: Generative masked modeling of 3d human motions","venue":null,"work_id":null,"year":1900},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.655390Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:467945876065db34388a240be03be5ec4225fc4a746233694b2c01dc01dd306c","observation_id":"e96f607d-db40-4bd0-8e15-07a904e7c950","resolution":{"observed_at":"2026-08-06T23:27:41.655390Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.896998Z","title":"Generating diverse and natural 3d human motions from text","venue":null,"work_id":"515b4386-ed6a-4870-8107-9f44e9c06a6c","year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.746624Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:8599fbfa8541f80c869cb9913514e80ac930317b8c7e4520e0112f41cd2cf70f","observation_id":"f6e7729f-74e0-4049-ba57-4573d64075bc","resolution":{"observed_at":"2026-08-06T23:27:45.902606Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:41.807511Z","title":"Action2motion: Conditioned generation of 3d human motions","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.807511Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:f740d912268834097821b5be34ac8df4b9e7ef241c959487b2d1bbf41d3da01e","observation_id":"053a7caf-3f12-40ed-b05e-44155b30ce1e","resolution":{"observed_at":"2026-08-06T23:27:41.807511Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.875086Z","title":"Stochastic scene-aware motion prediction","venue":null,"work_id":"6a8479f4-9594-4178-bcd6-316151dcafbe","year":2021},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:41.897057Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:30978772a8ebd96604167109dc20554bfdda3250755aebaa2d028444da416ed8","observation_id":"9d8e9e12-bc80-40e8-bf60-71324fefe8a8","resolution":{"observed_at":"2026-08-06T23:27:45.880994Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:42.014387Z","title":"Denoising diffusion probabilistic models.Advances in neural information processing systems, 33:6840–6851, 2020","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.014387Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:764be16e78e9c38146f1ea1ecb3577b231caadcbc52ff34ed3f63b9ebf98a524","observation_id":"2d05f54c-11dc-4994-bee4-fa4265f943a8","resolution":{"observed_at":"2026-08-06T23:27:42.014387Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.847618Z","title":"Avatarclip: zero-shot text-driven generation and animation of 3d avatars.ACM Transactions on Graphics (TOG), 41(4):1–19, 2022","venue":null,"work_id":"e1eeec82-0a72-4033-a2b7-35c5d51ca36a","year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.112072Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:5a6034d6397fcb28d6bf67ac433d66fee7486aa2bfebf2192ecb6b5228710353","observation_id":"5d02a334-0f00-49fd-b8eb-6c78d08a11b6","resolution":{"observed_at":"2026-08-06T23:27:45.852164Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:42.194414Z","title":"Motiongpt: Human motion as a foreign language.Advances in Neural Information Processing Systems, 36:20067–20079, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.194414Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:c08cee4ada6d2613fb8fc52ae30a46f01b1477b24b57356553c902a0e3c49bb2","observation_id":"2c1ea0fd-ffd0-4bb6-a55e-352daccc1e77","resolution":{"observed_at":"2026-08-06T23:27:42.194414Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1312.6114","last_updated":"2022-12-10T21:04:00Z","snapshot_observed_at":"2026-08-14T23:50:45.029465Z","submitted_at":"2013-12-20T20:58:10Z","title":"Auto-Encoding Variational Bayes","version":11},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1312.6114","snapshot_observed_at":"2026-08-06T23:27:42.307504Z","title":"Auto-encoding variational bayes.arXiv preprint arXiv:1312.6114, 2013","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.307504Z"},"links":{"cited_paper":"/paper/1312.6114","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:24e43bdbd63c5a70c56a25194193be3f05b72117966a1879cb6c1231609b1729","observation_id":"654560e6-27ae-4bfc-af7f-be440fa7ce2f","resolution":{"observed_at":"2026-08-06T23:27:42.307504Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.828477Z","title":"Generating images with multimodal language models.Advances in Neural Information Processing Systems, 36:21487–21506, 2023","venue":null,"work_id":"98c7a00c-8aac-41f2-8f33-1ba15374a50d","year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.414811Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:ad8b4f98ae341c5f220edf7f7e2c963b22ea988a6ad792bf5dd299d0030d6071","observation_id":"32780bfb-008c-42f1-a852-2bd3bf9cf993","resolution":{"observed_at":"2026-08-06T23:27:45.833196Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:42.495149Z","title":"Large language models are zero-shot reasoners.Advances in neural information processing systems, 35:22199–22213, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.495149Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:55d980436258a265fc3a33ae4f6eb7ffc552103446eb58eb13ab73a8cba50338","observation_id":"ff3a0ceb-25e9-42d8-abee-24ed46e8185d","resolution":{"observed_at":"2026-08-06T23:27:42.495149Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.801477Z","title":"Analyzing input and output representations for speech-driven gesture generation","venue":null,"work_id":"6a4201a4-62f5-4436-a6ee-c3ac71fb7b17","year":2019},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.550917Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:97efb1d04cd570be9a85ab9e1da83c56c993972eca744f8e6a6d8d442adde6f3","observation_id":"25609d26-928c-4375-bb12-5065c8a8095f","resolution":{"observed_at":"2026-08-06T23:27:45.807207Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.785550Z","title":"Au- dio2gestures: Generating diverse gestures from speech audio with conditional variational autoencoders","venue":null,"work_id":"e4c764b4-eee8-46d2-b767-20723385cf1f","year":2021},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.689624Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:f9d1ea317266fe770ea59165481d7423f5358b8e5a492e56a3fa27a47ba860d7","observation_id":"79e64c09-2848-47f4-9918-95d86b730bc5","resolution":{"observed_at":"2026-08-06T23:27:45.792121Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.08443","last_updated":"2024-07-12T07:12:05Z","snapshot_observed_at":"2026-08-17T09:11:59.372250Z","submitted_at":"2024-07-11T12:33:56Z","title":"Infinite Motion: Extended Motion Generation via Long Text Instructions","version":2},"cited_work":{"arxiv_id":"2407.08443","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.08443","snapshot_observed_at":"2026-08-06T23:27:45.401430Z","title":"Infinite Motion: Extended Motion Generation via Long Text Instructions","venue":"cs.CV","work_id":"5ed069bb-25a8-46eb-9432-7f6c0bb791c1","year":2024},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.736779Z"},"links":{"cited_paper":"/paper/2407.08443","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:97f5da1a299db754565f01d5b78c227d93ad2a380008a4ded787384d7d0488e1","observation_id":"d87e30b9-27d3-484c-babf-53422a71627e","resolution":{"observed_at":"2026-08-06T23:27:45.407544Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.771051Z","title":"Ai choreographer: Music conditioned 3d dance generation with aist++","venue":null,"work_id":"e3f45a79-a7c6-4721-a2a9-0cb8e86c61e0","year":2021},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.856241Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:a9029baf876e38d8e460ccf5a1773cb262058cd2426cc7d35974a28a2289a94a","observation_id":"ecdf3310-10a9-4923-b9ac-b9e703403049","resolution":{"observed_at":"2026-08-06T23:27:45.775916Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:42.936840Z","title":"A survey of multimodel large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:42.936840Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:5d91df552fd0c12a39d9f2d5ffd8f3c4e4c595305b83309fb78308a418394476","observation_id":"33abe0e4-c369-459e-a5be-8dd24423838e","resolution":{"observed_at":"2026-08-06T23:27:42.936840Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.02747","last_updated":"2023-02-08T15:46:05Z","snapshot_observed_at":"2026-08-16T02:30:42.660030Z","submitted_at":"2022-10-06T08:32:20Z","title":"Flow Matching for Generative Modeling","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.02747","snapshot_observed_at":"2026-08-06T23:27:43.009675Z","title":"Flow matching for generative modeling.arXiv preprint arXiv:2210.02747, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:43.009675Z"},"links":{"cited_paper":"/paper/2210.02747","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:0949f20a98f9314746318be9fedfb554a0dd1699a182f129a57bb1a293588a24","observation_id":"b00ffc27-4117-40dd-b262-b58435d78abd","resolution":{"observed_at":"2026-08-06T23:27:43.009675Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:43.127714Z","title":"Beat: A large-scale semantic and emotional multi-modal dataset for conversa- tional gestures synthesis","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:43.127714Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:263f3a3adc8f4b61c357f325f9286bf0ffcf2c569c8a72c24506858985c8fe72","observation_id":"8b592e95-dba2-4068-a28e-05b0e1a81b38","resolution":{"observed_at":"2026-08-06T23:27:43.127714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:43.227851Z","title":"Visual instruction tuning.Advances in neural information processing systems, 36:34892–34916, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:43.227851Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:f858dad8c9586dfa4313bccba19d899eb8ad5d0e1a65e237aa57c96616bfc66f","observation_id":"95e183fa-59da-4890-9eb4-b387a9854c97","resolution":{"observed_at":"2026-08-06T23:27:43.227851Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.737289Z","title":"Amass: Archive of motion capture as surface shapes","venue":null,"work_id":"486c4311-7057-47f9-9aa0-f6eef1989d6e","year":2019},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:43.370490Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:ab9e1ca1e60771ff12708b4ccaa94277d1dd0260848519316fd7a9a2c2d8f0bb","observation_id":"b074d865-de4b-473e-be3f-77e84cecbd71","resolution":{"observed_at":"2026-08-06T23:27:45.740960Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:43.440147Z","title":"The kit whole-body human motion database","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:43.440147Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:8068f44321189888a9180dd250f47abaab0093e6de6c8d42f6541754911b03de","observation_id":"dafbfee2-9fa1-492e-8948-23691e63c9d9","resolution":{"observed_at":"2026-08-06T23:27:43.440147Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.718580Z","title":"Gan-based reactive motion synthesis with class-aware discriminators for human–human interaction.Computers & Graphics, 102:634–645, 2022","venue":null,"work_id":"3b48402e-08e1-4b1e-a05e-b6499859d469","year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:43.571777Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:3a5122061c813b6d16d54994926620171dc3e17532ef961a81f06d60a1db26af","observation_id":"4f1350f8-20f1-47af-ad4b-0d602b3eb3c7","resolution":{"observed_at":"2026-08-06T23:27:45.722450Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.703925Z","title":"Action-conditioned 3d human motion synthesis with transformer vae","venue":null,"work_id":"1ec049f4-fd92-4083-80ea-ef0fe3018fc7","year":2021},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:43.675630Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:745a3508240671a82443c81fe57f7dd0a0b1adf1ffe7b252aa5aba18d03c2bd2","observation_id":"9c951a02-b31e-4534-b591-8f36d55fed0d","resolution":{"observed_at":"2026-08-06T23:27:45.708698Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.690094Z","title":"Bamm: Bidirectional autoregressive motion model","venue":null,"work_id":"c19ddb64-2573-4094-9923-954de309df86","year":2024},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:43.787534Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:bdca4bb3d53825d5389b6bd94fad5a16f98447460b0621737b18081a243853bf","observation_id":"e68d06b3-b7a4-4f46-9da2-5b8a3a7add91","resolution":{"observed_at":"2026-08-06T23:27:45.695175Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.676917Z","title":"Babel: Bodies, action and behavior with english labels","venue":null,"work_id":"3b82bcdb-9316-4d0f-9069-048cbcd5165d","year":2021},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:43.948822Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:d69152648d3dbd4fae99c893f0805d191488a99f757fd859cb27d50518fea0f6","observation_id":"d9dacdf2-b652-4c66-9f2a-a738cfa11e26","resolution":{"observed_at":"2026-08-06T23:27:45.681363Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:43.968345Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:43.968345Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:c6935f1656a67e49758f98ef9b1e16e9741c5a67f430ebd04465bddb4b452cfd","observation_id":"1f2469ff-d6d8-417d-bbb7-8ec95f704fde","resolution":{"observed_at":"2026-08-06T23:27:43.968345Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:44.043467Z","title":"U-net: Convolutional networks for biomedical image segmentation","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:44.043467Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:0d17c9dd7f0518d05a4dbb85993de0b39197f5e8f52d706e0bbc71d681d8f3b7","observation_id":"d82057e1-efa5-4331-92ed-b515759d9f69","resolution":{"observed_at":"2026-08-06T23:27:44.043467Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:44.142452Z","title":"Human motion diffusion as a generative prior","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:44.142452Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:343ee2372d2477cd4e70b16fe63423b85c568cd9ca833b0bc8f2570ace3cb1a6","observation_id":"31529926-e519-4b02-a244-9a8a2030b739","resolution":{"observed_at":"2026-08-06T23:27:44.142452Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2011.13456","last_updated":"2021-02-10T18:17:04Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-11-26T19:39:10Z","title":"Score-Based Generative Modeling through Stochastic Differential Equations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2011.13456","snapshot_observed_at":"2026-08-06T23:27:44.186577Z","title":"Score-based generative modeling through stochastic differential equations.arXiv preprint arXiv:2011.13456, 2020","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:44.186577Z"},"links":{"cited_paper":"/paper/2011.13456","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:b2c1f8a40e6b04d558af470a6dd6d24fd3369b8aa6989413989507050b1cc1c3","observation_id":"9bdc2a60-e9bc-497c-840b-5855fc38fc15","resolution":{"observed_at":"2026-08-06T23:27:44.186577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.644904Z","title":"Goal: Generating 4d whole-body motion for hand-object grasping","venue":null,"work_id":"51e0628c-8539-44fd-bc62-1f99628c4b65","year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:44.318731Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:88bdbece4be5d0b02c76cd4105d45421f8678c98e25d1549b93c344bdfab82f9","observation_id":"8bc70c01-ae56-4358-9f5a-7df27e6577ef","resolution":{"observed_at":"2026-08-06T23:27:45.648816Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.629556Z","title":"Motionclip: Exposing human motion generation to clip space","venue":null,"work_id":"124dedf5-41ce-479d-b3de-945367d28c06","year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:44.431258Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:13d033dfc46d7ebc7c78af74656128db53392b4276570b98be2e4dc3317455b3","observation_id":"d831c7a9-40fb-42f7-8dc7-45d51cdbcd4e","resolution":{"observed_at":"2026-08-06T23:27:45.636282Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:44.557955Z","title":"Human motion diffusion model","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:44.557955Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:bf5381f74188a60aed31a0ae3b7d2f0fbf6d9950c9f505fe7ddcc1cd245ac240","observation_id":"b8635515-863b-43b7-b7f6-aa9f4b806f47","resolution":{"observed_at":"2026-08-06T23:27:44.557955Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:44.720922Z","title":"Neural discrete representation learning.Advances in neural information processing systems, 30, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:44.720922Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:0d7c7100d80a6aab2aa038c5ecec3ce478f3c74695f18770e71bb58da46b42d6","observation_id":"4e695a6a-a812-4826-a7c8-001b8cafed4f","resolution":{"observed_at":"2026-08-06T23:27:44.720922Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.595298Z","title":"Scene-aware generative network for human motion synthesis","venue":null,"work_id":"c08be87b-bb3a-4091-bec6-fc2b3b129ba9","year":2021},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:44.891281Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:0eed31b4952dd85bf7feb5f3d0df7bd2d1fa7798837764f65ad9b776eb76086f","observation_id":"7a406796-3a34-4788-9783-7bbac4e63607","resolution":{"observed_at":"2026-08-06T23:27:45.600223Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.06805","last_updated":"2024-01-18T07:31:47Z","snapshot_observed_at":"2026-08-16T17:40:08.900487Z","submitted_at":"2024-01-10T15:29:21Z","title":"Exploring the Reasoning Abilities of Multimodal Large Language Models (MLLMs): A Comprehensive Survey on Emerging Trends in Multimodal Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.06805","snapshot_observed_at":"2026-08-06T23:27:45.028067Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.028067Z"},"links":{"cited_paper":"/paper/2401.06805","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:4c32e289f4dcbd16990674b208b6d9c4447bdc06f0c2df35ce9e27d5495375b7","observation_id":"44dab7cc-acd6-4a79-9d92-9c0cb37d0b3f","resolution":{"observed_at":"2026-08-06T23:27:45.028067Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21747","last_updated":"2026-06-08T02:14:45Z","snapshot_observed_at":"2026-08-16T13:05:06.623461Z","submitted_at":"2024-10-29T05:25:34Z","title":"MotionGPT-2: A General-Purpose Motion-Language Model for Motion Generation and Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21747","snapshot_observed_at":"2026-08-06T23:27:45.119692Z","title":"Motiongpt-2: A general-purpose motion-language model for motion generation and understanding.arXiv preprint arXiv:2410.21747, 2024","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.119692Z"},"links":{"cited_paper":"/paper/2410.21747","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:d599bf37b2c51a551d64927896c6874650a8e92087561b393815e976d739eab4","observation_id":"e9c64be4-d6a7-4840-b204-a360bf2d2205","resolution":{"observed_at":"2026-08-06T23:27:45.119692Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.07682","last_updated":"2022-10-26T05:06:24Z","snapshot_observed_at":"2026-08-17T00:14:01.413100Z","submitted_at":"2022-06-15T17:32:01Z","title":"Emergent Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.07682","snapshot_observed_at":"2026-08-06T23:27:45.217350Z","title":"Emergent abilities of large language models.arXiv preprint arXiv:2206.07682, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.217350Z"},"links":{"cited_paper":"/paper/2206.07682","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:08c69741b01cc69f62987b54159e957517232df4e5af6476d1b36d5ad035ade5","observation_id":"042684c2-8fef-4d0e-9957-3cc89415d60a","resolution":{"observed_at":"2026-08-06T23:27:45.217350Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.582189Z","title":"What are diffusion models?lilianweng","venue":null,"work_id":"225288fc-0104-451d-b759-6c0a8a14845c","year":2021},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.241093Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:32fe8c800fb80db79ae8cc2c40ab27bf5809573cd01054450b91fa034ea8088e","observation_id":"34dfba78-1941-413e-94be-e7b21e2a4062","resolution":{"observed_at":"2026-08-06T23:27:45.587048Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.569794Z","title":"Motion- agent: A conversational framework for human motion generation with LLMs","venue":null,"work_id":"ff061909-d6cf-4843-bd4a-a2bed324fc3f","year":2025},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.244205Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:44913b45fe802af5a6acb7ac6451fb004ae17a8f07b6cea78907c303534d4603","observation_id":"638ead98-4c4a-4941-ad2f-509994d3bf05","resolution":{"observed_at":"2026-08-06T23:27:45.573623Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.556633Z","title":"Saga: Stochastic whole-body grasping with contact","venue":null,"work_id":"d97e2c26-da50-4d6e-bb89-ee0207f3a0b0","year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.247621Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:008d609963e114d0f3713f746cbd108d2c03dae584bbf4381ec6ee8578a076be","observation_id":"c74c5ab4-2db8-40c0-aaf3-b51086f03b0d","resolution":{"observed_at":"2026-08-06T23:27:45.560929Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.540960Z","title":"Actformer: A gan-based transformer towards general action-conditioned 3d human motion generation","venue":null,"work_id":"a2e6e0ba-96c8-4b83-98ac-040f37c6d3ba","year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.250954Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:fb853258e1a790343d95da2b0bf69c83b70d1fc85ba321036367aabed0f21234","observation_id":"35a731e4-c5aa-450d-9f43-2249fe1a2292","resolution":{"observed_at":"2026-08-06T23:27:45.546818Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.527817Z","title":"Qpgesture: Quantization-based and phase-guided motion matching for natural speech- driven gesture generation","venue":null,"work_id":"9c593d31-d767-4014-b02e-b2aed0c1b5ce","year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.253922Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:a905cc65693894a620992088bcb349943a50c424368fe253c085b791d1862a42","observation_id":"3a5542d0-66f4-4b15-a776-fec27c11f0b1","resolution":{"observed_at":"2026-08-06T23:27:45.532734Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.514300Z","title":"Structure-aware human- action generation","venue":null,"work_id":"37e2cf10-6aef-488e-bb18-f4584b0340da","year":2020},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.257282Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:e5d89a5c1c9e4f69439ff9de25e94a4685bbefedcb8e5603fd31fa039dcc805d","observation_id":"77c05760-ddcd-4ade-84b1-06193add16bd","resolution":{"observed_at":"2026-08-06T23:27:45.518658Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11000","last_updated":"2023-05-19T14:41:16Z","snapshot_observed_at":"2026-08-16T15:31:59.521123Z","submitted_at":"2023-05-18T14:23:25Z","title":"SpeechGPT: Empowering Large Language Models with Intrinsic Cross-Modal Conversational Abilities","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11000","snapshot_observed_at":"2026-08-06T23:27:45.261034Z","title":"Speechgpt: Empowering large language models with intrinsic cross-modal conversational abilities.arXiv preprint arXiv:2305.11000, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.261034Z"},"links":{"cited_paper":"/paper/2305.11000","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:a054defc6da6e534ec5ae1cf700f2fbf889cc15ec2c7b1b4b53dcc1501ba92ae","observation_id":"02036f03-6e50-4726-acf0-08ef561ae583","resolution":{"observed_at":"2026-08-06T23:27:45.261034Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.264481Z","title":"T2m-gpt: Generating human motion from textual descriptions with discrete representations","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.264481Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:0cd6982084f142d0ad3c89098a6762e6342c18863da6a816be59acb7122fa120","observation_id":"1ed6f148-46f4-4da3-b350-93188cca6e25","resolution":{"observed_at":"2026-08-06T23:27:45.264481Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2208.15001","last_updated":"2022-08-31T17:58:54Z","snapshot_observed_at":"2026-08-16T16:35:36.335456Z","submitted_at":"2022-08-31T17:58:54Z","title":"MotionDiffuse: Text-Driven Human Motion Generation with Diffusion Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2208.15001","snapshot_observed_at":"2026-08-06T23:27:45.267627Z","title":"Motiondiffuse: Text-driven human motion generation with diffusion model.arXiv preprint arXiv:2208.15001, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.267627Z"},"links":{"cited_paper":"/paper/2208.15001","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:7c9f9d138c9f09c93aa2b942ffc92b11c11435ba9c5b314e06108f147a57c8bc","observation_id":"1534295e-ad2a-41a2-9952-083cc4e500c8","resolution":{"observed_at":"2026-08-06T23:27:45.267627Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.493412Z","title":"Remodiffuse: Retrieval-augmented motion diffusion model","venue":null,"work_id":"964fef7b-ec07-483e-90d0-1c86f12f5bf0","year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.271788Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:5073d9414d87231728afc3c1f0f8451e365c48932aef602d0f2b15f3d4d3e41b","observation_id":"d9d41000-93b0-4a9c-a6cb-628446c0ad75","resolution":{"observed_at":"2026-08-06T23:27:45.497864Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.275289Z","title":"Finemogen: Fine-grained spatio-temporal motion generation and editing.NeurIPS, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.275289Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:9b4a61b664c756d3e046619a715182c31d3860047cc9d475f3b31b2a2fbc89f7","observation_id":"0452802d-75d5-45c7-838f-1aba360c4ef1","resolution":{"observed_at":"2026-08-06T23:27:45.275289Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.279431Z","title":"Tinyllama: An open-source small language model, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.279431Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:bd57b96a1c3cca648003ce44e3f8b2f63ee1a77741d399aa8c0299f97419ebb5","observation_id":"0b83881c-b5db-45fa-8ec5-3f2e9e434c08","resolution":{"observed_at":"2026-08-06T23:27:45.279431Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.465730Z","title":"Large language models as commonsense knowledge for large-scale task planning.Advances in Neural Information Processing Systems, 36:31967– 31987, 2023","venue":null,"work_id":"dbbfac8e-9a11-4327-b14e-8a794cbf29fb","year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.282495Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:148d27df72725ee7923b6d00e4ed491c47ebda79eff1510ab92d70affd676bf9","observation_id":"7440acb5-c7bb-4952-844e-ea77fc968eba","resolution":{"observed_at":"2026-08-06T23:27:45.470701Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13372","last_updated":"2024-06-27T22:44:48Z","snapshot_observed_at":"2026-08-17T02:51:06.474773Z","submitted_at":"2024-03-20T08:08:54Z","title":"LlamaFactory: Unified Efficient Fine-Tuning of 100+ Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13372","snapshot_observed_at":"2026-08-06T23:27:45.285511Z","title":"Llamafactory: Unified efficient fine-tuning of 100+ language models.arXiv preprint arXiv:2403.13372, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.285511Z"},"links":{"cited_paper":"/paper/2403.13372","citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:4b7035a44fea959337a75cda8dc54d93d3035fbc0111b6a24738883aae90f3b7","observation_id":"e255058e-104e-46f3-92d1-93423f7ee31e","resolution":{"observed_at":"2026-08-06T23:27:45.285511Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:27:45.454944Z","title":"base” for “residual","venue":null,"work_id":"3ea6639b-dc15-44d6-a1f6-06274aa01f09","year":2023},"citing_paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-06T23:27:45.288936Z"},"links":{"citing_paper":"/paper/2506.17912"},"observation_digest":"sha256:75b31212f3b87882985d3d112e9b036f9ace05865cf0a2f971df04c7169d0968","observation_id":"4aab62f8-9254-4e92-8f11-be20860702c7","resolution":{"observed_at":"2026-08-06T23:27:45.458664Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.17912","last_updated":"2025-06-22T06:24:53Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-17T09:09:00.156035Z","submitted_at":"2025-06-22T06:24:53Z","title":"PlanMoGPT: Flow-Enhanced Progressive Planning for Text to Motion Synthesis"},"reference_resolution":{"displayed":63,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":33,"verified_exact":1,"verified_fuzzy":29},"total_outbound_references":63},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 63 of 63 outbound references and 1 inbound Pith citation observation for arXiv:2506.17912."}