{"as_of":"2026-08-23T01:31:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:480e7cd64cd89a12a040d6788b54df5360a8225331e2c35bed0e86d08fd79ba2","coverage":[{"denominator":55,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":55,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T23:58:48.341938Z","state":"measured"},{"denominator":60,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":60,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-22T06:32:14.747728+00:00","state":"measured"},{"denominator":5,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":5,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-14T04:18:05.065372Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T09:39:46.910380Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"cited_work":{"arxiv_id":"2506.15483","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.15483","snapshot_observed_at":"2026-07-04T09:39:46.910380Z","title":"Task-oriented human- object interactions generation with implicit neural representations","venue":null,"work_id":"886fb306-daad-46e3-9245-7a3020fd7314","year":2025},"citing_paper":{"arxiv_id":"2604.04843","last_updated":"2026-04-06T16:44:02Z","snapshot_observed_at":"2026-08-02T08:09:49.884908Z","submitted_at":"2026-04-06T16:44:02Z","title":"InfBaGel: Human-Object-Scene Interaction Generation with Dynamic Perception and Iterative Refinement","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T20:10:16.038861Z"},"links":{"cited_paper":"/paper/2506.15483","citing_paper":"/paper/2604.04843"},"observation_digest":"sha256:379ad36f0e65d2fbda477ff3ae9d342822b5e65e66cde58b534d4e641a8f5feb","observation_id":"93607517-7bf4-458b-976a-33425fd1a96f","resolution":{"observed_at":"2026-05-10T22:10:50.053682Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"cited_work":{"arxiv_id":"2506.15483","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.15483","snapshot_observed_at":"2026-07-04T09:39:46.910380Z","title":"Task-oriented human- object interactions generation with implicit neural representations","venue":null,"work_id":"886fb306-daad-46e3-9245-7a3020fd7314","year":2025},"citing_paper":{"arxiv_id":"2604.27491","last_updated":"2026-04-30T06:44:10Z","snapshot_observed_at":"2026-08-16T17:13:27.587995Z","submitted_at":"2026-04-30T06:44:10Z","title":"Uni-HOI:A Unified framework for Learning the Joint distribution of Text and Human-Object Interaction","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-07T09:40:37.124547Z"},"links":{"cited_paper":"/paper/2506.15483","citing_paper":"/paper/2604.27491"},"observation_digest":"sha256:8eee769e0d809d7faf104602d89ae5afeb0802a7f1da129701b102905bdd0a22","observation_id":"054e805a-e905-4a7f-b2fa-63ea0548617c","resolution":{"observed_at":"2026-05-12T09:41:27.149685Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"cited_work":{"arxiv_id":"2506.15483","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.15483","snapshot_observed_at":"2026-07-04T09:39:46.910380Z","title":"Task-oriented human- object interactions generation with implicit neural representations","venue":null,"work_id":"886fb306-daad-46e3-9245-7a3020fd7314","year":2025},"citing_paper":{"arxiv_id":"2606.22806","last_updated":"2026-06-22T03:32:35Z","snapshot_observed_at":"2026-08-15T19:52:00.128447Z","submitted_at":"2026-06-22T03:32:35Z","title":"Policy-as-Data: Learning Generalizable HOI Diffusion Models from Simulated Physics","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-26T09:40:37.183422Z"},"links":{"cited_paper":"/paper/2506.15483","citing_paper":"/paper/2606.22806"},"observation_digest":"sha256:380b9ffee6f853e1e1e9f2506196e817d1352de84e78f01f0d91fa1a15ab7699","observation_id":"d685115a-385d-4405-b716-2dc4a66975b5","resolution":{"observed_at":"2026-07-04T09:39:46.911736Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"cited_work":{"arxiv_id":"2506.15483","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.15483","snapshot_observed_at":"2026-07-04T09:39:46.910380Z","title":"Task-oriented human- object interactions generation with implicit neural representations","venue":null,"work_id":"886fb306-daad-46e3-9245-7a3020fd7314","year":2025},"citing_paper":{"arxiv_id":"2606.28215","last_updated":"2026-06-26T16:05:58Z","snapshot_observed_at":"2026-08-07T08:50:21.361337Z","submitted_at":"2026-06-26T16:05:58Z","title":"HAT-4D: Lifting Monocular Video for 4D Multi-Object Interactions via Human-Agent Collaboration","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-29T04:18:02.341742Z"},"links":{"cited_paper":"/paper/2506.15483","citing_paper":"/paper/2606.28215"},"observation_digest":"sha256:89889f1fecb24277aa115459e90467f07bae49cd5ef9c9aaaef27c50f6855f49","observation_id":"2f28fe85-7fee-4423-b8c4-a295a0a4703e","resolution":{"observed_at":"2026-07-01T17:05:50.601459Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.15483","snapshot_observed_at":"2026-08-14T04:18:05.065372Z","title":"arXiv preprint arXiv:2506.15483 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10162","last_updated":"2026-08-15T05:45:33Z","snapshot_observed_at":"2026-08-20T23:09:48.439097Z","submitted_at":"2026-08-10T19:30:59Z","title":"MAD-HOI: Masked Autoregressive Diffusion for Generating Articulated Hand Object Interactions from Text","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-14T04:18:05.065372Z"},"links":{"cited_paper":"/paper/2506.15483","citing_paper":"/paper/2608.10162"},"observation_digest":"sha256:6eefe4d4a5911810ac45460e1e40bc9c1374ff0638727627c108059e34479a04","observation_id":"8bdcee07-0bf2-434b-92c0-22ff9488fcce","resolution":{"observed_at":"2026-08-14T04:18:05.065372Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2506.15483/citation-record","integrity":"/paper/2506.15483/integrity","json":"/paper/2506.15483/citation-record.json","paper":"/paper/2506.15483"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:55.480912Z","title":"Behave: Dataset and method for tracking human object interactions","venue":null,"work_id":"3a47dbe5-9022-4088-bffe-9e7528afe23f","year":2022},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:42.002765Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:f9ca9663a9e88c97ad87974ca9eec7a31aca28719593e7bd9f8db5d8815b79f9","observation_id":"8bb89a2a-4470-4081-8065-abdc7807f47c","resolution":{"observed_at":"2026-08-06T23:58:55.610618Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:55.234785Z","title":"Text2hoi: Text-guided 3d motion generation for hand-object interaction","venue":null,"work_id":"ed57668f-5d92-4c4b-a5f6-8e313e5b2c6e","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:42.096650Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:fd3681c29b3404bea57750c008162e2d557dd3b6f74f00fc7e15e05e48bb0dfe","observation_id":"cf8b1e63-adba-4c43-8c9d-71b636943c0e","resolution":{"observed_at":"2026-08-06T23:58:55.353068Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.01291","last_updated":"2025-03-03T08:28:40Z","snapshot_observed_at":"2026-08-18T01:08:30.033776Z","submitted_at":"2025-03-03T08:28:40Z","title":"SemGeoMo: Dynamic Contextual Human Motion Generation with Semantic and Geometric Guidance","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.01291","snapshot_observed_at":"2026-08-06T23:58:42.146744Z","title":"Semgeomo: Dynamic contextual hu- man motion generation with semantic and geometric guidance.arXiv preprint arXiv:2503.01291, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:42.146744Z"},"links":{"cited_paper":"/paper/2503.01291","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:8895955813a3b1ff59deb1e868060d1b9058b390368345e964facdb30bddfb78","observation_id":"85e99873-00ef-4e60-ba51-2766fb97560d","resolution":{"observed_at":"2026-08-06T23:58:42.146744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:42.280842Z","title":"Human-object interaction with vision-language model guided relative movement dynamics","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:42.280842Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:3140ca3734489f81f0e648cee67451f52c670bf212b315bdfaca59d289897360","observation_id":"b9d195ff-beb5-49aa-a07d-b54d14f5bd41","resolution":{"observed_at":"2026-08-06T23:58:42.280842Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:55.051818Z","title":"Cg-hoi: Contact-guided 3d human-object interaction genera- tion","venue":null,"work_id":"23c4cb9b-17c4-498c-ac53-b7155ad2af50","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:42.413785Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:15478a301550cb5b771f4f3e9ad24c3d2d7c28c41561f40abb2d36a0db9ad710","observation_id":"9f4a522b-7c4f-4ae8-b817-7fe55c820f53","resolution":{"observed_at":"2026-08-06T23:58:55.130817Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:42.518312Z","title":"Arctic: A dataset for dexterous bimanual hand-object manipulation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:42.518312Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:ef592eefe507c8fc8c2b6f63cedcacbc153572d27f08b78ae11b2fbc412bd4f5","observation_id":"3fcb5930-6672-4c86-8c56-9cd9ada8aee2","resolution":{"observed_at":"2026-08-06T23:58:42.518312Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:42.617886Z","title":"3d-future: 3d furniture shape with texture","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:42.617886Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:63c674624f0ca008663211b3b41c1587b37d9dc44abf95cde632b415ab5a11fd","observation_id":"05febdda-f983-40b5-bd5e-68b04c6d02ae","resolution":{"observed_at":"2026-08-06T23:58:42.617886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:54.890420Z","title":"Coohoi: Learning cooperative human-object interaction with manipulated object dynamics","venue":null,"work_id":"6e7e604a-236e-4629-8f23-14764856f783","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:42.740041Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:d40d36a6da8c8b49d239c7e39c3d8998691745cfb84edfadca9d3e1a1fb3587e","observation_id":"7b00bdab-6a56-4c82-af4f-8367bfc1c1f9","resolution":{"observed_at":"2026-08-06T23:58:54.955569Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:54.643100Z","title":"Auto-regressive diffusion for generating 3d human-object interactions","venue":null,"work_id":"12dac137-c3da-4959-9e33-ed87209324e6","year":2025},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:42.857804Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:801cd3ec2db2720454817ab52947ba4d599f33b6f230f6b23fe74f5a4f7295ef","observation_id":"27d1151e-06cd-44fd-b7d3-a2a3b64e6a04","resolution":{"observed_at":"2026-08-06T23:58:54.784352Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:54.452817Z","title":"Imos: Intent-driven full-body motion synthesis for human-object interactions","venue":null,"work_id":"77c79de5-0cba-4fcb-baca-ff89f8d7dc7c","year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:42.981739Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:d757a9b6d60d6f71c7480b22006643912a56db302278198845040ff1df5c0c0f","observation_id":"cddcc914-3e35-4084-95ca-52ac5d239aed","resolution":{"observed_at":"2026-08-06T23:58:54.535978Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:54.197988Z","title":"Generating diverse and natural 3d human motions from text","venue":null,"work_id":"c60fb35b-30b1-4399-b38f-80f03da14d33","year":2022},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.117169Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:a35b48ccfc79cc75f8815768cb407f265488ff60e9195cd8613da4e17d3eda60","observation_id":"d41e7f0d-488a-4f5d-93da-11a84bd01880","resolution":{"observed_at":"2026-08-06T23:58:54.285016Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.06188","last_updated":"2025-05-09T17:25:34Z","snapshot_observed_at":"2026-08-16T13:36:03.593764Z","submitted_at":"2024-07-08T17:59:36Z","title":"CrowdMoGen: Zero-Shot Text-Driven Collective Motion Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.06188","snapshot_observed_at":"2026-08-06T23:58:43.178155Z","title":"Crowdmogen: Zero-shot text-driven collective motion generation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.178155Z"},"links":{"cited_paper":"/paper/2407.06188","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:06f90227c34c04fda40ed4e62a8e148eb6c9f9de9bc77993db1875a1eff5433b","observation_id":"8677cafd-01c2-46cf-a19f-f9085de28b16","resolution":{"observed_at":"2026-08-06T23:58:43.178155Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:53.945099Z","title":"Stochastic scene-aware motion prediction","venue":null,"work_id":"ca68ddc0-c448-498f-9c23-41898fea7136","year":2021},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.243697Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:daa56be82310d674cbd2575e2ebce13fafd81a09eab1057cb56ba027919393a9","observation_id":"d89d6b14-23ce-4a28-9cc2-b7bf1e2b18b6","resolution":{"observed_at":"2026-08-06T23:58:54.080788Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:53.680850Z","title":"Resolving 3d human pose ambiguities with 3d scene constraints","venue":null,"work_id":"a6efba66-c6c6-48ae-a383-dd5fe0e1e618","year":2019},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.295297Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:e84d5253414d228a70a6b3c7c9e60af902c38b98296d72d1dac59ba7acb0cb33","observation_id":"2db44fbe-bf52-4b32-95a1-ff7a1f0ed0d8","resolution":{"observed_at":"2026-08-06T23:58:53.767052Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:53.489310Z","title":"Nemf: Neural motion fields for kinematic animation","venue":null,"work_id":"f1f18322-eb29-4778-b0ee-8160953ab076","year":2022},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.377972Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:e6676cf1c8f05536deed471a5361cb42bedda8520ba37babdfab96c874fb6a57","observation_id":"8b84c955-aa78-4ce2-b618-49fca0fad081","resolution":{"observed_at":"2026-08-06T23:58:53.594127Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:43.450104Z","title":"Denoising diffusion probabilistic models","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.450104Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:17a8f2fc4d61a2f6037f9eea76a428228a32315caa6db4e45ffc15fd8d5e20f1","observation_id":"778e589c-1164-47db-9272-658214621178","resolution":{"observed_at":"2026-08-06T23:58:43.450104Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:53.335629Z","title":"Diffusion-based generation, optimization, and planning in 3d scenes","venue":null,"work_id":"c007a1ea-2f06-43be-9062-2c2a054884b1","year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.509467Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:13c106dd4183909d538fbc84953731bb3bc4631f3139f5b2e1e40245367e63f0","observation_id":"ec09d9a2-c466-4b4e-8df2-79a11d1514bb","resolution":{"observed_at":"2026-08-06T23:58:53.403773Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:53.121256Z","title":"Intercap: Joint markerless 3d tracking of humans and objects in interaction","venue":null,"work_id":"4e314c33-7900-4c38-a149-653d20359194","year":2022},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.613275Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:ad9f8520a3d67e78fce118d9ec83406df8cf1de35f7dc8e99d1d63b5d9c4c274","observation_id":"aca699ab-8897-4745-a08a-38de29cdaf2c","resolution":{"observed_at":"2026-08-06T23:58:53.189998Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:52.952801Z","title":"Full-body articulated human-object interaction","venue":null,"work_id":"3f79f7c5-d49f-4688-8668-67aa5cc135e4","year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.712282Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:f484676a7fff16bde4e35bcd202c15a3eb676a664a1bc474c37e0f975c7d4a31","observation_id":"8273493e-e640-4468-8a39-dc4c22ee0aa3","resolution":{"observed_at":"2026-08-06T23:58:53.041084Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.08333","last_updated":"2025-08-11T11:45:30Z","snapshot_observed_at":"2026-08-16T12:58:57.267113Z","submitted_at":"2025-01-14T18:59:59Z","title":"DAViD: Modeling Dynamic Affordance of 3D Objects Using Pre-trained Video Diffusion Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.08333","snapshot_observed_at":"2026-08-06T23:58:43.828049Z","title":"David: Modeling dynamic affordance of 3d objects using pre-trained video diffusion models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.828049Z"},"links":{"cited_paper":"/paper/2501.08333","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:c513151830f7afed5703a5cfb8f3af29fc2c16c378c0a70ccc3c0fe59b350b19","observation_id":"b0ae1742-aaad-4d24-970d-93668506cf04","resolution":{"observed_at":"2026-08-06T23:58:43.828049Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.18600","last_updated":"2025-03-21T16:17:28Z","snapshot_observed_at":"2026-08-20T14:08:55.582347Z","submitted_at":"2024-12-24T18:55:38Z","title":"ZeroHSI: Zero-Shot 4D Human-Scene Interaction by Video Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.18600","snapshot_observed_at":"2026-08-06T23:58:43.894378Z","title":"Zerohsi: Zero-shot 4d human-scene interaction by video generation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.894378Z"},"links":{"cited_paper":"/paper/2412.18600","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:37e6f959aa85afa202b6917fa59f5070fc8a8b7493029df2a7291b1f6ab244c6","observation_id":"4dade97a-83b4-4102-8f52-9fd0d3b347dd","resolution":{"observed_at":"2026-08-06T23:58:43.894378Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:52.834883Z","title":"Controllable human-object interaction synthesis","venue":null,"work_id":"a1d7b510-3a34-4546-8a6f-8068233089b6","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:43.996546Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:0ee6a73ed8454e51f4624d1c6eb81b29d75efcbf1be404d4350cdfdecb065488","observation_id":"e24f0525-b002-4026-8f61-c62222bba92d","resolution":{"observed_at":"2026-08-06T23:58:52.887739Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:52.579439Z","title":"Object motion guided human motion synthesis","venue":null,"work_id":"eaa4f3e8-018f-4abf-a737-6a28a56bf7e0","year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:44.468408Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:cd74d4ea33b3f58dc258856a5d764cc848cf55b7e9d90bc43259d9bcf855f76a","observation_id":"b76309c1-7e2f-4d7a-8331-0a9c4e8c6220","resolution":{"observed_at":"2026-08-06T23:58:52.738866Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:52.394032Z","title":"Intergen: Diffusion-based multi-human motion generation under complex interactions","venue":null,"work_id":"7e8b8e29-602a-47f9-9523-058e56a8b01e","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:45.214411Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:51cc5234cc3684b78b060b35047b76b8fe4637b4915d44f5f147b69197d34a5d","observation_id":"761e0590-d64d-4f3f-9027-ed3379687b2c","resolution":{"observed_at":"2026-08-06T23:58:52.468286Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:52.164233Z","title":"Learning basketball dribbling skills using trajectory optimization and deep reinforcement learning","venue":null,"work_id":"955c45ae-f6d0-480e-8bc8-c14b8385685a","year":2018},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:45.311591Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:62a59b14a430d93a6ff0d5afe691fdf95e8c5fb774fe659be0939a3b4afe89b4","observation_id":"69c81e82-d443-4cb7-a092-61be08398c86","resolution":{"observed_at":"2026-08-06T23:58:52.291512Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:51.988861Z","title":"Motion-x: A large-scale 3d expressive whole-body human motion dataset","venue":null,"work_id":"7cf00006-f6df-4ef2-8555-585bfa28ec6a","year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:45.390616Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:8520edbdc22333f266579fd2fb1af6d5cf64ae1233ca82a1ae8eb45828ef6ad5","observation_id":"ec0f6bfa-dac6-4331-a1fe-c1738896149f","resolution":{"observed_at":"2026-08-06T23:58:52.058585Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:51.871076Z","title":"Himo: A new benchmark for full-body human interacting with multiple objects","venue":null,"work_id":"6bdbe983-719e-4ec5-9c78-1416dbfcaf4f","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:45.551315Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:49aa59649c7bb0e9fc465f2d901052970bda6a0a73a97739c6f543aca47c14a0","observation_id":"d18afae9-064b-4a91-8855-275f85eaeecc","resolution":{"observed_at":"2026-08-06T23:58:51.922300Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:45.635299Z","title":"Amass: Archive of motion capture as surface shapes","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:45.635299Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:e759161944f60386973379a78f11fab52dd1a71b3227a6377e05df908c5aea70","observation_id":"6b622f35-1ffd-4bd6-93be-0fa579c47152","resolution":{"observed_at":"2026-08-06T23:58:45.635299Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:45.706998Z","title":"Expressive body capture: 3d hands, face, and body from a single image","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:45.706998Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:c585fcebcb737ae5d43e408817e988a35669ac83f9eae6439530c51ca2c05f02","observation_id":"2237fa89-5080-4be3-bb9c-ca82cbde5b1d","resolution":{"observed_at":"2026-08-06T23:58:45.706998Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.06553","last_updated":"2025-07-07T05:09:32Z","snapshot_observed_at":"2026-08-20T10:40:30.650361Z","submitted_at":"2023-12-11T17:41:17Z","title":"HOI-Diff: Text-Driven Synthesis of 3D Human-Object Interactions using Diffusion Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.06553","snapshot_observed_at":"2026-08-06T23:58:45.792596Z","title":"Hoi-diff: Text-driven synthesis of 3d human-object interactions using diffusion models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:45.792596Z"},"links":{"cited_paper":"/paper/2312.06553","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:910b89daf37f6afbf84d4e01466f21c39af00ea7c13b4b65da254b2c646cbde8","observation_id":"31686fc4-93d8-4b42-87e7-0e45c9473103","resolution":{"observed_at":"2026-08-06T23:58:45.792596Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:51.683730Z","title":"Deepmimic: Example- guided deep reinforcement learning of physics-based character skills","venue":null,"work_id":"679962ac-8ebb-4382-a0aa-0b1b2a1ac510","year":2018},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:45.860094Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:ec575b8a0a9a73b92d3689ef3f49412dfb44f0d7b835bad7e1d0eaa58a151997","observation_id":"ddf6243b-079a-4509-9228-7b4cd588e7c0","resolution":{"observed_at":"2026-08-06T23:58:51.783146Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:45.947221Z","title":"Amp: Adversarial motion priors for stylized physics-based character control","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:45.947221Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:1fa1f439673a6e9395be11d826f76fb16e31909792d99e7c4e8f2b1379d6bd4f","observation_id":"f0441dec-e8db-44b1-b207-d67e72be647e","resolution":{"observed_at":"2026-08-06T23:58:45.947221Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:51.453636Z","title":null,"venue":null,"work_id":"cdab25e0-2e8e-4512-92a9-a1b0f05ec825","year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:46.027688Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:ac41e632f9f909f0c88d6815c1d80369f48300b3bef40669473d1d8086a5b417","observation_id":"9d5a696d-fdde-4695-be08-3e2489242943","resolution":{"observed_at":"2026-08-06T23:58:51.555638Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:46.235580Z","title":"Pointnet++: Deep hierarchical feature learning on point sets in a metric space","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:46.235580Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:5b868e0a2654565c1fee8fe2bba029928fee6b1fa4f593dca6cb681470ff49ee","observation_id":"e44631c2-c95a-4c59-b02c-7ec88d863f98","resolution":{"observed_at":"2026-08-06T23:58:46.235580Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:46.339867Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:46.339867Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:29c1d6c6e32f6c53ce1b1da32e9410c8808920b494d18de55ac824d422769b28","observation_id":"81d9f47f-e7df-496a-8f97-b85bb41a3f67","resolution":{"observed_at":"2026-08-06T23:58:46.339867Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:51.306395Z","title":"Hoianimator: Generating text-prompt human-object animations using novel perceptive diffusion models","venue":null,"work_id":"dc684ed0-1b91-42b1-a2d4-3ccf5f473e98","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:46.422003Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:ffa16905a2f67c4acc7ffd2def9e8c3ba25907e14c933321b77ab5b9ef4e0368","observation_id":"ddfc4ec0-aed6-4b92-b569-fb4a694e7267","resolution":{"observed_at":"2026-08-06T23:58:51.361816Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:46.501297Z","title":"A survey on human interaction motion generation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:46.501297Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:ca0d92d2f630bbfadca79c8340f71063099cc272561108bc508b9927b1bdbb5a","observation_id":"69e86349-d96d-4a6e-9f02-a34d0a55d574","resolution":{"observed_at":"2026-08-06T23:58:46.501297Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:46.569130Z","title":"Grab: A dataset of whole-body human grasping of objects","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:46.569130Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:be2f0ab1d255c92566053ce9ea98ac2ef26815736a621cdf0338111c456870e9","observation_id":"ebfd29de-58b8-439a-b51a-4fc1b7828270","resolution":{"observed_at":"2026-08-06T23:58:46.569130Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:51.087817Z","title":"Human motion diffusion model","venue":null,"work_id":"b191eedf-b073-44b5-ade0-f11b7ab17014","year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:46.654005Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:0c531be6c31f8d7eea2189f3f4f13e1d747ef756695b55b54fb8c01d165b7932","observation_id":"27b22746-7cdd-4a96-8c9b-3b0206e71a5a","resolution":{"observed_at":"2026-08-06T23:58:51.194489Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.04393","last_updated":"2023-12-07T16:06:31Z","snapshot_observed_at":"2026-08-20T05:06:00.838468Z","submitted_at":"2023-12-07T16:06:31Z","title":"PhysHOI: Physics-Based Imitation of Dynamic Human-Object Interaction","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.04393","snapshot_observed_at":"2026-08-06T23:58:46.767124Z","title":"Physhoi: Physics-based imitation of dynamic human-object interaction","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:46.767124Z"},"links":{"cited_paper":"/paper/2312.04393","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:c97abbf80d3414f1965c38cfa8efa8af0a0c82828990c425581f65dec684c5a0","observation_id":"e943a2a1-1745-454c-a6ab-e1fac9ac3ad4","resolution":{"observed_at":"2026-08-06T23:58:46.767124Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:50.873507Z","title":"Move as you say interact as you can: Language-guided human motion generation with scene affordance","venue":null,"work_id":"ca962f40-3d7e-4e73-b6bf-7081cc9ec50a","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:46.836988Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:56f0e481a095a2968bd621819c0112756df9d00c79f62668ecc81a16bf9e204a","observation_id":"81684676-5982-45ac-82fb-997deb5a267b","resolution":{"observed_at":"2026-08-06T23:58:50.974359Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:50.635286Z","title":"Humanise: Language-conditioned human motion generation in 3d scenes","venue":null,"work_id":"32fbec90-8e75-4d13-aa81-e12963364567","year":2022},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:46.949651Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:f91f1844d2ab0eda71e78652d3a3ac5b07d99c70dbc30dff8c4cd6e471f0fac2","observation_id":"bbff33ce-185e-4677-a112-366ae08ab09c","resolution":{"observed_at":"2026-08-06T23:58:50.772239Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.15898","last_updated":"2025-03-20T06:50:18Z","snapshot_observed_at":"2026-08-18T12:36:57.430158Z","submitted_at":"2025-03-20T06:50:18Z","title":"Reconstructing In-the-Wild Open-Vocabulary Human-Object Interactions","version":1},"cited_work":{"arxiv_id":"2503.15898","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.15898","snapshot_observed_at":"2026-08-06T23:58:48.756726Z","title":"Reconstructing In-the-Wild Open-Vocabulary Human-Object Interactions","venue":"cs.CV","work_id":"89abc6b2-1fdc-49d1-b6cb-474788fffb3f","year":2025},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:47.018564Z"},"links":{"cited_paper":"/paper/2503.15898","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:78a142d08b3184c3b73a805a3fde02275a0ec31e7e52f2fe049b1aadeccb6cd2","observation_id":"15d450c7-0e08-4df3-84ef-3cb8fb5baa37","resolution":{"observed_at":"2026-08-06T23:58:48.860854Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.11208","last_updated":"2024-03-17T13:17:25Z","snapshot_observed_at":"2026-08-19T21:36:26.361410Z","submitted_at":"2024-03-17T13:17:25Z","title":"THOR: Text to Human-Object Interaction Diffusion via Relation Intervention","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.11208","snapshot_observed_at":"2026-08-06T23:58:47.146503Z","title":"Thor: Text to human-object interaction diffusion via relation intervention","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:47.146503Z"},"links":{"cited_paper":"/paper/2403.11208","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:fb83aedcc482f14228ad495f877ce0c254f701b6216abab12272cf90425b6cbb","observation_id":"e04b1c83-bc29-4261-8147-14d46a3cbd63","resolution":{"observed_at":"2026-08-06T23:58:47.146503Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.17840","last_updated":"2025-08-21T08:23:55Z","snapshot_observed_at":"2026-08-18T02:06:39.510317Z","submitted_at":"2024-06-25T17:46:28Z","title":"Human-Object Interaction from Human-Level Instructions","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.17840","snapshot_observed_at":"2026-08-06T23:58:47.255699Z","title":"Human-object interaction from human-level instructions","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:47.255699Z"},"links":{"cited_paper":"/paper/2406.17840","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:87dd0a8e0fc98741bdfecaf0e445a9f142ff176dee7c5eeeabcfe32d35cc0381","observation_id":"0f7732d0-78c0-4d75-a9ef-811590ab6eea","resolution":{"observed_at":"2026-08-06T23:58:47.255699Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:50.302520Z","title":"Inter-x: Towards versatile human-human interac- tion analysis","venue":null,"work_id":"e08bcefe-cf22-4d0b-a2fc-b88443187e15","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:47.324819Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:cda12b078df2463b3a6db35d5d5f8049c187cb64d0de358571b027cf1486110c","observation_id":"f07d24a9-71a9-4ef3-a05b-5df159e40401","resolution":{"observed_at":"2026-08-06T23:58:50.467461Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:47.416700Z","title":"Interdiff: Generating 3d human-object interactions with physics-informed diffusion","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:47.416700Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:80fd97a7549a02c82ad8ff9b821ebdae9dc4dd3192ad99177128a6004dffca12","observation_id":"fc17981d-cd17-4caf-b811-aa47b6550f6e","resolution":{"observed_at":"2026-08-06T23:58:47.416700Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:47.522418Z","title":"Intermimic: Towards universal whole-body control for physics-based human-object interactions","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:47.522418Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:543322548fcd40184aedbf291d44a7d9826f4e6e07244942f1c3ff06cfa9e9f0","observation_id":"0fb00f35-0c3e-45db-8298-9f68ec685030","resolution":{"observed_at":"2026-08-06T23:58:47.522418Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:47.624950Z","title":"Interdreamer: Zero-shot text to 3d dynamic human-object interaction","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:47.624950Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:82e8163d6c4086108275690dc4f103a1f1a8d3979d344c9ed6735cc281ffb914","observation_id":"577b72e0-40c9-4ea6-9920-caa88a87fb22","resolution":{"observed_at":"2026-08-06T23:58:47.624950Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:50.073823Z","title":"F-hoi: Toward fine- grained semantic-aligned 3d human-object interactions, 2024","venue":null,"work_id":"c3dc3b0e-562c-418b-bfd9-848f04b7e73b","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:47.757343Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:02e6f1341edbb76210accbb67e060979bcde1cc31dfbd8a0a805d308c38b5a52","observation_id":"d804603e-af5e-4944-b6bd-b0195460fa42","resolution":{"observed_at":"2026-08-06T23:58:50.171636Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:49.796406Z","title":"Generating human interaction motions in scenes with text control","venue":null,"work_id":"10f06c59-9e6a-40a4-8c3a-79e55af476fe","year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:47.846996Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:b5c446dcd2d1fe24803a6cb25b54978f684a47464d018ea6c177400cd243ae56","observation_id":"aab35056-7b66-4679-850e-2ddf317a14c8","resolution":{"observed_at":"2026-08-06T23:58:49.923478Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.13130","last_updated":"2025-03-17T12:55:34Z","snapshot_observed_at":"2026-08-20T11:54:19.149834Z","submitted_at":"2025-03-17T12:55:34Z","title":"ChainHOI: Joint-based Kinematic Chain Modeling for Human-Object Interaction Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.13130","snapshot_observed_at":"2026-08-06T23:58:47.939099Z","title":"Chainhoi: Joint-based kinematic chain modeling for human-object interaction generation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:47.939099Z"},"links":{"cited_paper":"/paper/2503.13130","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:ffeeef9b2084d333891fc6b3b075bbb5341f79820f906f2e5780b189404dc7bf","observation_id":"6017ce03-d665-4fa8-9eb8-de93d2096b17","resolution":{"observed_at":"2026-08-06T23:58:47.939099Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.19353","last_updated":"2024-12-24T09:33:24Z","snapshot_observed_at":"2026-08-19T11:12:35.394889Z","submitted_at":"2024-06-27T17:32:18Z","title":"CORE4D: A 4D Human-Object-Human Interaction Dataset for Collaborative Object REarrangement","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.19353","snapshot_observed_at":"2026-08-06T23:58:48.047771Z","title":"Core4d: A 4d human- object-human interaction dataset for collaborative object rearrangement","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:48.047771Z"},"links":{"cited_paper":"/paper/2406.19353","citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:7a6c97f557bcf0934a26b4124655ab5cd5ade691a45682da0db528fd9422e63c","observation_id":"e8fe0b82-aab9-4a9d-9ad1-a6b1fe69443d","resolution":{"observed_at":"2026-08-06T23:58:48.047771Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:49.462939Z","title":"Couch: Towards controllable human-chair interactions","venue":null,"work_id":"5eba9a11-1bde-41e2-b7eb-86747b10d624","year":2022},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:48.200420Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:44d0770f4421bf1a4e64d8089b63ef7610edb4cc187ea3fc63b75997ed4d8631","observation_id":"4ae07036-8fec-41fe-95b4-ff15fcf28ac9","resolution":{"observed_at":"2026-08-06T23:58:49.635573Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:58:49.193097Z","title":"On the continuity of rotation representations in neural networks","venue":null,"work_id":"bea441c5-89ed-4f05-90dd-5f39d1b8ef79","year":2019},"citing_paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T23:58:48.341938Z"},"links":{"citing_paper":"/paper/2506.15483"},"observation_digest":"sha256:6a01cb6c2deec5e3fd9a965d856af1d6e55b6c4be4f91142c2beaa66897192f0","observation_id":"e9a529ed-d606-4a58-8b38-ed3e9a39a955","resolution":{"observed_at":"2026-08-06T23:58:49.333983Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.15483","last_updated":"2025-06-18T14:17:53Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-18T01:09:03.624677Z","submitted_at":"2025-06-18T14:17:53Z","title":"GenHOI: Generalizing Text-driven 4D Human-Object Interaction Synthesis for Unseen Objects"},"reference_resolution":{"displayed":55,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":25,"verified_exact":1,"verified_fuzzy":29},"total_outbound_references":55},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"thesis":"As of 23 August 2026, this Paper Citation Record lists 55 of 55 outbound references and 5 inbound Pith citation observations for arXiv:2506.15483."}