{"as_of":"2026-08-13T16:00:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8d6370ba3654d02f6f3b7f1c647593866088e99b2f9775eb0675800973bb5c5a","coverage":[{"denominator":37,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":37,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T23:10:00.733804Z","state":"measured"},{"denominator":38,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":38,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T13:00:00.743937Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-12T13:00:02.209848Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"cited_work":{"arxiv_id":"2506.19331","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.19331","snapshot_observed_at":"2026-08-12T13:00:02.209848Z","title":"Segment Any 3D-Part in a Scene from a Sentence","venue":"cs.CV","work_id":"988e8e1a-de64-44d5-986f-0469fefaa47d","year":2025},"citing_paper":{"arxiv_id":"2608.10981","last_updated":"2026-08-11T14:36:16Z","snapshot_observed_at":"2026-08-13T15:32:56.955366Z","submitted_at":"2026-08-11T14:36:16Z","title":"ThinkAfford: Affordance-Centric Reasoning for Fine-Grained 3D Grounding in Cluttered Scenes","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-12T13:00:00.743937Z"},"links":{"cited_paper":"/paper/2506.19331","citing_paper":"/paper/2608.10981"},"observation_digest":"sha256:8ec3237b2f99a03fb64a21d56e83f865618a399062347ad287099a66e70d74df","observation_id":"2d5f75ae-749e-4e5a-9982-45b250d87065","resolution":{"observed_at":"2026-08-12T13:00:02.218262Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.19331/citation-record","integrity":"/paper/2506.19331/integrity","json":"/paper/2506.19331/citation-record.json","paper":"/paper/2506.19331"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:07.491763Z","title":"Arkitscenes - a diverse real-world dataset for 3d indoor scene understanding using mobile RGB-d data","venue":null,"work_id":"94b32f2b-5e23-4d90-89db-0884beb52583","year":2021},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:56.874736Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:ef426eff8b223d2eebf078d538c894518a064e896560836742114ec223bb5037","observation_id":"dae0ab43-4a6e-4838-80cd-b8d55f5878bf","resolution":{"observed_at":"2026-08-06T23:10:07.564645Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:07.279019Z","title":"Matterport3d: Learning from rgb-d data in indoor environments","venue":null,"work_id":"07c956eb-4ead-418b-b506-4313ee039349","year":2017},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:56.931179Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:30832c9f166941782e0888eb71f6cc4f066bec11f95abfd445df4e0fb6575e70","observation_id":"2e8cfcec-b24e-4c8b-868f-5c36f514bce8","resolution":{"observed_at":"2026-08-06T23:10:07.402156Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:07.123063Z","title":"Shapenet: An information-rich 3d model repository","venue":null,"work_id":"b67aeefb-44af-4fe2-a776-497d711dcc0d","year":2015},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:57.052097Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:c37f57d9262de6a95cfaa2d033fba6ba399f29f6b683b6b3d06af0803e824d24","observation_id":"0e736a7b-b27e-4e85-a192-ab04b6a4d903","resolution":{"observed_at":"2026-08-06T23:10:07.206013Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:06.975054Z","title":"Clip2scene: Towards label-efficient 3d scene understanding by clip","venue":null,"work_id":"c354d667-cf56-4ccb-8bfb-bfc4b78fc3fa","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:57.163655Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:842b614d89a5ceb74e11eb17758cf13c58d7dbcf297e567a2921f37d0acea5d6","observation_id":"ab56bf19-561b-4173-bc5d-794f672657c3","resolution":{"observed_at":"2026-08-06T23:10:07.057792Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:09:57.290866Z","title":"Scannet: Richly-annotated 3d reconstructions of indoor scenes","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:57.290866Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:7076343fbdd84e74ffc5abf069c444ed3cb7f9a6e0c460f7505d62e788c54c61","observation_id":"b058d641-64ec-423e-b79e-3d7db3e0e918","resolution":{"observed_at":"2026-08-06T23:09:57.290866Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:06.776239Z","title":"SceneFun3D: Fine-grained functionality and affordance understanding in 3d scenes","venue":null,"work_id":"0e3c10f0-f41b-40d6-ba6f-77e129796db7","year":2024},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:57.375246Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:73d852d32ae7dbb423f54888a64a3cecf3b50a36a1da21cc09c68eb183cbd834","observation_id":"a9299ddf-3890-49be-aa81-60677c50d841","resolution":{"observed_at":"2026-08-06T23:10:06.851373Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:06.599457Z","title":"Pla: Language-driven open-vocabulary 3d scene understanding","venue":null,"work_id":"0795dfb7-a238-4975-9d02-6b119ee4b093","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:57.531272Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:aee8898ed3d1af006ac1774e8882a5c662c1332f7675c042f1de7526ae1a99d8","observation_id":"a5854ae3-edf0-479f-b640-c1aa42e6c588","resolution":{"observed_at":"2026-08-06T23:10:06.669803Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:06.438586Z","title":"Lowis3d: Language- driven open-world instance-level 3d scene understanding","venue":null,"work_id":"a50f6280-3782-40ae-8b02-d9fdff413ff7","year":2024},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:57.662839Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:981f42d8ab502e4a74d4ef7606b778e9631bd747423471d57302844e5d53744f","observation_id":"b01d799b-dc3e-4204-b967-6e8040a48344","resolution":{"observed_at":"2026-08-06T23:10:06.523170Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:06.278675Z","title":"Openins3d: Snap and lookup for 3d open-vocabulary instance segmentation","venue":null,"work_id":"f1c1a2e4-26aa-46c4-8f77-7ee783af8acd","year":2024},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:57.800464Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:e1d04db5d1f7b17eedaaca9e952b9f2b3b542f1d6bb5ea75fcba010a35b6c1e1","observation_id":"c9b86397-a68e-46fe-84a1-5dc7c9a17b0f","resolution":{"observed_at":"2026-08-06T23:10:06.343766Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:06.079228Z","title":"Clevr: A diagnostic dataset for compositional language and elementary visual reasoning","venue":null,"work_id":"053b1cc8-d1fc-4005-b85b-392630af0a51","year":2017},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:57.888345Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:063a0bbb3dfb51967273a94f7f0a9b39a4b39549c1a00c907def2969bc8c24ea","observation_id":"3cb568bb-c9a5-497d-8b35-94dbff1bd772","resolution":{"observed_at":"2026-08-06T23:10:06.181432Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:05.899555Z","title":"Lerf: Language embedded radiance fields","venue":null,"work_id":"31f58704-bdc1-4f11-b9ed-87aa3d26dfd5","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:58.007139Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:33c634deab67be04b3a0135a106ff58cfb2b016e82e49db7afe0a41523a05b9c","observation_id":"31590d88-9b23-4d68-a0e0-2239bd42e312","resolution":{"observed_at":"2026-08-06T23:10:05.976482Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:09:58.144205Z","title":"Segment anything","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:58.144205Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:7dfc6c08abebab20fee93cdc8b548ffaf52232fe13f201e650105709eece1032","observation_id":"44c9f49a-a902-426d-b3d4-98ab31da056f","resolution":{"observed_at":"2026-08-06T23:09:58.144205Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:05.720632Z","title":"Cut pursuit: Fast algorithms to learn piecewise constant functions on general weighted graphs","venue":null,"work_id":"3aae4ee5-a53e-437f-9b48-890a77212ee3","year":2017},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:58.268608Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:e5629a89a1dbd08d3a4df446bb6850f3534b5625c604b3834be8b69de6c7011f","observation_id":"05ca29fb-dfae-4146-a5dc-ba6effc0023e","resolution":{"observed_at":"2026-08-06T23:10:05.811775Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:05.563288Z","title":"Large-scale point cloud semantic segmentation with superpoint graphs","venue":null,"work_id":"99d3ef2a-1aae-43ed-b73b-973f46a55d8b","year":2018},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:58.419400Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:350c713685ac8ed592b506cf4de6e5f5a0e87862d3f69d70ca34896b64f37c62","observation_id":"97c34d83-5d01-4c2c-ac10-3e8c7031cf13","resolution":{"observed_at":"2026-08-06T23:10:05.624719Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:05.372872Z","title":"Grounding dino: Marrying dino with grounded pre-training for open-set object detection","venue":null,"work_id":"1fd45c21-065c-43d9-abb2-8f4976f51a83","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:58.578103Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:08ea8c1188a97dff5a592d977a439aa12e8778a63371b4ff829faaae14b4a903","observation_id":"e8ec7941-ff0b-41e6-9428-5659c9487461","resolution":{"observed_at":"2026-08-06T23:10:05.438205Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:05.193035Z","title":"Multiscan: Scalable rgbd scanning for 3d environments with articulated objects","venue":null,"work_id":"c02ded02-9cda-4a63-8f04-b8d02d2cf336","year":2022},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:58.649119Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:89e19e307c0e172dfed8a6a26cc6d5a30613f98694ad095fcf3f2e4aceb28128","observation_id":"0b7a2467-36ed-4e8a-a912-05153fdd6b6e","resolution":{"observed_at":"2026-08-06T23:10:05.279387Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:05.015689Z","title":"Partnet: A large-scale benchmark for fine-grained and hierarchical part-level 3d object understanding","venue":null,"work_id":"f8277a50-180f-4f02-a045-c479d7d5aa3e","year":2019},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:58.783571Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:519ece4259a22ce25d4054fbc44656a341e36d1e23455524424c1c583f90805a","observation_id":"e1e30ffd-5061-4c0f-8f59-75b613a656df","resolution":{"observed_at":"2026-08-06T23:10:05.083897Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:04.872852Z","title":"Open3DIS: Open-vocabulary 3d instance segmentation with 2d mask guidance","venue":null,"work_id":"ff70e984-f0a8-4a75-aaaf-a3020727ca4a","year":2024},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:58.937755Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:08b925cfa39869e4f21d4116f9b70d31e4ccf5b0c9e169668a1e6ba921f353d0","observation_id":"b80f952f-c4ff-4700-bc31-e34a383fbbb0","resolution":{"observed_at":"2026-08-06T23:10:04.937474Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:04.700012Z","title":"Openscene: 3d scene understanding with open vocabularies","venue":null,"work_id":"132c2391-fdf1-4b01-a67d-c2b168c145b7","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.065667Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:8b62b6774e97427c6d6f3a47dd5d22d93723987998465c07ceb6224fa585784b","observation_id":"54e0be65-031a-428d-b501-4ab7ba43ba3c","resolution":{"observed_at":"2026-08-06T23:10:04.760947Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:04.527642Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":"38f20f1d-3522-4248-930e-e033ba4541da","year":2021},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.151631Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:51f4daa736077f467e128313ab800f4369b4294b6f960bb6dc6d8e6e118f9b58","observation_id":"580d45d1-5053-428b-9a46-e90994b107a6","resolution":{"observed_at":"2026-08-06T23:10:04.608034Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:04.326534Z","title":"Sam 2: Segment anything in images and videos","venue":null,"work_id":"6c8a49f3-4dc1-4464-80a4-d416ce962184","year":2024},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.261400Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:5b97a40978c385ce693bf31d6b14706e5c6a1541fc8d066b846a035a21f03f1e","observation_id":"3c6d0a51-f3d4-438f-a8ea-0a28108a10f8","resolution":{"observed_at":"2026-08-06T23:10:04.409658Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:04.052072Z","title":"Texture: Text-guided texturing of 3d shapes","venue":null,"work_id":"4e8706db-b44d-4df4-ad8a-a5911fa9f892","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.354109Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:c4de5d3d8aef9e667c5f29653c762857e5922b52ab5b2868005dd2140c6ee180","observation_id":"2f34b79e-06e6-4c4a-94b2-0051abc6e921","resolution":{"observed_at":"2026-08-06T23:10:04.190407Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:03.688343Z","title":"Language-grounded indoor 3d semantic segmentation in the wild","venue":null,"work_id":"cc8bbc5c-6244-4546-b76e-f0568d9fb9ed","year":2022},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.448232Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:adfba244ebbd3daa7e31e56d2e72216e5d4079ac218f51a5dbfcbcae20391f0e","observation_id":"e98c3e71-6447-4c58-a3e0-b7a3ca1f6330","resolution":{"observed_at":"2026-08-06T23:10:03.860724Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:03.333538Z","title":"Mask3d: Mask transformer for 3d semantic instance segmentation","venue":null,"work_id":"2cf5d079-b37f-42ed-860f-87bdb09c4dba","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.580812Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:920d786910e2f9073cd34f2cf13c579d4df4548411dacdd48b9a64398a9c254d","observation_id":"345809f0-3298-46aa-8a86-35098ff0e5fe","resolution":{"observed_at":"2026-08-06T23:10:03.533309Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:03.064435Z","title":null,"venue":null,"work_id":"b2d1f341-0b31-4782-ba39-d25c37459dc0","year":2019},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.679299Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:8f48a1d1dbbcdb685403d0ad22297134af1967a01d19f39642a61df1fccfafee","observation_id":"7e86ef37-9a08-4211-a1d7-cdb5b2d9bee6","resolution":{"observed_at":"2026-08-06T23:10:03.158906Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:09:59.750397Z","title":"Going denser with open-vocabulary part segmentation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.750397Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:256a090b27d95ce573102025ff56344d8baf1226f1d99fbeeeed653ef9e91d7c","observation_id":"55f46e5c-140a-468e-bdcd-aabd11ffccfb","resolution":{"observed_at":"2026-08-06T23:09:59.750397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:02.845630Z","title":"OpenMask3D: Open-vocabulary 3d instance segmentation","venue":null,"work_id":"2d669c52-a15a-4394-bb7f-dbd6e6ff6332","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.858547Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:885d49eaa46c275d4d97ed1c31729f7880339a9372a1a4c29381153177e4be77","observation_id":"2592679b-ca88-45d6-95e8-66661d24ac69","resolution":{"observed_at":"2026-08-06T23:10:02.936618Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:02.614759Z","title":"Search3d: Hierarchical open-vocabulary 3d segmentation","venue":null,"work_id":"97aedc7a-9eaf-4dde-b077-38ddbb5df909","year":2025},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.925603Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:6657c6d32a9eb4844220de0f6e0cdc7dfb9b04ee9a7e47aa616f032c375f3b74","observation_id":"6bf993da-3bda-4f88-9260-978fd75b06fc","resolution":{"observed_at":"2026-08-06T23:10:02.731010Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:02.370407Z","title":"Partdistill: 3d shape part segmentation by vision-language model distillation","venue":null,"work_id":"8fae10a8-51d2-4da8-adaf-226aeef4aea2","year":2024},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T23:09:59.995835Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:f63d0ecd5d0be85d940ddc7013f02946348dd6a21ec4463e31378aa3d591543b","observation_id":"e3869d9c-c99b-4076-b951-ba6216a0332d","resolution":{"observed_at":"2026-08-06T23:10:02.481935Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:02.180328Z","title":"Chang, Leonidas J","venue":null,"work_id":"8a200a4e-cee3-42e2-beae-9065b4d1a5ab","year":2020},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:00.114269Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:4adefbf744995bb1e0e1e6aaf6b1685b92ec062e0b57278ec24c97d3fce5ad1e","observation_id":"f05ac0e4-7d78-4d82-9a98-896d17bc2a99","resolution":{"observed_at":"2026-08-06T23:10:02.309690Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:02.035706Z","title":"Florence-2: Advancing a unified representation for a variety of vision tasks","venue":null,"work_id":"2492dbc4-e7f9-492a-bb3e-e833d7e27b8d","year":2024},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:00.191217Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:e7698a4614a25af87421e9b80bc8f3d3ed218c096e8f31156bc2171907d30b24","observation_id":"17a2f6dc-e69b-4315-8b5f-7785020f9c0f","resolution":{"observed_at":"2026-08-06T23:10:02.088956Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:01.866311Z","title":"Regionplc: Regional point-language contrastive learning for open-world 3d scene understanding","venue":null,"work_id":"2aa9e512-a3eb-47f1-84a8-151299fdbc3b","year":2024},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:00.313713Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:19d54c3383ed2c7b670073f9f94798be8c37a4b9163fcb64c04443d0542d8df9","observation_id":"4de23616-b85b-4ee4-b4df-24c416b01ba4","resolution":{"observed_at":"2026-08-06T23:10:01.968203Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:01.674438Z","title":"Scannet++: A high-fidelity dataset of 3d indoor scenes","venue":null,"work_id":"aa3dd261-cb0d-4c60-9b9c-de48eb0fe69f","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:00.374797Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:898534b65025b2ed2dfde5be4f2008a65b1aaf0008b3048c7f81ddb990adc480","observation_id":"a228de60-5945-4eca-b93a-244783672087","resolution":{"observed_at":"2026-08-06T23:10:01.734933Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:01.515245Z","title":"Sample elimination for generating poisson disk sample sets","venue":null,"work_id":"79610d29-abef-4f79-abf9-c44a548733cd","year":2015},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:00.466180Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:e17c1c76c3e5c1d0e802a1e55fea43730f08a7879d222c4cea4240e7f11e4997","observation_id":"3be96a80-ba32-48b7-9d35-68f0a84fcf50","resolution":{"observed_at":"2026-08-06T23:10:01.587981Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:01.334377Z","title":"Clip2: Contrastive language-image-point pretraining from real-world point cloud data","venue":null,"work_id":"935246c7-3e84-4b49-90b5-d4c6289b2e0c","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:00.557074Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:110ca33293a2f3f717b197d27ed51fe22c029d9f2d92c60fbbe5126d2601c66e","observation_id":"9d54488b-34b6-444b-83c5-a9cd65f5af31","resolution":{"observed_at":"2026-08-06T23:10:01.402629Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:01.100297Z","title":"A simple framework for open-vocabulary segmentation and detection","venue":null,"work_id":"0bed8e4a-4be9-4f71-bf30-54031f456d14","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:00.651883Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:fbb4eabc9a895cd50d11096e3b8b390e37f82d7e68940e8f8bfd956c9ff842de","observation_id":"47c179f1-f792-47ed-97f1-17ac7df5f486","resolution":{"observed_at":"2026-08-06T23:10:01.234291Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:00.900458Z","title":"Uni3d: Exploring unified 3d representation at scale","venue":null,"work_id":"057a91f7-e68a-4825-af90-93f976408fbf","year":2023},"citing_paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:00.733804Z"},"links":{"citing_paper":"/paper/2506.19331"},"observation_digest":"sha256:9577581541523ee828cb7434ecdf5eae0d9fb2188dbd6fb11fbf5690ae5b1e63","observation_id":"1610aca7-88e5-4811-836b-1028c3f47055","resolution":{"observed_at":"2026-08-06T23:10:00.978669Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.19331","last_updated":"2025-06-24T05:51:22Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-12T22:39:50.680724Z","submitted_at":"2025-06-24T05:51:22Z","title":"Segment Any 3D-Part in a Scene from a Sentence"},"reference_resolution":{"displayed":37,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":4,"verified_exact":0,"verified_fuzzy":33},"total_outbound_references":37},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 37 of 37 outbound references and 1 inbound Pith citation observation for arXiv:2506.19331."}