{"as_of":"2026-08-10T12:07:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:f30ac5609122d0ade78348bd994bee5761662d7dfda9c2a478d06af35d9769f0","coverage":[{"denominator":55,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":55,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T05:23:20.346385Z","state":"measured"},{"denominator":55,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":55,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2502.03266/citation-record","integrity":"/paper/2502.03266/integrity","json":"/paper/2502.03266/citation-record.json","paper":"/paper/2502.03266"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.534106Z","title":"Unseen object instance segmentation for robotic environments,","venue":null,"work_id":"6156523b-f0a7-4157-b375-50e0b8679f05","year":2021},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.112165Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:2fd00e1c56d95a2416ce9573bd8a4112c07871b9cb85ae193e55b985d6a71a85","observation_id":"6476cf82-976f-4d62-9052-2faffe87043e","resolution":{"observed_at":"2026-08-09T05:23:21.538866Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.748200Z","title":"Taylor neural network for unseen object instance segmentation in hierarchical grasping,","venue":null,"work_id":"32adba12-a56d-4ffd-9663-c367e9f42e5a","year":2024},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.117293Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:8f59c8c763f73a6dc3af138c9618d12381ae764bd1215d715fa4000528ac30c1","observation_id":"a9ff49f9-caeb-4e95-9fc9-6544cfc46715","resolution":{"observed_at":"2026-08-09T05:23:20.753148Z","resolver_source":"arxiv_id_nonexistent","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.520347Z","title":"Segmentation of unknown objects in indoor environments,","venue":null,"work_id":"54bb4855-9233-425b-a470-18fe2a69ab18","year":2012},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.122095Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:965e2137cc9ba5bf0c529e1ddc59bf9aa80ebc24636b3736a6a3cbc0691f44a7","observation_id":"0f5676cf-fa96-47d6-88de-31117bfd6235","resolution":{"observed_at":"2026-08-09T05:23:21.524604Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.505305Z","title":"Learn fast, segment well: Fast object segmentation learning on the icub robot,","venue":null,"work_id":"a5618987-dc33-4832-b689-5619267b26a3","year":2022},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.127034Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:ce3c456584253f917d9e829a6c3f974f7d8ab3f11a709740bdbff5cc0b33cc71","observation_id":"32bbd310-4b23-4ccb-bb53-50c4a2a28aea","resolution":{"observed_at":"2026-08-09T05:23:21.510313Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.489631Z","title":"Learning rgb-d feature embeddings for unseen object instance segmentation,","venue":null,"work_id":"1abeb93c-2714-429d-9af6-eef8f4af893d","year":2021},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.132220Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:a08c422a512504bcac6de794e3309b692191ba2822ce6989b0ddfddfd9358b9b","observation_id":"701e7110-bfa6-4cea-86c6-0562c7f722d7","resolution":{"observed_at":"2026-08-09T05:23:21.495522Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.474602Z","title":"Stow: Discrete-frame segmentation and tracking of unseen objects for warehouse picking robots,","venue":null,"work_id":"330f2372-bbec-4a71-a0df-b6c9d219b8e9","year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.137425Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:bd8f14b3291ad4599a46b32a92bedefdd7b8129ba6f39ad54139715556c3c254","observation_id":"89b0c793-5613-46c7-8530-513209a7a5ef","resolution":{"observed_at":"2026-08-09T05:23:21.479687Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.459150Z","title":"Mask r-cnn,","venue":null,"work_id":"2061056d-b34a-431c-b5d0-cc7209d0fe56","year":2017},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.142287Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:fde94624580b17483d19b8e33e7d233c3f724ac760b72b417fd469c584368f77","observation_id":"56fd16fb-ca0a-4e9f-b58b-c3df087b9fb4","resolution":{"observed_at":"2026-08-09T05:23:21.464135Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.443643Z","title":"Image segmentation using deep learning: A survey,","venue":null,"work_id":"48e98f57-715d-4084-bbf0-9dcaf70771f2","year":2022},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.146496Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:27e6954241f7338192a8874fd59ae67f72322b3b718840f0c52ecf6a6df02cf6","observation_id":"e91b3374-1d2b-40b6-a32b-7a6ae767de25","resolution":{"observed_at":"2026-08-09T05:23:21.448449Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.429236Z","title":"Imagenet large scale visual recognition challenge,","venue":null,"work_id":"73d803c0-5426-443f-b22e-537b58e3b18f","year":2015},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.150680Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:43f5723f3819e999e1d3dba77bedc2263db53b16894a1ce513007ccaec22220e","observation_id":"156047b9-bce1-4afe-898f-91216fbd73b9","resolution":{"observed_at":"2026-08-09T05:23:21.434253Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.414894Z","title":"Microsoft coco: Common objects in context,","venue":null,"work_id":"7e434783-01ba-436c-af2f-b5760d5d2711","year":2014},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.154763Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:e89f5341a8589cb7cd39f5d8d8aba6809cb540d500973d2a342f98008a1ef301","observation_id":"6b348b9e-beca-40f6-8e75-ab294a2a174b","resolution":{"observed_at":"2026-08-09T05:23:21.419748Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.400710Z","title":"Unseen object amodal instance segmentation via hierarchical occlusion mod- eling,","venue":null,"work_id":"543c85b9-03ba-4b12-bfdf-5010fd8c9504","year":2022},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.158973Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:309bd729b0bed4cc129daedd918064470fd35fa7e8b42cd835e428986df9a13b","observation_id":"7b88ad0c-6a85-4e6c-88c4-e3f8135b33c5","resolution":{"observed_at":"2026-08-09T05:23:21.405133Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.386890Z","title":"Unknown object segmentation from stereo images,","venue":null,"work_id":"7d41e3ea-6777-4efc-bb15-4e82b3730dfb","year":2021},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.162892Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:642be1c868d0a34ba5c02900d5c43634ea95b1edebd4ce6fe249c1e1f7d11053","observation_id":"1b422727-4637-4608-a5f0-b54a6d5076bd","resolution":{"observed_at":"2026-08-09T05:23:21.391631Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.372726Z","title":"The best of both modes: Separately leveraging rgb and depth for unseen object instance segmentation,","venue":null,"work_id":"8c3c13ce-13bc-46ff-ac6d-d693e6cc71ce","year":2020},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.167161Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:81a8fddaabe8b0fdce0e977a6860ac6744da97aa9425f66624a23ab25ba80dab","observation_id":"5288e27f-be1a-4548-8949-cd5d4a7df62c","resolution":{"observed_at":"2026-08-09T05:23:21.377280Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.358505Z","title":"Segmenting unknown 3d objects from real depth images using mask r-cnn trained on synthetic data,","venue":null,"work_id":"540ebade-562a-4059-b611-251fdf86e68e","year":2019},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.171495Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:c2ae939db83d3996ff0097ecf0253de687f7704131199ac984eaecaa042b29eb","observation_id":"e81d89b0-0161-428a-b49d-365064c0f7ed","resolution":{"observed_at":"2026-08-09T05:23:21.363133Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.344287Z","title":"Unseen object instance segmentation with fully test-time rgb-d embeddings adaptation,","venue":null,"work_id":"534763a4-f6fb-47f9-9dc4-86cb469787a1","year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.175694Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:3c8d2999808d7a8cc679d80afab0799fe403081fe12087106368b8ce9b601a59","observation_id":"3c3a61b3-9ab9-42fa-9fa4-0bdebabd04f9","resolution":{"observed_at":"2026-08-09T05:23:21.349197Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.03793","last_updated":"2023-02-07T23:11:29Z","snapshot_observed_at":"2026-07-06T14:49:30.554619Z","submitted_at":"2023-02-07T23:11:29Z","title":"Self-Supervised Unseen Object Instance Segmentation via Long-Term Robot Interaction","version":1},"cited_work":{"arxiv_id":"2302.03793","doi":null,"metadata_source":"pith","pith_arxiv_id":"2302.03793","snapshot_observed_at":"2026-08-09T05:23:20.505867Z","title":"Self-Supervised Unseen Object Instance Segmentation via Long-Term Robot Interaction","venue":"cs.RO","work_id":"317da509-e9d7-4e05-9333-67b2614f09e1","year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.179793Z"},"links":{"cited_paper":"/paper/2302.03793","citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:4f0780be3184dabb9f7b8dbaeba8c504b257bfd1ef63806c46f30f3d1ec5941e","observation_id":"54a555ba-77d1-47c8-8ee3-4fa17c30d32e","resolution":{"observed_at":"2026-08-09T05:23:20.511015Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.223938Z","title":"Self-supervised interactive object segmentation through a singulation-and-grasping approach,","venue":null,"work_id":"872078d1-60ee-43e0-93e4-bc5048132abd","year":2022},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.184496Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:41284ddb7d0b7154eb6c4312af8b8ed1dd51e3a269888450aa806e0e546ee0f1","observation_id":"e742de18-d91b-4578-ab4f-44137aa2c639","resolution":{"observed_at":"2026-08-09T05:23:21.228756Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.209572Z","title":"Self-supervised transfer learning for instance segmentation through physical interaction,","venue":null,"work_id":"d2c962fa-a097-42f3-bdf8-a476f413f277","year":2019},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.188506Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:c09cd0f8e19332667b4ecb4d871f2833d0276e4e9057d8225af9d34260139811","observation_id":"e7dcfb8a-91a4-45df-a579-bd74a435c361","resolution":{"observed_at":"2026-08-09T05:23:21.214304Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.02643","last_updated":"2023-04-05T17:59:46Z","snapshot_observed_at":"2026-08-08T05:14:59.435033Z","submitted_at":"2023-04-05T17:59:46Z","title":"Segment Anything","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.02643","snapshot_observed_at":"2026-08-09T05:23:20.192654Z","title":"Segment anything,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.192654Z"},"links":{"cited_paper":"/paper/2304.02643","citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:ff091a64ea330c0a321579931588654abf68ce73ce16a144d1175d7f0355a2c0","observation_id":"95507e09-a7f3-4c99-a672-9eb5ee5a684e","resolution":{"observed_at":"2026-08-09T05:23:20.192654Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.194468Z","title":"Seggpt: Segmenting everything in context,","venue":null,"work_id":"09c8f74e-cd23-46b5-8f18-c19d0bce87a5","year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.197147Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:d33d7c3347d933788fa5742a0c225668b53a3d18037dcb520f2f35925f29ccb1","observation_id":"c93e2336-db87-44bb-b734-2b53dbea134c","resolution":{"observed_at":"2026-08-09T05:23:21.199626Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.180775Z","title":"Semantic segment anything,","venue":null,"work_id":"7f3ea808-2c19-484d-826c-a35ec5d230e6","year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.201368Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:a0b5fbe63c64446d4a95a4f6381e5f6cbf0896debc2a47de82098af25afaa62c","observation_id":"be13fd24-5fac-462d-b587-dfac89a9a6a6","resolution":{"observed_at":"2026-08-09T05:23:21.185246Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.165737Z","title":"Segment everything everywhere all at once,","venue":null,"work_id":"c257fc15-7c26-4b3d-a4ca-a8fa0b02e4b4","year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.205463Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:7449e2174a557f08e5c2869481231f8aed2765b7955e72efc21224496b7f23a2","observation_id":"11f14516-8011-4a2d-a08c-4ad1cfb20a61","resolution":{"observed_at":"2026-08-09T05:23:21.170197Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.150398Z","title":"Matplotlib: A 2d graphics environment,","venue":null,"work_id":"d253a867-ece8-4892-ba50-ac3172be829e","year":2007},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.209557Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:6ab22d3fd03b2b4800b63d59cd5c1d906b72382cdfd604c226ec7b3775f31d9a","observation_id":"c3ebe5a3-47e2-4d2f-a81e-bcd11c7ebae1","resolution":{"observed_at":"2026-08-09T05:23:21.155592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.132743Z","title":"Dinov2: Learning robust visual features without supervision,","venue":null,"work_id":"05f30a2e-030d-4fa1-b4e6-4eed7da6e36d","year":2024},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.213572Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:4f1d5ad0cdb760eb3f4cb1616b9c984d3250a360b6e92f7293755e665b836d75","observation_id":"a6af7ca2-ad80-4888-aca4-828a58c583ca","resolution":{"observed_at":"2026-08-09T05:23:21.138516Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.113229Z","title":"Rice: Refining instance masks in cluttered environments with graph neural networks,","venue":null,"work_id":"3fb586eb-daeb-4f4a-a34c-c78699116a1c","year":2022},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.217667Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:3777e5784c9e4100d56788fe3145af0061a981a916e84768cc27d90041a9d11a","observation_id":"9c715dcf-c307-4f6e-852e-8d745e795e56","resolution":{"observed_at":"2026-08-09T05:23:21.119147Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17934","last_updated":"2024-09-29T05:56:47Z","snapshot_observed_at":"2026-08-01T18:35:31.455107Z","submitted_at":"2023-05-29T07:54:04Z","title":"ZeroPose: CAD-Prompted Zero-shot Object 6D Pose Estimation in Cluttered Scenes","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.17934","snapshot_observed_at":"2026-08-09T05:23:20.221643Z","title":"3d model-based zero-shot pose estimation pipeline,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.221643Z"},"links":{"cited_paper":"/paper/2305.17934","citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:354ab0a310c3ca1feab45aeab96b3f125ca00565bba99ca15ebb3466fb20ffe5","observation_id":"ca2cf670-a8f0-4e9e-ab94-0ceb2148eeac","resolution":{"observed_at":"2026-08-09T05:23:20.221643Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.093676Z","title":"Cnos: A strong baseline for cad-based novel object segmentation,","venue":null,"work_id":"c34b3030-649f-4747-bde8-b807d7053a66","year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.226107Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:651c2043005bd9ccff6afabb2031d2e2f3608d6a201f9d107b485f8f184d4f6d","observation_id":"8edf030d-e922-4308-acc7-59dba7b8c472","resolution":{"observed_at":"2026-08-09T05:23:21.100270Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.077223Z","title":"Panoptic seg- mentation,","venue":null,"work_id":"3aab29f3-5976-4472-820b-c0e5a641054a","year":2019},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.230228Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:a4518516128c0bd6612eb09a9dd489a5de200372f60cf6e78acb3f548c754cf3","observation_id":"b0e594ad-1030-4e50-b31b-b500dd096b6f","resolution":{"observed_at":"2026-08-09T05:23:21.082316Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.060902Z","title":"Yolact: Real-time instance segmentation,","venue":null,"work_id":"433fd321-7516-454b-8296-fd82d45df06e","year":2019},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.234472Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:7517b97220ee5dc3cf18d32e0b339e34607eb299e1525af54ffb0365b186a2fe","observation_id":"f36af7de-0bf3-46fc-838a-0d4000357252","resolution":{"observed_at":"2026-08-09T05:23:21.066464Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.238745Z","title":"Fully convolutional networks for semantic segmentation,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.238745Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:d660fc70357ee55179b2d881ac3b1c4643dae4b0a7db0718d8f6a5262a15731c","observation_id":"a1f3a908-cb16-4706-a202-47ff3b62f363","resolution":{"observed_at":"2026-08-09T05:23:20.238745Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.13721","last_updated":"2023-07-25T17:59:18Z","snapshot_observed_at":"2026-08-07T23:40:08.386706Z","submitted_at":"2023-07-25T17:59:18Z","title":"Foundational Models Defining a New Era in Vision: A Survey and Outlook","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.13721","snapshot_observed_at":"2026-08-09T05:23:20.242929Z","title":"Foundational models defining a new era in vision: A survey and outlook,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.242929Z"},"links":{"cited_paper":"/paper/2307.13721","citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:4fb80315d4dd3aca66e6237e63c5626c58f20cb4ac6bba3d1b55650032da3231","observation_id":"4518110c-7b4c-41a6-8550-d103d3dac445","resolution":{"observed_at":"2026-08-09T05:23:20.242929Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.09324","last_updated":"2023-05-05T18:58:52Z","snapshot_observed_at":"2026-08-09T21:11:38.511413Z","submitted_at":"2023-04-18T22:16:49Z","title":"Computer-Vision Benchmark Segment-Anything Model (SAM) in Medical Images: Accuracy in 12 Datasets","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.09324","snapshot_observed_at":"2026-08-09T05:23:20.248723Z","title":"Accuracy of segment- anything model (sam) in medical image segmentation tasks,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.248723Z"},"links":{"cited_paper":"/paper/2304.09324","citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:c581d6873b490577227969d70fe1f8b704b6e3e4dd2e640418607b04c85a5148","observation_id":"412b69be-84f8-43fe-aa63-2a92b85082fb","resolution":{"observed_at":"2026-08-09T05:23:20.248723Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.06558","last_updated":"2023-05-11T04:33:08Z","snapshot_observed_at":"2026-08-03T19:49:16.800693Z","submitted_at":"2023-05-11T04:33:08Z","title":"Segment and Track Anything","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.06558","snapshot_observed_at":"2026-08-09T05:23:20.253205Z","title":"Segment and track anything,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.253205Z"},"links":{"cited_paper":"/paper/2305.06558","citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:bef2e7a499270b569c7cc5e81155e61ac91598c949f881b25cafbf2d5dbd2d40","observation_id":"ed34e449-ddcb-432b-8bfb-19ae0bc3e577","resolution":{"observed_at":"2026-08-09T05:23:20.253205Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.06790","last_updated":"2023-04-13T19:23:52Z","snapshot_observed_at":"2026-08-05T23:54:40.289344Z","submitted_at":"2023-04-13T19:23:52Z","title":"Inpaint Anything: Segment Anything Meets Image Inpainting","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.06790","snapshot_observed_at":"2026-08-09T05:23:20.257724Z","title":"Inpaint anything: Segment anything meets image inpainting,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.257724Z"},"links":{"cited_paper":"/paper/2304.06790","citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:3c19e2115024beec4c981d06a4c5f226612050b09c1e4348dc11eaba81c19022","observation_id":"b7fe7a0e-b6f6-4d57-99a2-be9a911bc198","resolution":{"observed_at":"2026-08-09T05:23:20.257724Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.032862Z","title":"All-in-sam: from weak annotation to pixel-wise nuclei segmentation with prompt-based finetuning,","venue":null,"work_id":"6ebd55cb-e6fd-4718-810e-c5d250ad1293","year":2024},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.261966Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:2fe2a10dccdc385638782428907b16af0e92864f270e25a05234a2db22cb8c06","observation_id":"ecd83b3c-2fcb-4424-8d9d-316b72ab7518","resolution":{"observed_at":"2026-08-09T05:23:21.038696Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.12659","last_updated":"2025-07-08T12:38:02Z","snapshot_observed_at":"2026-07-06T15:30:26.626858Z","submitted_at":"2023-05-22T03:03:29Z","title":"UVOSAM: A Mask-free Paradigm for Unsupervised Video Object Segmentation via Segment Anything Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.12659","snapshot_observed_at":"2026-08-09T05:23:20.266079Z","title":"Uvosam: A mask- free paradigm for unsupervised video object segmentation via segment anything model,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.266079Z"},"links":{"cited_paper":"/paper/2305.12659","citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:21bd891e56628eb7755702192ac170b26a3b9d9e222ddb17c19a09ff380e90ff","observation_id":"f781c3d2-c427-4d55-9b9d-3fe9f2feb0d0","resolution":{"observed_at":"2026-08-09T05:23:20.266079Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.270649Z","title":"Attention is all you need,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.270649Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:5975a39455d99b158054f1ac048b2645fd04220acef45306e8fe4b0c30bfa964","observation_id":"be8d51ea-8808-441f-a96f-bad992dbb5f5","resolution":{"observed_at":"2026-08-09T05:23:20.270649Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:21.005454Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale,","venue":null,"work_id":"d6162a2b-9fbc-470e-82b3-435f3a21baee","year":2021},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.274554Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:4b27d66330663c71cd2e7d49ac846df1742435641798f99bdfad361a18ad40cf","observation_id":"4269a717-34a4-45f6-93c1-24e7d34fbf3e","resolution":{"observed_at":"2026-08-09T05:23:21.010904Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.989018Z","title":"Cross-level multi-modal features learning with transformer for rgb-d object recognition,","venue":null,"work_id":"280d7fd5-6b4d-4883-9719-64f902aebc27","year":2023},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.278667Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:cca5c2953b2b189dda5325f6459e43545b7c1147726443bfc319f026aae2fee7","observation_id":"20554afa-5199-4459-9c72-eaa5c52757e3","resolution":{"observed_at":"2026-08-09T05:23:20.994413Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.970355Z","title":"An empirical study of training self- supervised vision transformers,","venue":null,"work_id":"949b73db-cd55-4b16-86b5-ddeea8144f34","year":2021},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.282805Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:2110acd10ee8a2613b0f084f211be42796f6e4e9f5b2c805adcbe6a885a7004d","observation_id":"c9031c0f-ebf8-40a0-8268-ec954db31cc9","resolution":{"observed_at":"2026-08-09T05:23:20.975933Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.286973Z","title":"Emerging properties in self-supervised vision transformers,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.286973Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:ad8de9e7bcb41d704e03fed4f12ade90f0e006ce4acaaeb96890076e67e8a2a9","observation_id":"28042dbc-bfda-4ca7-80c6-b00ac6c5133a","resolution":{"observed_at":"2026-08-09T05:23:20.286973Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.942315Z","title":"Mst: Masked self-supervised transformer for visual representation,","venue":null,"work_id":"c1892a38-73c6-4e9a-ae71-476e49ceedc0","year":2021},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.291068Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:f2149ebacfe9b60a5de79962b4ff4f2e5729ecbd7118a91693fa172900b30658","observation_id":"a832af2f-8466-4a1f-8945-73f18be8d8ff","resolution":{"observed_at":"2026-08-09T05:23:20.947889Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.920828Z","title":"Bert: Pre-training of deep bidirectional transformers for language understanding,","venue":null,"work_id":"a717eaf9-fad5-4c79-af5a-075eaff3bef3","year":2019},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.295130Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:437eca6adb340bc48ea94d59e78bb7c3d0f45f0a6a2e4d894645580ed3908e28","observation_id":"332eaa78-d809-4aff-b197-324ab4bcbaca","resolution":{"observed_at":"2026-08-09T05:23:20.925421Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.905334Z","title":"Beit: Bert pre-training of image transformers,","venue":null,"work_id":"32e395aa-24e2-4c56-a5b3-ef5a2f80bb2e","year":2022},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.299376Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:d04e8c223ef158aa059ba945f39a12e4d7ab83b8d54f670758ad04c91fd60ffb","observation_id":"38131734-cc85-4094-b2e9-d47bfabdcf96","resolution":{"observed_at":"2026-08-09T05:23:20.910615Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.889421Z","title":"Masked au- toencoders are scalable vision learners,","venue":null,"work_id":"c264d8b9-cf06-405d-bad7-09eaaf5b8d5a","year":2022},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.303847Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:93955687f52b1802ade73d41901eb0048c74bebffff5c038827c95a1c3354f1c","observation_id":"b0c9f7ec-a1b4-45c0-9037-355706538299","resolution":{"observed_at":"2026-08-09T05:23:20.893921Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.874168Z","title":"Deep vit features as dense visual descriptors,","venue":null,"work_id":"315a9a3d-09e2-476c-8ce3-445630d90859","year":2022},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.308105Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:af5d046c16e24a6286115623a8383ffccc5d709f8656bf9ddba7e5d02df61d21","observation_id":"b88df0a6-50b2-486c-b9a4-a3bba402e3ec","resolution":{"observed_at":"2026-08-09T05:23:20.878823Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.859294Z","title":"Easylabel: A semi- automatic pixel-wise object annotation tool for creating robotic rgb-d datasets,","venue":null,"work_id":"2c3d0bf7-726c-45b2-9a80-9358c53562f2","year":2019},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.312262Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:eff56f8a93fb7f50a6cb242bca62a05e3baf9473d40f12dce7be05fc538426d8","observation_id":"7f9315ac-fb86-49a7-a18d-0a36d066e887","resolution":{"observed_at":"2026-08-09T05:23:20.864158Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.844566Z","title":"A simple and fast algorithm for k-medoids clustering,","venue":null,"work_id":"03b9b005-c10e-4b74-b115-6a43e299ef82","year":2009},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.316393Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:e824d385e25f4fa8f55b9016b312b49cb07617a35080c35951c958c48bf43a17","observation_id":"94527a10-78ad-427a-a527-c295b525417a","resolution":{"observed_at":"2026-08-09T05:23:20.849368Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.829881Z","title":"Vision transformers need registers,","venue":null,"work_id":"411de786-2669-4085-9ab5-a559d734479c","year":2024},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.320617Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:65180a4cd5991b6fb10a1373115dbc83792438cf6d5f53b28822ce8a78661d1b","observation_id":"25f14822-4517-42fc-a867-bf4e7cc5dec9","resolution":{"observed_at":"2026-08-09T05:23:20.834488Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.325001Z","title":"Pytorch: An imperative style, high-performance deep learning library,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.325001Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:83a8c463cf9d17ed0b54f40625c4c4277376c056dca3de7cc52edb4eda46e40c","observation_id":"7095e00d-a8d3-42ff-a15f-5cc52d5bed30","resolution":{"observed_at":"2026-08-09T05:23:20.325001Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.805317Z","title":"Towards segmenting anything that moves,","venue":null,"work_id":"28997b37-c3f5-4662-9a52-d7de98004926","year":2019},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.329254Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:1e04a1f2589b3a4b4b96d26dad3832bdaf41956987a166d50f92ece503352054","observation_id":"e7774d24-fdbe-4001-8428-e53298ab36dc","resolution":{"observed_at":"2026-08-09T05:23:20.810502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.11679","last_updated":"2023-09-21T23:04:42Z","snapshot_observed_at":"2026-07-06T14:21:18.799177Z","submitted_at":"2022-11-21T17:47:48Z","title":"Mean Shift Mask Transformer for Unseen Object Instance Segmentation","version":3},"cited_work":{"arxiv_id":"2211.11679","doi":null,"metadata_source":"pith","pith_arxiv_id":"2211.11679","snapshot_observed_at":"2026-08-09T05:23:20.381508Z","title":"Mean Shift Mask Transformer for Unseen Object Instance Segmentation","venue":"cs.CV","work_id":"eb19d770-3875-461a-8a5f-3053e33bb171","year":2022},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.333378Z"},"links":{"cited_paper":"/paper/2211.11679","citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:0a19be691b8ad3923701a88bd5da1ee2a18ae6301da905ebd1f96d3a622cb8a9","observation_id":"59362ca4-dd32-4592-90de-a181d5af94f6","resolution":{"observed_at":"2026-08-09T05:23:20.388116Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.337953Z","title":"Contact- graspnet: Efficient 6-dof grasp generation in cluttered scenes,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.337953Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:67b87e557ee5570c651d644f7ed3426f3f3a7fde9b0f3476e2a340abec57d5c1","observation_id":"c42e1cbc-0e51-4ebc-9d04-84dfdf1de024","resolution":{"observed_at":"2026-08-09T05:23:20.337953Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.781188Z","title":"Safe and efficient robot manipulation: Task-oriented environment modeling and object pose estimation,","venue":null,"work_id":"00efa7b2-5f99-4500-af63-78a77e47308b","year":2021},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.342220Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:9482d7a68a31df44c44504b08a119016b5a473d47a6238648c49774b45f057cc","observation_id":"ebbd8654-fb32-41b3-a956-c6a916b427b4","resolution":{"observed_at":"2026-08-09T05:23:20.785710Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T05:23:20.764758Z","title":"Moveit![ros topics],","venue":null,"work_id":"ec1ac014-22ae-49df-99bc-4923bde29a5b","year":2012},"citing_paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-09T05:23:20.346385Z"},"links":{"citing_paper":"/paper/2502.03266"},"observation_digest":"sha256:4733bd7470552cd838a81feafd4573bece382f5ceb6990a2e030345494e7efbc","observation_id":"2ca36434-fd16-4e2a-a0e8-49436e369490","resolution":{"observed_at":"2026-08-09T05:23:20.769931Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.03266","last_updated":"2025-02-05T15:22:20Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-09T21:11:58.038369Z","submitted_at":"2025-02-05T15:22:20Z","title":"ZISVFM: Zero-Shot Object Instance Segmentation in Indoor Robotic Environments with Vision Foundation Models"},"reference_resolution":{"displayed":55,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":12,"verified_exact":3,"verified_fuzzy":40},"total_outbound_references":55},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 55 of 55 outbound references and 0 inbound Pith citation observations for arXiv:2502.03266."}