{"as_of":"2026-08-14T00:33:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:457376b69eda9d1615c7afa2d55225459d8c2e10fe7bf05efa825de2a057c476","coverage":[{"denominator":39,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":39,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T10:29:14.944575Z","state":"measured"},{"denominator":39,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":39,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2502.02972/citation-record","integrity":"/paper/2502.02972/integrity","json":"/paper/2502.02972/citation-record.json","paper":"/paper/2502.02972"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.528549Z","title":"Semi-supervised active learning for semantic segmentation in unknown environments using informative path planning,","venue":null,"work_id":"7dd2e585-8bff-4c43-8dac-758f859b6e10","year":2024},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.760092Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:90f0a3725902e2454032a36bcf82ba4cd0bf28486625323af57e3bb6b342180f","observation_id":"dfcc9967-2229-420c-a2aa-b3d884b7c84a","resolution":{"observed_at":"2026-08-09T10:29:15.533681Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.513044Z","title":"Lightweight semantic segmentation network for semantic scene understanding on low-compute devices,","venue":null,"work_id":"e6a05029-44f9-4017-8f80-1d78971ebba5","year":2023},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.765425Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:c185abeb649bfb14ff741d22d2cae56781ed726159fc45bbff1e4d3d48169765","observation_id":"3c7e04f9-2119-47ae-8912-7f9ce6aa8898","resolution":{"observed_at":"2026-08-09T10:29:15.518099Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.01103","last_updated":"2024-07-01T09:10:35Z","snapshot_observed_at":"2026-08-12T23:31:28.633016Z","submitted_at":"2024-07-01T09:10:35Z","title":"FedRC: A Rapid-Converged Hierarchical Federated Learning Framework in Street Scene Semantic Understanding","version":1},"cited_work":{"arxiv_id":"2407.01103","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.01103","snapshot_observed_at":"2026-08-09T10:29:15.146833Z","title":"FedRC: A Rapid-Converged Hierarchical Federated Learning Framework in Street Scene Semantic Understanding","venue":"cs.RO","work_id":"986cf820-362a-40eb-beb4-cc6fa362da4c","year":2024},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.770160Z"},"links":{"cited_paper":"/paper/2407.01103","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:051b68f2c3814073471a1e48147ea94c9d5f7e6d469eee2e1b4486bdbc23fc65","observation_id":"2b0823e9-3e45-4094-a31d-08ba3676e1ad","resolution":{"observed_at":"2026-08-09T10:29:15.152199Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.497374Z","title":"Motionsc: Data set and network for real- time semantic mapping in dynamic environments,","venue":null,"work_id":"9d0aeb5f-e76a-43f2-8ae0-fab509f1758b","year":2022},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.775541Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:ef8f2a57689d4487fc505ef7e2e58602ffe806bb375239e1e5dcd4707a13ba67","observation_id":"31c5ce28-1be1-4fa0-b286-51057cd676da","resolution":{"observed_at":"2026-08-09T10:29:15.502477Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.481124Z","title":"pfedlvm: A large vision model (lvm)-driven and latent feature-based personalized federated learning framework in autonomous driving,","venue":null,"work_id":"3d16f983-dcdc-4707-bb81-4e97125db83c","year":2024},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.780977Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:9550ba8c827ece1ab4a26cf7edfe760f7d1d14c02bb520d941ca7e5a995efd26","observation_id":"7dfd1de9-5b93-45ad-9911-a7fd31fdf862","resolution":{"observed_at":"2026-08-09T10:29:15.485966Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.465783Z","title":"Cekd: Cross-modal edge-privileged knowledge distillation for semantic scene understanding using only thermal images,","venue":null,"work_id":"18886212-b941-44bb-9207-8ac4ebe31b83","year":2023},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.786534Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:5a2b17dcb6565d83ea5c406f600dc39bd548b526382283b21e4858d2a24d1441","observation_id":"1a2ab054-1b63-4332-8124-fb69ab683570","resolution":{"observed_at":"2026-08-09T10:29:15.471006Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.01710","last_updated":"2025-05-30T23:06:21Z","snapshot_observed_at":"2026-08-12T08:56:55.607706Z","submitted_at":"2025-01-03T09:10:56Z","title":"Enhancing Large Vision Model in Street Scene Semantic Understanding through Leveraging Posterior Optimization Trajectory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.01710","snapshot_observed_at":"2026-08-09T10:29:14.792526Z","title":"Enhancing large vision model in street scene semantic understanding through leveraging posterior optimization trajectory,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.792526Z"},"links":{"cited_paper":"/paper/2501.01710","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:d006fe04fd57478fa55e9a371bd3f7dbdc1776cd3badb905a23759b5303c286f","observation_id":"ecf2b56e-24da-421a-b0c4-f4744d7dfd36","resolution":{"observed_at":"2026-08-09T10:29:14.792526Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.450170Z","title":"Under- standing bird’s-eye view of road semantics using an onboard camera,","venue":null,"work_id":"660f1724-48b3-477d-8857-7b9fe6b3c3b3","year":2022},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.798286Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:24eff95ef752e55d81a8c374e1e166d8e897cd76c95d68bd6e87b31bfffe0cd9","observation_id":"b5b39902-b2ab-4115-a948-cd5220846623","resolution":{"observed_at":"2026-08-09T10:29:15.455282Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.14737","last_updated":"2025-01-09T07:09:44Z","snapshot_observed_at":"2026-08-12T22:39:52.006599Z","submitted_at":"2024-09-23T06:33:52Z","title":"Generalizable Autonomous Driving System across Diverse Adverse Weather Conditions","version":3},"cited_work":{"arxiv_id":"2409.14737","doi":null,"metadata_source":"pith","pith_arxiv_id":"2409.14737","snapshot_observed_at":"2026-08-09T10:29:15.109384Z","title":"Generalizable Autonomous Driving System across Diverse Adverse Weather Conditions","venue":"cs.RO","work_id":"c4230674-c4da-48fb-95e2-198ec7a34e81","year":2024},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.802924Z"},"links":{"cited_paper":"/paper/2409.14737","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:7fd2e2f617e52c73ab719c539f3d601b386b44522f1ab7bcd2b55644fd01008f","observation_id":"56c9d69a-cda8-4e98-9dfb-bc48b4ae18b8","resolution":{"observed_at":"2026-08-09T10:29:15.114985Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:14.808342Z","title":"Towards compact autonomous driving perception with balanced learning and multi-sensor fusion,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.808342Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:87ef4240b2975468d3abbe230b3f574521669239aad0b0bd4fca081899dd0a4a","observation_id":"59db7d5d-5739-4283-a677-5e15da736ee7","resolution":{"observed_at":"2026-08-09T10:29:14.808342Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.19560","last_updated":"2024-09-29T05:27:40Z","snapshot_observed_at":"2026-08-12T22:35:17.326329Z","submitted_at":"2024-09-29T05:27:40Z","title":"Fast-Convergent and Communication-Alleviated Heterogeneous Hierarchical Federated Learning in Autonomous Driving","version":1},"cited_work":{"arxiv_id":"2409.19560","doi":null,"metadata_source":"pith","pith_arxiv_id":"2409.19560","snapshot_observed_at":"2026-08-09T10:29:15.085114Z","title":"Fast-Convergent and Communication-Alleviated Heterogeneous Hierarchical Federated Learning in Autonomous Driving","venue":"cs.LG","work_id":"54370f14-eca9-4590-863a-4954c6a0eb74","year":2024},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.813578Z"},"links":{"cited_paper":"/paper/2409.19560","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:1c35f56816e39e8fc0f543c1c939bb16b5826f203425b9288cd7c47131f22bc5","observation_id":"a1ae3bbf-5b05-46ac-a95c-8dc900a48cb4","resolution":{"observed_at":"2026-08-09T10:29:15.092727Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.02643","last_updated":"2023-04-05T17:59:46Z","snapshot_observed_at":"2026-08-08T05:14:59.435033Z","submitted_at":"2023-04-05T17:59:46Z","title":"Segment Anything","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.02643","snapshot_observed_at":"2026-08-09T10:29:14.818947Z","title":"Segment anything,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.818947Z"},"links":{"cited_paper":"/paper/2304.02643","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:08feec99a98cf50ae04b95bb9509ac16588bdbf5faa683fb390e95ab065ecb6e","observation_id":"1f14fe9d-301d-4889-b7d0-8161407ca57b","resolution":{"observed_at":"2026-08-09T10:29:14.818947Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-13T14:19:26.598265Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-09T10:29:14.824130Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale,","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.824130Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:af1c337974bb5dfae4c2935fdefb64f02f3445f8e6354497b97cf627758a107c","observation_id":"e45f044e-102c-4da8-b970-00199173b4f5","resolution":{"observed_at":"2026-08-09T10:29:14.824130Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.425303Z","title":"Carla: An open urban driving simulator,","venue":null,"work_id":"1178ca2a-7982-456c-b320-bcb21f85e6cb","year":2017},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.828648Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:abffeb5b19b1511559d1b7d7837684b86b0f2de54e59448b475643b8c339b637","observation_id":"1c83a78c-839c-4b0e-bc4c-0ba2c9b7478a","resolution":{"observed_at":"2026-08-09T10:29:15.430102Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:14.832984Z","title":"Scaling up visual and vision-language representation learning with noisy text supervision,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.832984Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:d6fb33cbeac52109b71ff5997761f748847259f5092293bedd1c73b55d1d9941","observation_id":"a6b6c010-a191-4af1-b573-9817a052bf61","resolution":{"observed_at":"2026-08-09T10:29:14.832984Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:14.837873Z","title":"Learning transferable visual models from natural language supervision,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.837873Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:0bf571c06c7335c093c462135e028aa4c69f654953d59f52431343ba9ba7ccdb","observation_id":"1c9de113-a94b-4c79-a4fb-8c2513e0f498","resolution":{"observed_at":"2026-08-09T10:29:14.837873Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.389332Z","title":"Vilt: Vision-and-language transformer without convolution or region supervision,","venue":null,"work_id":"915d0fb3-ffa8-4987-9eda-b03286669b40","year":2021},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.842396Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:6767910eb4b52c4ab62bbe191fc6b10ece0fef78ffe16335c1df2e7384893e65","observation_id":"9bf32a30-4f75-480d-ad0c-657053fc37e5","resolution":{"observed_at":"2026-08-09T10:29:15.394788Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.373984Z","title":"Vlmo: Unified vision-language pre- training with mixture-of-modality-experts,","venue":null,"work_id":"5d563ac3-1fb7-41fb-bf43-fc1853cf61d2","year":2022},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.847046Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:6f42a7c5afa1f5c54aac4f09e067ac069bc73af30a5588bb4bfcc8a8cd77f0b9","observation_id":"5a422b35-3bfd-4b42-a604-4e0d789d4550","resolution":{"observed_at":"2026-08-09T10:29:15.379020Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:14.851509Z","title":"Blip: Bootstrapping language- image pre-training for unified vision-language understanding and gen- eration,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.851509Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:c6dda26f6099b55c54e62cedb024324bbabaed258b48905859d9b1273360516c","observation_id":"e47951c7-a150-458b-80ec-3c6b47095261","resolution":{"observed_at":"2026-08-09T10:29:14.851509Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:14.856350Z","title":"Lit: Zero-shot transfer with locked-image text tuning,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.856350Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:09e290312eb353badb56106786432693608d7dcdbc5665800c0ea23ff403c787","observation_id":"7295529f-95a3-409e-929e-4bd76fe6c1c5","resolution":{"observed_at":"2026-08-09T10:29:14.856350Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.12597","last_updated":"2023-06-15T07:57:29Z","snapshot_observed_at":"2026-08-12T12:01:54.105712Z","submitted_at":"2023-01-30T00:56:51Z","title":"BLIP-2: Bootstrapping Language-Image Pre-training with Frozen Image Encoders and Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.12597","snapshot_observed_at":"2026-08-09T10:29:14.860765Z","title":"Blip-2: Bootstrapping language- image pre-training with frozen image encoders and large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.860765Z"},"links":{"cited_paper":"/paper/2301.12597","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:24a81c817e30310bc24bcd9f7a0f49a444d910d6435ffbb0a10effb61719dda7","observation_id":"ba576611-5226-4ae2-ac76-daa16684796c","resolution":{"observed_at":"2026-08-09T10:29:14.860765Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.01313","last_updated":"2024-01-08T16:19:17Z","snapshot_observed_at":"2026-07-06T17:10:56.398607Z","submitted_at":"2024-01-02T17:56:30Z","title":"A Comprehensive Survey of Hallucination Mitigation Techniques in Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.01313","snapshot_observed_at":"2026-08-09T10:29:14.865485Z","title":"A comprehensive survey of hallucination mitigation tech- niques in large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.865485Z"},"links":{"cited_paper":"/paper/2401.01313","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:5e782ed449ab181dfa2ca363a361371a98fd3b8825a5da1a71ae15e8d2d12eda","observation_id":"2af0d4f8-21dd-44ce-820a-0304ae942299","resolution":{"observed_at":"2026-08-09T10:29:14.865485Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:14.870799Z","title":"Graspgpt: Leveraging semantic knowledge from a large language model for task- oriented grasping,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.870799Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:bc738e145847b67b0ca9ffe12aec1f0396a5650d563224685696dfe17b3aaaab","observation_id":"1fa3abd3-2c8c-452b-b164-0ef455583eb1","resolution":{"observed_at":"2026-08-09T10:29:14.870799Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.329421Z","title":"Zero-shot open-vocabulary tracking with large pre- trained models,","venue":null,"work_id":"8edc3675-c6f3-47e5-a14d-3c1d5021bccc","year":2024},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.875246Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:ab4e68750711e7c1bf4a76ab80236e193964ac10709f96279fd8bd30e199d17e","observation_id":"e8b5b47d-e035-4972-a510-f42ced572335","resolution":{"observed_at":"2026-08-09T10:29:15.335181Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.313642Z","title":"Prompt, plan, perform: Llm-based humanoid control via quantized imitation learning,","venue":null,"work_id":"b440ab17-23fd-43b5-9806-5668a95adb61","year":2024},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.879825Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:005e70f30410893db1a1f25cd0685fb3c2c9d005ad2b45ab8f8d6809353df876","observation_id":"a1fdfb13-62af-42b7-b6ac-0dd7c7e8695a","resolution":{"observed_at":"2026-08-09T10:29:15.318407Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.15012","last_updated":"2024-10-08T16:30:11Z","snapshot_observed_at":"2026-08-12T23:59:24.071076Z","submitted_at":"2024-05-23T19:35:03Z","title":"Extracting Prompts by Inverting LLM Outputs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.15012","snapshot_observed_at":"2026-08-09T10:29:14.884066Z","title":"Extracting prompts by inverting llm outputs,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.884066Z"},"links":{"cited_paper":"/paper/2405.15012","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:0784d1513627c4c1c467936d34ab5698b0471cec05d0715672adaa5aa85196bf","observation_id":"9db8788d-3c96-4788-b889-1e0816698f6a","resolution":{"observed_at":"2026-08-09T10:29:14.884066Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10382","last_updated":"2024-10-06T21:19:20Z","snapshot_observed_at":"2026-08-12T23:42:24.514769Z","submitted_at":"2024-06-14T19:24:00Z","title":"Efficient Prompting for LLM-based Generative Internet of Things","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10382","snapshot_observed_at":"2026-08-09T10:29:14.888993Z","title":"Efficient prompting for llm-based generative internet of things,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.888993Z"},"links":{"cited_paper":"/paper/2406.10382","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:a4df34756c328a8d62b6035c412e5fb68a78ec0554e71b4db674c1cbc6b1978d","observation_id":"0e8643b5-d7f3-4526-ab36-b143d6e49b8d","resolution":{"observed_at":"2026-08-09T10:29:14.888993Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.298733Z","title":"A simple zero-shot prompt weighting technique to improve prompt ensembling in text-image models,","venue":null,"work_id":"c74d57df-586a-49bb-8a72-d4d0cb67181a","year":2023},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.894207Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:751cbba24b97bcd3289afb9a41beaac5de33976551596c304636581c30a3cf2e","observation_id":"d342d9ab-7d62-4f68-ab05-7b3cc407cdcc","resolution":{"observed_at":"2026-08-09T10:29:15.303502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2104.08786","last_updated":"2022-03-03T12:10:58Z","snapshot_observed_at":"2026-07-06T11:01:05.577957Z","submitted_at":"2021-04-18T09:29:16Z","title":"Fantastically Ordered Prompts and Where to Find Them: Overcoming Few-Shot Prompt Order Sensitivity","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.08786","snapshot_observed_at":"2026-08-09T10:29:14.898680Z","title":"Fantasti- cally ordered prompts and where to find them: Overcoming few-shot prompt order sensitivity,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.898680Z"},"links":{"cited_paper":"/paper/2104.08786","citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:9f7f7f23e892ffe7374d333143902deabdb5319dc18dc1c2cf15af34513f7770","observation_id":"1e2df27d-6e39-4457-af29-dda49117160f","resolution":{"observed_at":"2026-08-09T10:29:14.898680Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.283440Z","title":"Design guidelines for prompt engineering text-to-image generative models,","venue":null,"work_id":"1bae40b7-b556-4e10-8ae9-562fab4495c4","year":2022},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.904154Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:45ff28c8fc589e8abe598e1be2009fd45edfe4e6a2c251b37a84da1363baa6e8","observation_id":"bbe02eb5-f5ad-4ff8-aed1-24a6a67252b0","resolution":{"observed_at":"2026-08-09T10:29:15.288189Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:14.908652Z","title":"The cityscapes dataset for semantic urban scene understanding,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.908652Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:1de0aa547972f1944b74cf4f5dfe00ddd6d562655752950d35428b0a23cc0201","observation_id":"08b173d0-42b2-47fc-80dc-f3fe37245cb1","resolution":{"observed_at":"2026-08-09T10:29:14.908652Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.258233Z","title":"Segmentation and recognition using structure from motion point clouds,","venue":null,"work_id":"b8e189fc-131c-4926-98ed-8012ee435fb1","year":2008},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.913132Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:fc0562f285e8ba7468a552173f3640e3492089e048c855e49cf8ab872fabf07b","observation_id":"65fb5263-df78-4547-a74b-cb3a893da0a1","resolution":{"observed_at":"2026-08-09T10:29:15.263843Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.243046Z","title":null,"venue":null,"work_id":"ebd87d99-ff15-4830-9ae1-7c019aafd667","year":2008},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.917689Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:52a0c50c503cc2f13681445f197d0427cea12a0454ea5d358a3cebf24c5468ad","observation_id":"2e247956-cf90-4bfd-9338-fd6eadb1a552","resolution":{"observed_at":"2026-08-09T10:29:15.247762Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.227873Z","title":"The apolloscape open dataset for autonomous driving and its application,","venue":null,"work_id":"87522734-cbe4-4574-83b4-a494188d2083","year":2019},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.922160Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:1befd773a148666d5ad84fb78466306c51ea41177d7cebd8e42c3b248f04e508","observation_id":"10199100-27ca-4efb-86da-809cbdf2abd0","resolution":{"observed_at":"2026-08-09T10:29:15.232777Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.212373Z","title":"Bisenet v2: Bilateral network with guided aggregation for real-time semantic segmentation,","venue":null,"work_id":"46f739ae-62a0-4599-9b4f-68e25dc14058","year":2021},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.926421Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:869845f5d696e79a7b32f5aabe5603b0ecf0caf3764532a9a0e96e852de30420","observation_id":"ba85a5f6-9296-4027-9f00-005e511e0a37","resolution":{"observed_at":"2026-08-09T10:29:15.217132Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:14.931254Z","title":"Segnet: A deep convolutional encoder-decoder architecture for image segmentation,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.931254Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:70a2f5a42c2c876649e8462b8e01064d88c1ac0662936bc70a1db563f4343604","observation_id":"252f2cb5-bc94-4ade-b462-1b3184fac473","resolution":{"observed_at":"2026-08-09T10:29:14.931254Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.187297Z","title":"Encoder-decoder with atrous separable convolution for semantic image segmentation,","venue":null,"work_id":"1eee9124-1231-42af-8b28-8d9caf0c1e5d","year":2018},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.935380Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:c13a76f8e8636e102a232112df0540bb573e25c06832280f214583fd541b76a0","observation_id":"141bd968-ae3f-4167-89e5-1e1e85a9e6d4","resolution":{"observed_at":"2026-08-09T10:29:15.192146Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:15.172360Z","title":"Segformer: Simple and efficient design for semantic segmen- tation with transformers,","venue":null,"work_id":"d3892ac4-958f-4c4e-ba5a-015f042d2e7b","year":2021},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.940207Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:552d5e7a26efc62f942823e0c35e9ea0f81b13e542c413c2c8d80f5d57611e9d","observation_id":"e21c183a-37de-4476-8935-089d9c3b7e72","resolution":{"observed_at":"2026-08-09T10:29:15.177218Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:29:14.944575Z","title":"Deeplab: Semantic image segmentation with deep convolutional nets, atrous convolution, and fully connected crfs,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-09T10:29:14.944575Z"},"links":{"citing_paper":"/paper/2502.02972"},"observation_digest":"sha256:7b5711a02be76ada74cd7420578055844f08d9db76613510745be555918e206a","observation_id":"403c0e97-dcad-40c2-b37f-3792fdb16de8","resolution":{"observed_at":"2026-08-09T10:29:14.944575Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2502.02972","last_updated":"2025-02-05T08:14:52Z","latest_version":1,"primary_category":"cs.RO","snapshot_observed_at":"2026-08-12T08:55:42.824309Z","submitted_at":"2025-02-05T08:14:52Z","title":"Label Anything: An Interpretable, High-Fidelity and Prompt-Free Annotator"},"reference_resolution":{"displayed":39,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":18,"verified_exact":2,"verified_fuzzy":18},"total_outbound_references":39},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 39 of 39 outbound references and 0 inbound Pith citation observations for arXiv:2502.02972."}