{"as_of":"2026-08-11T00:07:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:561b26c26875ea77a78ed31384d56c8ce840e44ee5693620c99a931920159ab1","coverage":[{"denominator":22,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":22,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T14:27:14.180633Z","state":"measured"},{"denominator":23,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":23,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-27T03:40:26.565029Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-03T17:48:46.256458Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"cited_work":{"arxiv_id":"2501.15326","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.15326","snapshot_observed_at":"2026-07-03T17:48:46.256458Z","title":"Recognize any surgical object: Unleashing the power of weakly- supervised data.arXiv preprint arXiv:2501.15326, 2025","venue":null,"work_id":"5148ed58-4ca5-4d73-934e-a8cd700d3908","year":2025},"citing_paper":{"arxiv_id":"2606.17279","last_updated":"2026-06-15T20:40:45Z","snapshot_observed_at":"2026-08-06T02:55:29.523710Z","submitted_at":"2026-06-15T20:40:45Z","title":"Training LLMs with Reinforcement Learning over Digital Twin Representations for Reasoning-Intensive Surgical VideoQA","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-27T03:40:26.565029Z"},"links":{"cited_paper":"/paper/2501.15326","citing_paper":"/paper/2606.17279"},"observation_digest":"sha256:c89cd97a0935c49f7c4df04daf91f6a1061eb54238969dc7125e9159521804d5","observation_id":"4e9ccbeb-ff57-4cd9-b07c-a01f9202be7d","resolution":{"observed_at":"2026-07-03T17:48:46.258354Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2501.15326/citation-record","integrity":"/paper/2501.15326/integrity","json":"/paper/2501.15326/citation-record.json","paper":"/paper/2501.15326"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T14:27:14.489532Z","title":"Whisperx: Time-accurate speech transcription of long-form audio","venue":null,"work_id":"9d8c6e0d-5d23-4f3e-aaa3-e8b1536c31a9","year":2023},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.108760Z"},"links":{"citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:7fb7f6424892ae444fb713acbccdc483be282663f34a966f3cfecde6b8f2ae15","observation_id":"cea29bd6-d807-447e-8473-51fda6a68ee4","resolution":{"observed_at":"2026-08-10T14:27:14.494082Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2005.14165","last_updated":"2020-07-22T19:47:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-05-28T17:29:03Z","title":"Language Models are Few-Shot Learners","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.14165","snapshot_observed_at":"2026-08-10T14:27:14.113049Z","title":"Language models are few-shot learners","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.113049Z"},"links":{"cited_paper":"/paper/2005.14165","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:b941bb0e992975d5d19d3aee59fb65aa3fe138523ed519dfe2bda9fb8b454113","observation_id":"665fcf13-ea0f-4a09-87b2-404b01573c47","resolution":{"observed_at":"2026-08-10T14:27:14.113049Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T14:27:14.475370Z","title":"Pubmedclip: How much does clip benefit visual question answering in the medical domain? In Findings of the Association for Computa- tional Linguistics: EACL 2023, pp","venue":null,"work_id":"f2d287a2-2b96-4a37-9f51-5135cb7fb727","year":2023},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.122834Z"},"links":{"citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:453fdd03c038383b6d1a3e20934a4ccdc0c83f999b2e839e135eb5ea31e47c52","observation_id":"b7fa4c21-f84c-4790-856c-c2715ff22b51","resolution":{"observed_at":"2026-08-10T14:27:14.480061Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1512.03385","last_updated":"2015-12-10T19:51:55Z","snapshot_observed_at":"2026-07-06T04:39:28.429064Z","submitted_at":"2015-12-10T19:51:55Z","title":"Deep Residual Learning for Image Recognition","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1512.03385","snapshot_observed_at":"2026-08-10T14:27:14.126892Z","title":"Deep residual learning for image recog- nition","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.126892Z"},"links":{"cited_paper":"/paper/1512.03385","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:5c4b411e05a0c925aa5165b27d8460dea4e07e02a406621d6cb57a71a0cccaa2","observation_id":"e15500c1-0dfd-4576-bc4d-657da58421db","resolution":{"observed_at":"2026-08-10T14:27:14.126892Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T14:27:14.447329Z","title":"A comprehensive study of gpt-4v’s multimodal capabilities in medical imaging","venue":null,"work_id":"cf6ae88a-a364-4821-9eda-e0abe4f72711","year":2023},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.139780Z"},"links":{"citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:26efec69c8e17ed174c8d1703e98bdd61e9a40691f49773b194163a1620dc69f","observation_id":"77dee18d-5e96-4ccb-8939-22736ce84f81","resolution":{"observed_at":"2026-08-10T14:27:14.451869Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T14:27:14.144036Z","title":"Microsoft coco: Common objects in context","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.144036Z"},"links":{"citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:d0be792479595755b05eb1f6b86af1c979f9ceb91d7fc0dc96196f934e34be17","observation_id":"2c8cc461-0293-444b-bbf9-0d407d06ab20","resolution":{"observed_at":"2026-08-10T14:27:14.144036Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.05499","last_updated":"2024-07-19T06:00:41Z","snapshot_observed_at":"2026-07-06T15:00:58.804337Z","submitted_at":"2023-03-09T18:52:16Z","title":"Grounding DINO: Marrying DINO with Grounded Pre-Training for Open-Set Object Detection","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.05499","snapshot_observed_at":"2026-08-10T14:27:14.148033Z","title":"Grounding dino: Marrying dino with grounded pre-training for open-set object detection","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.148033Z"},"links":{"cited_paper":"/paper/2303.05499","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:5b71f514433057dbaef8d99146f1bc3a2f0d049fd4d4af472bafb36c34f1b9f9","observation_id":"2e0350b9-7e4d-4d5d-9f73-9cd16242fe34","resolution":{"observed_at":"2026-08-10T14:27:14.148033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2204.05235","last_updated":"2023-02-28T17:12:35Z","snapshot_observed_at":"2026-08-09T18:23:04.295870Z","submitted_at":"2022-04-11T16:32:25Z","title":"Data Splits and Metrics for Method Benchmarking on Surgical Action Triplet Datasets","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.05235","snapshot_observed_at":"2026-08-10T14:27:14.156617Z","title":"doi: 10.18653/v1/W19-5034","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.156617Z"},"links":{"cited_paper":"/paper/2204.05235","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:df49af9a59b77f9ae34b5783a1e649183658743f9c007867d214204daf71e111","observation_id":"78b30a2b-3263-40bf-b165-8e456066929d","resolution":{"observed_at":"2026-08-10T14:27:14.156617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T14:27:14.424354Z","title":"Recognition of instrument-tissue interactions in endoscopic videos via action triplets","venue":null,"work_id":"59e574dc-d11d-464f-83dd-db8fbc387e96","year":2020},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.160603Z"},"links":{"citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:b0b7f5eee898420477f701eafd85012f3edd769d19dab01dcc24627285478c6e","observation_id":"deee046e-fd95-4189-b1f0-f1f55c9365d4","resolution":{"observed_at":"2026-08-10T14:27:14.428876Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T14:27:14.410366Z","title":"websurg.com","venue":null,"work_id":"18359b07-2072-45e6-883b-413d73b30f57","year":2024},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.164447Z"},"links":{"citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:193e7782cd1f2fc8bfab6bbea87a9f4de3eb89c91c2248de1d9bf6a6e53ea373","observation_id":"1ea11d78-d165-42dc-853c-9e5e0fe4b934","resolution":{"observed_at":"2026-08-10T14:27:14.414753Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15220","last_updated":"2025-06-16T19:07:50Z","snapshot_observed_at":"2026-07-06T15:59:30.956483Z","submitted_at":"2023-07-27T22:38:12Z","title":"Learning Multi-modal Representations by Watching Hundreds of Surgical Video Lectures","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15220","snapshot_observed_at":"2026-08-10T14:27:14.167961Z","title":"Learning multi-modal representations by watching hundreds of surgical video lectures","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.167961Z"},"links":{"cited_paper":"/paper/2307.15220","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:cf63bbe8287ebfb0a4f1692d82eef2398731685d984cf695bfe4a69436bab148","observation_id":"ae4da659-3a01-4e87-b81c-d92547942d0d","resolution":{"observed_at":"2026-08-10T14:27:14.167961Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.10075","last_updated":"2025-03-13T15:27:41Z","snapshot_observed_at":"2026-07-06T18:15:16.164110Z","submitted_at":"2024-05-16T13:14:43Z","title":"HecVL: Hierarchical Video-Language Pretraining for Zero-shot Surgical Phase Recognition","version":2},"cited_work":{"arxiv_id":"2405.10075","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.10075","snapshot_observed_at":"2026-08-10T14:27:14.229762Z","title":"HecVL: Hierarchical Video-Language Pretraining for Zero-shot Surgical Phase Recognition","venue":"cs.CV","work_id":"c6d76a8f-1aa5-47f8-a04b-14f1bfcd5502","year":2024},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.172270Z"},"links":{"cited_paper":"/paper/2405.10075","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:3b7c9d69c1b9185334cc053d01ef1d6f3201b743b2c58210d1cb4ccd7c7dd453","observation_id":"152bfeaf-8740-44a6-8d86-944b53861717","resolution":{"observed_at":"2026-08-10T14:27:14.236326Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T14:27:14.396104Z","title":"Surgicalsam: Efficient class promptable surgical instrument segmentation","venue":null,"work_id":"7931fbbe-4fa1-4884-904c-e77a375ce690","year":2025},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.176634Z"},"links":{"citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:e3532be6f4e2360e8f6cefaaf6f222eef85c555f42043dffc81d0f6cd07b69a0","observation_id":"b58a7a3b-bda2-4135-8e86-8aab6368aaea","resolution":{"observed_at":"2026-08-10T14:27:14.400957Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.00915","last_updated":"2025-01-08T22:58:51Z","snapshot_observed_at":"2026-07-06T14:57:39.647497Z","submitted_at":"2023-03-02T02:20:04Z","title":"BiomedCLIP: a multimodal biomedical foundation model pretrained from fifteen million scientific image-text pairs","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.00915","snapshot_observed_at":"2026-08-10T14:27:14.180633Z","title":"Biomedclip: a multimodal biomedical foundation model pretrained from fifteen million scientific image-text pairs","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.180633Z"},"links":{"cited_paper":"/paper/2303.00915","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:ba1de3f77d3d7e84d187334b47a01164625f370f4bec6970111e561d10a28593","observation_id":"15b7b0f1-3663-4563-8b18-60a708a441e1","resolution":{"observed_at":"2026-08-10T14:27:14.180633Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.07981","last_updated":"2024-08-15T07:00:20Z","snapshot_observed_at":"2026-07-06T19:00:57.743738Z","submitted_at":"2024-08-15T07:00:20Z","title":"LLaVA-Surg: Towards Multimodal Surgical Assistant via Structured Surgical Video Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.07981","snapshot_observed_at":"2026-08-10T14:27:14.135386Z","title":"Llava-surg: Towards multimodal surgical assistant via structured surgical video learning","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":2004,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.135386Z"},"links":{"cited_paper":"/paper/2408.07981","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:4931766e90ab28763333ce7a9dc5803f7abdc48b504c666468938ca5ac145d65","observation_id":"dc62d15e-98cc-40fe-b6ab-fb250a69c777","resolution":{"observed_at":"2026-08-10T14:27:14.135386Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-10T01:12:16.468283Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-10T14:27:14.117774Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":2015,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.117774Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:8ea20cab2f79d2b8c151071a13830b26180d0c439d6a3a85f50707a99303551a","observation_id":"1bbe3414-6d98-46c5-8313-e000484b6097","resolution":{"observed_at":"2026-08-10T14:27:14.117774Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2001.11190","last_updated":"2020-08-03T01:55:24Z","snapshot_observed_at":"2026-08-05T21:16:20.262505Z","submitted_at":"2020-01-30T06:37:07Z","title":"2018 Robotic Scene Segmentation Challenge","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2001.11190","snapshot_observed_at":"2026-08-10T14:27:14.094340Z","title":"2018 robotic scene segmentation challenge","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.094340Z"},"links":{"cited_paper":"/paper/2001.11190","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:74df3c2b7aa74ee3f1c2e19753c7f5d42cd34c066b373753d477850a322cccb5","observation_id":"9520d660-dde8-4799-a280-fde59746ee74","resolution":{"observed_at":"2026-08-10T14:27:14.094340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T14:27:14.503914Z","title":"Matis: Masked- attention transformers for surgical instrument segmentation","venue":null,"work_id":"a66604dd-7478-4590-8928-62cd677cf3bb","year":2023},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.099088Z"},"links":{"citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:cf646be4b5a7b6a76cb2dc5615c6f9713b2c1fde2399aa59b13459bae2d76fc3","observation_id":"bb15bb22-7f97-4d8d-85f3-a64027b7e04e","resolution":{"observed_at":"2026-08-10T14:27:14.507908Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1711.05101","last_updated":"2019-01-04T21:01:49Z","snapshot_observed_at":"2026-08-09T20:34:52.923500Z","submitted_at":"2017-11-14T14:24:06Z","title":"Decoupled Weight Decay Regularization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.05101","snapshot_observed_at":"2026-08-10T14:27:14.152397Z","title":"Decoupled weight decay regularization","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.152397Z"},"links":{"cited_paper":"/paper/1711.05101","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:1350165aeaf23afdbdca314f53fc1b4e3ae29cb09bca42aa019389dd0f4eeab3","observation_id":"a1c91e10-93f3-472d-b7b6-497027d8a265","resolution":{"observed_at":"2026-08-10T14:27:14.152397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1902.06426","last_updated":"2019-02-21T17:01:02Z","snapshot_observed_at":"2026-08-07T15:18:47.460768Z","submitted_at":"2019-02-18T07:08:36Z","title":"2017 Robotic Instrument Segmentation Challenge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1902.06426","snapshot_observed_at":"2026-08-10T14:27:14.088840Z","title":"2017 robotic instrument segmentation challenge","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.088840Z"},"links":{"cited_paper":"/paper/1902.06426","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:dadbc64683bed2c1f618d43e5d53e260b7887595a69dc8a27a74199acf528149","observation_id":"7ea4ecf4-d277-442a-add6-e01113fc7677","resolution":{"observed_at":"2026-08-10T14:27:14.088840Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.11174","last_updated":"2024-12-27T01:39:43Z","snapshot_observed_at":"2026-08-06T20:36:27.875441Z","submitted_at":"2024-01-20T09:09:52Z","title":"Pixel-Wise Recognition for Holistic Surgical Scene Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.11174","snapshot_observed_at":"2026-08-10T14:27:14.103774Z","title":"Pixel-wise recognition for holistic surgical scene understanding","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.103774Z"},"links":{"cited_paper":"/paper/2401.11174","citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:88be412a6c5b0d5fdff7d0b5730edae195c091c43b600524374ecad6dd8007e8","observation_id":"9a79688c-e3f5-49ed-8050-a1cbb363b90a","resolution":{"observed_at":"2026-08-10T14:27:14.103774Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T14:27:14.461227Z","title":"Ophnet: A large-scale video benchmark for ophthalmic surgical workflow understanding","venue":null,"work_id":"21c50511-97ab-4afb-b0fc-34fec11f037b","year":2025},"citing_paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-10T14:27:14.131305Z"},"links":{"citing_paper":"/paper/2501.15326"},"observation_digest":"sha256:f8a3dd850b08572ea80e738cd1e01c3236d4e24bc235c639ddaf323b639677b3","observation_id":"10d5ca56-e559-48ee-9169-822454a3290b","resolution":{"observed_at":"2026-08-10T14:27:14.466073Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2501.15326","last_updated":"2025-05-06T03:57:31Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-10T14:21:38.172280Z","submitted_at":"2025-01-25T21:01:52Z","title":"Recognize Any Surgical Object: Unleashing the Power of Weakly-Supervised Data"},"reference_resolution":{"displayed":22,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":13,"verified_exact":0,"verified_fuzzy":8},"total_outbound_references":22},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 11 August 2026, this Paper Citation Record lists 22 of 22 outbound references and 1 inbound Pith citation observation for arXiv:2501.15326."}