{"as_of":"2026-08-06T07:02:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:21e0739175c54a03ca85d8bad4a0cc64c2d51ffe9fe840467b4df451e9db6282","coverage":[{"denominator":48,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":48,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-16T07:16:29.588452Z","state":"measured"},{"denominator":50,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":50,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-06T06:34:29.942622+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T23:28:14.974682Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-06T00:16:16.768632Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2602.05638","snapshot_observed_at":"2026-08-01T11:30:57.042487Z","title":"& Others UniSurg: A Video-Native Foundation Model for Universal Understanding of Surgical Videos.ArXiv Preprint ArXiv:2602.05638","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19889","last_updated":"2026-07-22T08:20:52Z","snapshot_observed_at":"2026-08-05T17:22:37.189170Z","submitted_at":"2026-07-22T08:20:52Z","title":"LAVIFT: Latent-Action-Guided Vision Fine-Tuning for Surgical Interaction Recognition","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-01T11:30:57.042487Z"},"links":{"cited_paper":"/paper/2602.05638","citing_paper":"/paper/2607.19889"},"observation_digest":"sha256:2d8ade1f21cd8b282f15e887b805cd27a76776d4fcafd68a84c5e9947d642d8f","observation_id":"0ba885f5-863f-45d2-9ca0-d3b792d444ab","resolution":{"observed_at":"2026-08-01T11:30:57.042487Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"cited_work":{"arxiv_id":"2602.05638","doi":"10.48550/arxiv.2602.05638","metadata_source":"pith","pith_arxiv_id":"2602.05638","snapshot_observed_at":"2026-08-06T00:16:16.768632Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","venue":"cs.CV","work_id":"7644f6e1-208e-4065-9d9b-82b58253afcf","year":2026},"citing_paper":{"arxiv_id":"2608.03211","last_updated":"2026-08-04T06:49:04Z","snapshot_observed_at":"2026-08-06T06:21:06.304090Z","submitted_at":"2026-08-04T06:49:04Z","title":"CrossScope: A Role-Asymmetric World Model for Joint Dual-Scope Surgical Video Prediction","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-05T23:28:14.974682Z"},"links":{"cited_paper":"/paper/2602.05638","citing_paper":"/paper/2608.03211"},"observation_digest":"sha256:c5c852a91ab66f063e61024f9043b899aa5c1991f0b1bf68bda416deeb59a13f","observation_id":"9cf3481b-1f16-4c10-9caf-95b29adc03d8","resolution":{"observed_at":"2026-08-05T23:28:15.122891Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2602.05638/citation-record","integrity":"/paper/2602.05638/integrity","json":"/paper/2602.05638/citation-record.json","paper":"/paper/2602.05638"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2304.07193","last_updated":"2024-02-02T10:24:09Z","snapshot_observed_at":"2026-08-06T05:58:29.182448Z","submitted_at":"2023-04-14T15:12:19Z","title":"DINOv2: Learning Robust Visual Features without Supervision","version":2},"cited_work":{"arxiv_id":"2304.07193","doi":"10.48550/arxiv.2304.07193","metadata_source":"pith","pith_arxiv_id":"2304.07193","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DINOv2: Learning Robust Visual Features without Supervision","venue":"cs.CV","work_id":"26b304e5-b54a-4f26-be7e-83299eca52e4","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/2304.07193","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:d18caa14b361efdb6a7e0dd2ce2c3806f3110baf4aa1ba7955703fb59292773e","observation_id":"65c8424e-6828-4ea6-8cf6-2bab48d67022","resolution":{"observed_at":"2026-05-16T07:17:30.329076Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T18:32:48.038151Z","title":"Masked autoencoders are scalable vision learners","venue":null,"work_id":"23a19f52-833f-4385-a893-cba047433888","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:7036396c2a93fa0e4f7b38f293577226cf207504a117f78f5c8948b2125d9639","observation_id":"c7eaac3f-bcd2-4ab7-9889-f3bd8f9f0e80","resolution":{"observed_at":"2026-05-16T07:17:31.042736Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Masked autoencoders as spatiotemporal learners","venue":null,"work_id":"865e299e-505f-4a01-893f-33cacf47e6e3","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:18bcc33748bcf04f16a4eb110d0dfcbd82b5a728f7879dec579439aab69d3430","observation_id":"5a94e565-6873-4479-8a3f-d1166deb91bb","resolution":{"observed_at":"2026-05-16T07:17:31.049984Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Endovit: pretraining vision transformers on a large collection of endoscopic images","venue":null,"work_id":"4d636846-6483-451c-875f-5c26818de446","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:486fce793476a6061a30188a19d57f28df2a568b6f40138c3b1755e990b4abcf","observation_id":"a8c9586c-5198-4534-9ffa-53d9cfbe7d17","resolution":{"observed_at":"2026-05-16T07:17:31.045207Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Foundation model for endoscopy video analysis via large- scale self-supervised pre-train","venue":null,"work_id":"a6f381e7-048a-47af-9a80-d5db99acd397","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:7682077e3c793abda1e74581aa9753b1638561c40a11912c8b71d82741f4548c","observation_id":"f0317426-8f8a-4b0f-8a0b-9561d3a8b46c","resolution":{"observed_at":"2026-05-16T07:17:31.052145Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05949","last_updated":"2024-04-12T22:30:54Z","snapshot_observed_at":"2026-08-05T14:07:32.632054Z","submitted_at":"2024-03-09T16:02:46Z","title":"General surgery vision transformer: A video pre-trained foundation model for general surgery","version":3},"cited_work":{"arxiv_id":"2403.05949","doi":"10.48550/arxiv.2403.05949","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.05949","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Schmidgall, J","venue":"arXiv (Cornell University)","work_id":"e879f7aa-b06c-432c-949c-e12cf1d790ee","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/2403.05949","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:cb0f11989feafbd42b951e8bdc729f76ff7e4713f348bf92ec919a9296688d52","observation_id":"bf41add1-0ecf-4883-a5fa-aa24103cf21a","resolution":{"observed_at":"2026-05-16T07:17:30.324703Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Videomae: Masked autoencoders are data-efficient learners for self-supervised video pre-training","venue":null,"work_id":"e9210d53-bd15-4edf-8dbe-3aab9d819a81","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:ea85f5a9972ba58269511b59b9712c894c29238b2f64a971d15fefabfa5068e1","observation_id":"87a0759a-26b3-4e95-bce9-2fc11c56f802","resolution":{"observed_at":"2026-05-16T07:17:31.054747Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Videomae v2: Scaling video masked autoencoders with dual masking","venue":null,"work_id":"21d4295f-c66e-4abf-9660-50a659f6e94f","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:7d5f027d1326811365c2c314df78735a823b57cac91531627a1a6d6cab244bff","observation_id":"d34de7fd-0b84-4b29-b33b-c8fcc05c2a96","resolution":{"observed_at":"2026-05-16T07:17:31.056970Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Dissecting self-supervised learning methods for surgical computer vision","venue":null,"work_id":"a124c126-f3b3-416c-b4e6-745c77523762","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:c3f04dc411ff70cfa79953202214e5b19cf992b3a1cf425ea7ea85baa29f8eed","observation_id":"098ed1b4-6051-48cc-9cfb-d45625b74eef","resolution":{"observed_at":"2026-05-16T07:17:31.047502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Endonet: a deep architecture for recognition tasks on laparoscopic videos","venue":null,"work_id":"d77e719b-f339-4187-96b3-0407d900dc23","year":2016},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:53b38656516702d0a4068b1e2eb51caef481b371faff72ba90a11e9ee0d5eb5a","observation_id":"71913d2d-5ac9-4783-a5dc-4b4654ef6414","resolution":{"observed_at":"2026-05-16T07:17:31.059469Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pitvis-2023 challenge: Workflow recognition in videos of endoscopic pituitary surgery","venue":null,"work_id":"a4eed530-a7a2-48ed-b460-6b07bcf62c05","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:c9015d4b07faf81f07d3fdbd35b8abd27c6386c64fae14f16c2c032226a80775","observation_id":"100195a1-baa4-47bc-a1c2-67a51dad23bf","resolution":{"observed_at":"2026-05-16T07:17:31.016048Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Egosurgery-phase: a dataset of surgical phase recognition from egocentric open surgery videos","venue":null,"work_id":"9c176e39-5383-40ac-aded-11d59f7e4edf","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:a41d618ae02c98a95b0e2f8c088442f61f01967fe2c35bf6f4fd38f4b365fe0c","observation_id":"b2615ac1-326c-4387-9d6a-beaf33b18a48","resolution":{"observed_at":"2026-05-16T07:17:31.026298Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.08471","last_updated":"2024-02-15T18:59:11Z","snapshot_observed_at":"2026-08-02T05:40:31.086014Z","submitted_at":"2024-02-15T18:59:11Z","title":"Revisiting Feature Prediction for Learning Visual Representations from Video","version":1},"cited_work":{"arxiv_id":"2404.08471","doi":"10.1145/3178876.3185996","metadata_source":"pith","pith_arxiv_id":"2404.08471","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Revisiting Feature Prediction for Learning Visual Representations from Video","venue":"cs.CV","work_id":"f7251dcf-5341-4915-bfe7-27812387b61a","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/2404.08471","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:eb6e88cccacf0a495a7ea185a955a4cb9799c0554b194ac6b971a6575b6483de","observation_id":"b991f193-b445-45b3-8239-99d6d20e37d5","resolution":{"observed_at":"2026-05-16T07:17:30.319773Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-05-23T01:54:34.489718+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T01:54:34.489718+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09985","last_updated":"2025-06-11T17:57:09Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-11T17:57:09Z","title":"V-JEPA 2: Self-Supervised Video Models Enable Understanding, Prediction and Planning","version":1},"cited_work":{"arxiv_id":"2506.09985","doi":"10.48550/arxiv.2506.09985","metadata_source":"pith","pith_arxiv_id":"2506.09985","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"V-JEPA 2: Self-Supervised Video Models Enable Understanding, Prediction and Planning","venue":"cs.AI","work_id":"a9c28401-f16a-4933-89f0-788e2f94e52b","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/2506.09985","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:12ff2e35b009b8f1979357c34a779757835a733fc4bf820858e9376195f476b5","observation_id":"b5c27669-8051-4f81-9fca-a33b4e8e9d6e","resolution":{"observed_at":"2026-05-16T07:17:30.309716Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-03T19:08:37.559938+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-03T19:08:37.559938+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Bootstrap your own latent: A new approach to self-supervised learn- ing","venue":null,"work_id":"fa22a7fc-4e5f-4eef-a115-8d3356ff9f10","year":2020},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:d7282b609d197e3d137680d99f067cbcd09f0ee3604df6f34de86a2722e0efbc","observation_id":"1f629bb9-cf2f-445b-a19a-b81cd415a272","resolution":{"observed_at":"2026-05-16T07:17:31.021009Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Internvideo2: Scaling video foundation models for multimodal video understanding","venue":null,"work_id":"5cf2349d-b2fe-4140-a4e1-2a382a56f596","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:11212dfaca3b872bcc92a9f810534723f9905b29715a3de0c4a0d1e60da17835","observation_id":"4e21a4e5-bfbb-4062-9c7e-20bf4afa6a3a","resolution":{"observed_at":"2026-05-16T07:17:31.018463Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2512.01342","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T14:28:32.049805Z","title":"Internvideo-next: Towards general video foundation models without video-text supervision","venue":null,"work_id":"d65f61ad-e3f7-4e3f-9b08-da7a6e8b0595","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:7572954a8e9d71de6ee9a3ecdd04b93e33720591bcf99259f55188dfe21090d0","observation_id":"7214fd01-073e-495c-ba87-88203e7ec15b","resolution":{"observed_at":"2026-05-16T07:17:30.296414Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T23:44:22.359690Z","title":"Emerging properties in self-supervised vision transformers","venue":null,"work_id":"c47546d4-4e91-494a-a5e3-a1e01ff6556f","year":2021},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:0ad4774b505344c01f7cbb7ec058a9d8d629f5504f65a30c45ae6fc3dbbffad2","observation_id":"8d73de36-f6d3-4ab2-bc36-7e6bf6778573","resolution":{"observed_at":"2026-05-16T07:17:31.023621Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.10104","last_updated":"2025-08-13T18:00:55Z","snapshot_observed_at":"2026-07-06T22:12:35.584339Z","submitted_at":"2025-08-13T18:00:55Z","title":"DINOv3","version":1},"cited_work":{"arxiv_id":"2508.10104","doi":"10.1055/a-2487-1252","metadata_source":"pith","pith_arxiv_id":"2508.10104","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DINOv3","venue":"cs.CV","work_id":"c8b07deb-8fe7-4e18-9620-f3569d3529ce","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/2508.10104","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:a22cb36a0c477a2e8d8bec3c811bc292e7274bf4b1c0c6108e8ef28790d85cc4","observation_id":"1021f588-4279-4938-b861-5e2df567ecaf","resolution":{"observed_at":"2026-05-16T07:17:30.314610Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gastronet-5m: A multicenter dataset for developing foundation models in gastrointestinal endoscopy","venue":null,"work_id":"5b482dd4-c8fb-4307-8e07-c3f44db1548c","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:9c57e523b46ac148019df68c980fc261957752ac1d4650f92f8944fa4c86d0a6","observation_id":"4e6d5886-4dd6-4273-a38c-99dd91c43f7c","resolution":{"observed_at":"2026-05-16T07:17:31.028649Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Self-supervised learning for endoscopic video analysis","venue":null,"work_id":"70739a11-dcec-434d-b6cc-b8cbefba7c3e","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:f9938303b533efb687d5f5de396617d1cac95b003992e15b99706fe5ebeb8d3f","observation_id":"1c07827e-2596-4bfa-be7d-42d423e629cc","resolution":{"observed_at":"2026-05-16T07:17:31.011155Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Endomamba: an efficient founda- tion model for endoscopic videos via hierarchical pre-training","venue":null,"work_id":"bfe27351-6c16-4dbd-87f9-0bda669fd6ba","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:46c8b6496ca0793e3b850f2949b4512c532c412f97f871a59f78242902c134d2","observation_id":"ebc9fd4d-172e-470d-8857-de06490ac087","resolution":{"observed_at":"2026-05-16T07:17:31.005691Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2501.09436","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Scaling up self-supervised learning for improved surgical foundation models","venue":null,"work_id":"4271cc8e-489b-44fa-8dfd-c0ae7720647e","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:891b6ff75ee537c28223535301e8f002b2a37d7a02a910fea48eb9bcde8a5bdd","observation_id":"6fd0697f-7fca-40d5-a3fa-587ab0e66800","resolution":{"observed_at":"2026-05-16T07:17:30.292026Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Learn- ing multi-modal representations by watching hundreds of surgical video lectures","venue":null,"work_id":"c2417f7b-a452-4aa6-8abb-f9b5fecdee35","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:f3d9800b8a238be7cde58a03a55def9cb4f6abf6a9c1d3c3e7ef981189290a5c","observation_id":"a0969ef8-c80e-4c83-8510-7030bec8b518","resolution":{"observed_at":"2026-05-16T07:17:31.002969Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1610.09278","last_updated":"2017-08-31T14:27:37Z","snapshot_observed_at":"2026-08-03T02:45:46.755540Z","submitted_at":"2016-10-28T15:36:58Z","title":"The TUM LapChole dataset for the M2CAI 2016 workflow challenge","version":2},"cited_work":{"arxiv_id":"1610.09278","doi":null,"metadata_source":"pith","pith_arxiv_id":"1610.09278","snapshot_observed_at":"2026-07-04T20:10:07.403035Z","title":"The TUM LapChole dataset for the M2CAI 2016 workflow challenge","venue":"cs.CV","work_id":"020971ab-45bc-489b-9897-f0255dc6d71f","year":2016},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/1610.09278","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:163da7d47f01e6dbac0810d2c1e80618c0b5802204d15053d308438bd2574024","observation_id":"ae41b1ce-250c-4a0b-86ad-af55f87cc6f9","resolution":{"observed_at":"2026-05-16T07:17:30.299900Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Rendezvous: Attention mechanisms for the recognition of surgical action triplets in endoscopic videos","venue":null,"work_id":"117d7265-b07e-4307-a8e9-a450e4519ef0","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:369d2ba22e5de82bbed920f0b4a46adbe7af3cfdc434b723d6921d2dcefe933b","observation_id":"881fa467-b602-416e-b807-a77cf03bdad8","resolution":{"observed_at":"2026-05-16T07:17:31.000856Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Autolaparo: A new dataset of integrated multi-tasks for image-guided surgical automation in laparoscopic hysterectomy","venue":null,"work_id":"f9426651-8798-47fd-8b49-0e36d6f648dd","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:5678924d30d9942509e04278611d668922a75f88de282a68f8465e8337ea2651","observation_id":"7ee17f12-9e50-4697-b874-622379a6e908","resolution":{"observed_at":"2026-05-16T07:17:30.998477Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Surgical workflow recognition and blocking effectiveness detection in laparoscopic liver resection with pringle maneuver","venue":null,"work_id":"0048b71b-f79f-4709-aa9a-0ac2c7ef5012","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:4ca3614df9856bdeaad0c0d3aa1f3baceed0722c82a2ea58d19a5d723200574d","observation_id":"91c4640e-79da-496f-a89b-21038eaea380","resolution":{"observed_at":"2026-05-16T07:17:31.031094Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Ophnet: A large-scale video benchmark for ophthalmic surgical workflow understanding","venue":null,"work_id":"a6fb6497-8c30-4650-b48c-5d17e575ae25","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:160f163435b6db0238c2901e8f528dc7f8e6f1a28c2fe818ea2815e8a9bf06d5","observation_id":"bd78fb99-e91a-4db8-b9b3-a2cb15f72130","resolution":{"observed_at":"2026-05-16T07:17:31.013524Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Analyzing surgical technique in diverse open surgical videos with multitask machine learning","venue":null,"work_id":"b5fb16e1-78bd-470f-96d2-9dd6a7e6a033","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:d50a7d7466f4fb72eb4f662b790ed478b4cb6c39ba1e142b621f6e9b7a7892e4","observation_id":"f4bda7a8-62b1-4c61-89a2-f4f9cae6c375","resolution":{"observed_at":"2026-05-16T07:17:30.985662Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A dataset and benchmarks for segmentation and recognition of gestures in robotic surgery","venue":null,"work_id":"5cf765af-4ac2-4125-8fe3-0181e9c9092b","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:f228a42b76aa32185c087418dd47da8466132612f528f9acebb0e13ac9bfaf6b","observation_id":"bcaee95c-e5e9-48f3-97a5-39cdcec28c13","resolution":{"observed_at":"2026-05-16T07:17:31.038165Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Aixsuture: vision-based assessment of open suturing skills","venue":null,"work_id":"28a4f930-c4b9-49db-94ce-274bc2709727","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:99459ff382c78e0ff82d0c6f3baf91642e8b87d29329efe98060f1270ed6060d","observation_id":"416044c4-e24f-4216-8ad4-2198076a5ebd","resolution":{"observed_at":"2026-05-16T07:17:31.033391Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Video retrieval in laparoscopic video recordings with dynamic content descriptors","venue":null,"work_id":"3b024d39-879a-492e-bb91-feb1f3e58b77","year":2018},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:20fe37b94e89290ea127446db5e89555c2dba3f9b878bf5f2ad27710ddb1779c","observation_id":"7f67398d-310d-4c54-844f-8d8438a43f0a","resolution":{"observed_at":"2026-05-16T07:17:30.983458Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Contrastive transformer- based multiple instance learning for weakly supervised polyp frame detection","venue":null,"work_id":"3720a984-cc0e-4076-98df-dbb137e4825c","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:f5058e9bdf94b5bc36c13eadb40a011b3ba2861955bd771bc4d779741e2370d2","observation_id":"0e5c71cd-c19b-4398-98b6-8ffae8dba498","resolution":{"observed_at":"2026-05-16T07:17:31.035865Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Wm-dova maps for accurate polyp highlighting in colonoscopy: Validation vs. saliency maps from physicians","venue":null,"work_id":"8432b16d-9f06-464d-829d-921f9ccea4a6","year":2015},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:1ece45980aafb62b920be360a284f9d9c85901a916b517198e04ec2de972d27a","observation_id":"f98b61b0-9374-4ff1-a737-679468ca8c54","resolution":{"observed_at":"2026-05-16T07:17:30.989951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Implicit domain adaptation with conditional generative adversarial networks for depth prediction in endoscopy","venue":null,"work_id":"2b36d126-2df8-4676-8fa8-a55a58f1334c","year":2019},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:1b5040dbf1d961b7294f8bd2d0bb8c0dc91988cf1a86fd15a2ab0d10044ee38b","observation_id":"2cc3a194-d676-4fae-88de-d27188eeae67","resolution":{"observed_at":"2026-05-16T07:17:31.040582Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Colonoscopy 3d video dataset with paired depth from 2d-3d registration","venue":null,"work_id":"f5bcba8c-59c3-47ce-ae77-f34ec63c7d76","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:b9576c745e2de661b1a371c1b4e40adae80af23c05fd47b40a4d4fafd6e696ac","observation_id":"7df0091d-7d66-432f-a3d1-bdbd0e6f5b4f","resolution":{"observed_at":"2026-05-16T07:17:31.008591Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.21227/ac97-8m18","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Cataracts: Challenge on automatic tool annotation for cataract surgery","venue":"IEEE DataPort","work_id":"0f81b61c-8a7b-4a7b-a2f2-9da6d7c7291d","year":2019},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:243e7777a2702ca5ed9b1f98ca21954b5c460511b2c676d184e45d4514fcc601","observation_id":"7a94bfda-bcd3-46bc-a550-8a6d427f0949","resolution":{"observed_at":"2026-05-16T07:17:30.179924Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Challenges in multi-centric generalization: phase and step recog- nition in roux-en-y gastric bypass surgery","venue":null,"work_id":"5e4ca10e-ec87-4245-aaa0-72e344cb63d1","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:a3893f634ee97da4b0c8323e847eca3b2acfa7732b2afc6d5ba278bd99c2c752","observation_id":"eef2e08a-40b3-4bf1-b948-f20326cd3f64","resolution":{"observed_at":"2026-05-16T07:17:30.976472Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Copesd: A multi- level surgical motion dataset for training large vision-language models to co-pilot endoscopic submucosal dissection","venue":null,"work_id":"8d7ea69c-56c6-4396-8019-6fc40289e1d9","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:ea8a1f7f136bc80ba7a244f23d06788dd2949260cddb7e7e6a8b556e0227bbd8","observation_id":"2d9333d2-f622-4027-8bdd-c18428c6e1b9","resolution":{"observed_at":"2026-05-16T07:17:30.978987Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Towards holistic surgical scene understanding","venue":null,"work_id":"76c9e713-3f89-4ecd-ba64-4bcbc18ecbf9","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:da7bc166bc299523f9bd725819ffeadf60c47c6a0fb9153cd167747708951dbe","observation_id":"fe24adf8-d2a4-435a-b1e2-1636d1db2711","resolution":{"observed_at":"2026-05-16T07:17:30.974178Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Kvasir-seg: A segmented polyp dataset","venue":null,"work_id":"b48e79c0-6a76-47e0-9a9a-400f1b927c27","year":2020},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:62ff48001d35d22fd7be0d444c9e850a7b45dab316303dac050d593d307b49a9","observation_id":"00a9411f-b337-447d-82ab-0ab0a47cb83b","resolution":{"observed_at":"2026-05-16T07:17:30.971963Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A benchmark for endoluminal scene segmentation of colonoscopy images","venue":null,"work_id":"b15940e6-b237-4970-ad13-87c03c4f35bf","year":2017},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:b32489da48f4dd6faeb5fe71167472020082efa7fef26294f8240421324e402f","observation_id":"10bfce8b-2270-4931-b1ae-c7ab493c2d0e","resolution":{"observed_at":"2026-05-16T07:17:30.981215Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Towards automatic polyp detection with a polyp appearance model","venue":null,"work_id":"2b1717a8-acd8-4c3d-ad9c-e1782d9a8db5","year":2012},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:841949bddff561bcada0729ce5ec92b872f7198fa566efbfb41bb5525275ca0a","observation_id":"0d5a1ad6-0bd0-4022-9370-4de51db83397","resolution":{"observed_at":"2026-05-16T07:17:30.987823Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Toward embedded detection of polyps in wce images for early diagnosis of colorectal cancer","venue":null,"work_id":"5334d03e-687a-4d7e-9e11-36e593cb3a94","year":2014},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:9e62880cf6b5a7f0c99d977e59e11553b8a672754acf1e17d9302ed166b02e86","observation_id":"59a4ce96-7c01-40eb-8b5b-5b2c83f43849","resolution":{"observed_at":"2026-05-16T07:17:30.992302Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pranet: Parallel reverse attention network for polyp segmentation","venue":null,"work_id":"ed1aa94e-7244-4e6f-8ca2-1d780406bdf4","year":2020},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:e02ca95c282ffe46b09f936a0a2d2136b419b8ba8273a0390c69c86804644105","observation_id":"91b31e61-a7de-4ceb-8c61-a2b21be828ba","resolution":{"observed_at":"2026-05-16T07:17:30.969766Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Uacanet: Uncertainty augmented context attention for polyp segmentation","venue":null,"work_id":"3fbf992d-3eff-4ef6-b6be-7dc37bfb8ddb","year":2021},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:25d91648de544442e2c3872b673262d9f0320a3f5a692363e425b3fc7db767a8","observation_id":"4d2cb680-6211-4d5a-b3bc-b180c933ea31","resolution":{"observed_at":"2026-05-16T07:17:30.967308Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2504.10986","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pranet-v2: Dual-supervised reverse attention for medical image segmentation","venue":null,"work_id":"8d8b11c2-95c7-4fb3-a135-4a75a0d83f52","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:95cf86203d46b20a2777cd3acc89bea8d4fd14cf2acf20b0f46fc3622448cd9b","observation_id":"243a59d5-e604-4e35-950e-9cf958af9d13","resolution":{"observed_at":"2026-05-16T07:17:30.304570Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-04T22:30:41.280585Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos"},"reference_resolution":{"displayed":48,"state_counts":{"malformed_identifier":0,"metadata_mismatch":2,"parse_uncertain":0,"unresolved":0,"verified_exact":8,"verified_fuzzy":38},"total_outbound_references":48},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"thesis":"As of 6 August 2026, this Paper Citation Record lists 48 of 48 outbound references and 2 inbound Pith citation observations for arXiv:2602.05638."}