{"as_of":"2026-08-10T12:07:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c9546c21b7aab4fe44873ccffe66c7a7ad0f3fb67f1d94f55310a7fa62d54134","coverage":[{"denominator":89,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":89,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-08T12:28:35.008604Z","state":"measured"},{"denominator":92,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":92,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-31T04:55:07.002233Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-03T16:38:39.514289Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"cited_work":{"arxiv_id":"2605.06537","doi":null,"metadata_source":"pith","pith_arxiv_id":"2605.06537","snapshot_observed_at":"2026-07-03T16:38:39.514289Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","venue":"cs.CV","work_id":"ad226d8b-9b32-43cb-bee6-e0d5069860a2","year":2026},"citing_paper":{"arxiv_id":"2607.01751","last_updated":"2026-07-02T06:07:44Z","snapshot_observed_at":"2026-08-04T13:31:05.282842Z","submitted_at":"2026-07-02T06:07:44Z","title":"MedStreamBench: A Time-Aware Benchmark for Streaming and Proactive Medical Video Understanding","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-07-03T16:37:09.666491Z"},"links":{"cited_paper":"/paper/2605.06537","citing_paper":"/paper/2607.01751"},"observation_digest":"sha256:e1960c85084b25eaebdad1715a511b05e51f6f477a9e02db2c422558923734d1","observation_id":"2f7d6379-5e45-4668-92e5-300186793dfa","resolution":{"observed_at":"2026-07-03T16:38:39.515751Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.06537","snapshot_observed_at":"2026-07-14T14:52:42.813309Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.09880","last_updated":"2026-07-10T18:13:54Z","snapshot_observed_at":"2026-08-09T13:05:07.464071Z","submitted_at":"2026-07-10T18:13:54Z","title":"CLIR-Bench: Benchmarking Multimodal Question Answering over Irregular Clinical Time Series","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-14T14:52:42.813309Z"},"links":{"cited_paper":"/paper/2605.06537","citing_paper":"/paper/2607.09880"},"observation_digest":"sha256:9ce0375101118c65363b144ef6fafab01adf63c348915c61c28533281024f364","observation_id":"c44e243c-2bba-46b7-8550-f3038b76a576","resolution":{"observed_at":"2026-07-14T14:52:42.813309Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.06537","snapshot_observed_at":"2026-07-31T04:55:07.002233Z","title":"Alessandro Favero, Luca Zancato, Matthew Trager, Siddharth Choudhary, Pramuditha Perera, Alessandro Achille, Ashwin Swaminathan, and Stefano Soatto","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28516","last_updated":"2026-07-30T16:55:00Z","snapshot_observed_at":"2026-08-09T19:19:14.012446Z","submitted_at":"2026-07-30T16:55:00Z","title":"Beyond Frame Selection: Generative Latent Evidence Aggregation for Long-Video Understanding","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-31T04:55:07.002233Z"},"links":{"cited_paper":"/paper/2605.06537","citing_paper":"/paper/2607.28516"},"observation_digest":"sha256:a031e92b8b6eef6f005713039f919ee5a467f43a81b33a83fde77976563bc375","observation_id":"23c65698-d66c-49e6-bdc3-1f9c9aa562da","resolution":{"observed_at":"2026-07-31T04:55:07.002233Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2605.06537/citation-record","integrity":"/paper/2605.06537/integrity","json":"/paper/2605.06537/citation-record.json","paper":"/paper/2605.06537"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pixel-wise recognition for holistic surgical scene understanding","venue":null,"work_id":"531624f2-8157-46a6-a222-19b01bbdeb65","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:791673c198378ed3f07adfad4c49afe7427ab6f9742c92d8d3a9a62dcd48e13e","observation_id":"a718ec8a-ec67-4d4a-8aac-406108aa6e9d","resolution":{"observed_at":"2026-05-26T13:17:49.931456Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Real-colon: A dataset for developing real-world ai applications in colonoscopy","venue":null,"work_id":"c26c1da7-e6f1-4499-9d2d-e7c99d0f3275","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:e64f4a2963e27cb4ea5c5fa95505d3eaa2e512296dbf3f92950e6e7e66419e06","observation_id":"f563f633-34ba-4454-98ce-52fc6058e9bc","resolution":{"observed_at":"2026-05-26T13:17:50.026003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Artificial intelligence for surgical scene understanding: a systematic review and reporting quality meta-analysis","venue":null,"work_id":"7d8a12cc-bf5c-43ea-9071-328d30b24a72","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:8d5c7e50e4cd92424646a107d40a345996dd11c5d4ad2a5495d03242e80a474e","observation_id":"a6a44dd9-4f53-4e83-b302-9b027de06801","resolution":{"observed_at":"2026-05-26T13:17:50.033555Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Internvl: Scaling up vision foundation models and aligning for generic visual-linguistic tasks","venue":null,"work_id":"13871f3e-9e66-4ca8-97a0-a787cfdbab37","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:b52ad2a2f0e5167524813f4a95cb0e6295edc3cf30c262ac0e0b0046626ccc63","observation_id":"a5511266-0db8-4b9c-a8ed-e5a7f73f6a87","resolution":{"observed_at":"2026-05-26T13:17:50.040290Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Corley, Christopher D","venue":null,"work_id":"e68ab848-8163-4134-a5c4-96b91e91ce25","year":2014},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:705a45b8beb8451cb6739148f2810fcc823a9d10b5810c07a5e83535f9074ebb","observation_id":"997f73c3-bf99-407e-ba97-005929064c7b","resolution":{"observed_at":"2026-05-26T13:17:50.037030Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Lovat, Peter Mountney, et al","venue":null,"work_id":"dd812922-b681-479a-81dd-bd0dab921c02","year":2023},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:9331d66517d123bdecea86b26592d17cd48ec38f25c99c3d426a1886e30dca94","observation_id":"c660e2ec-7425-4a18-bab7-29eedcb6c27d","resolution":{"observed_at":"2026-05-26T13:17:50.054794Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Deep learning in surgical workflow analysis: A review of phase and step recognition","venue":null,"work_id":"3f306160-555f-44ce-8317-d1dc90313571","year":2023},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:7086f5d04a75cf36bea4c1b08f77dccf1847b38f07605f5d732230fd2cf8f856","observation_id":"b0c069f8-c698-4bd8-ba6b-d32eb6c5f5e4","resolution":{"observed_at":"2026-05-26T13:17:50.072712Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Standardized cine-loop documentation in abdominal ultrasound facilitates offline image interpretation","venue":null,"work_id":"1d8680a3-8963-4b63-8fb2-b9afe622efbc","year":2015},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:2184669d69ee372c2b0fe9153bd85685292f0d2a983427fa5bfea30d55ea300c","observation_id":"8c27e3e4-d092-4c5e-ab0c-b9855d2689d9","resolution":{"observed_at":"2026-05-26T13:17:49.998353Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"We're expanding our Gemini 2.5 family of models","venue":null,"work_id":"bc096867-75a5-491f-a596-605f3a1b05a6","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:8006efcb7e9e7b0a7efde88c4f05307ee521375556a99240a14ace58d20baf03","observation_id":"bd9fd3c4-93bb-4819-99a7-b4910ec7fdf6","resolution":{"observed_at":"2026-05-26T13:17:50.061692Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Yadlapati, Mark Benson, Andrew J","venue":null,"work_id":"cf940d7a-6e8d-40d2-9c12-e64a0976754a","year":2019},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:c86dbc517750aa08b135c58731dfb981e218ec1ae1566758085a8d84f6c392ba","observation_id":"0843f89c-cb6e-47de-bfaa-409b970b8dbe","resolution":{"observed_at":"2026-05-26T13:17:50.057687Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Real-time ultrasound demonstration of uterine isthmus contractions during pregnancy","venue":null,"work_id":"f1aa16fe-bc2c-4f16-8442-f476a0c42da8","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:b9433e08df3ac78af5d3e087e42682139bb6dbb48f564a6cc726e26e64f9f18e","observation_id":"fc665a48-946a-4192-bd85-39d52d93525c","resolution":{"observed_at":"2026-05-26T13:17:49.896244Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Video-mme: The first-ever comprehensive evaluation benchmark of multi-modal llms in video analysis","venue":null,"work_id":"803e6ad7-22f7-4d45-b1f9-274b99ad601a","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:0918dc4a08b9ee275fd2f7a732d08f1f31a16b3af74e26c132fd2162bcdc7e1e","observation_id":"beed10f3-921c-49ab-a8f0-eb48e4a5469c","resolution":{"observed_at":"2026-05-26T13:17:49.944640Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Khan, Sophia Bano, Hani J","venue":null,"work_id":"29e0cdcb-9594-42a1-8edd-d1287376bae7","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:c4b782032ec9321aa181d73cf19b91cf99eee9edd067a8be8c6f98d086297e54","observation_id":"808cda99-e8ae-4155-8085-8e4590fd4dd6","resolution":{"observed_at":"2026-05-26T13:17:49.889557Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Ophnet: A large-scale video benchmark for ophthalmic surgical workflow understanding","venue":null,"work_id":"5fa69f59-e0e3-45ce-83f6-8745152f6a54","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:a58dcbf9ecd61c7f1ef5a4f30e969dac6038d9e019c3a1b7d4e71cd834ff0531","observation_id":"2735fd1d-0f0f-4a4c-a76d-2348acc336df","resolution":{"observed_at":"2026-05-26T13:17:49.910547Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A new dataset and versatile multi-task surgical workflow analysis framework for thoracoscopic mitral valvuloplasty","venue":null,"work_id":"db2627f8-300c-47f5-92bc-d338234723a5","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:17509f0242e5f0a8f0bcda3b02beb28afbcaef42d22d3a92230d17f539ef2f38","observation_id":"8569740b-622e-44c3-ac3c-b35824ea44d9","resolution":{"observed_at":"2026-05-26T13:17:50.008378Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Lang, Luigi P","venue":null,"work_id":"06bf5224-f3a6-4b29-8649-5ed6a062e65a","year":2015},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:93e239d000a284459c981640a6fcc87d037c6dce2e06039ed0d12148513f36f3","observation_id":"df7d7a5c-411f-4533-88b7-7f0294074669","resolution":{"observed_at":"2026-05-26T13:17:50.064944Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"e l L Lavanchy, Sanat Ramesh, Diego Dall’Alba, Cristians Gonzalez, Paolo Fiorini, Beat P M \\","venue":null,"work_id":"5123b9ce-b6b7-41b0-9441-d6b6e365e521","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:294fdada509e8570ba9c8cf51eda1057ff54b93a315cb6de311400a9fe948697","observation_id":"af8f6b2c-b691-4372-b6e0-76162ef5a6c2","resolution":{"observed_at":"2026-05-26T13:17:50.068892Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Galar-a large multi-label video capsule endoscopy dataset","venue":null,"work_id":"5bd38c2c-119d-4cd0-85c4-76b53f66ea5c","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:b56ecb6361cb3d9bbadba82cb1e265c8189b45e16d79ab14575b8eccff6f14dc","observation_id":"436aaeef-a3d2-498b-ae19-6c9e0914e19a","resolution":{"observed_at":"2026-05-26T13:17:50.084053Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Detecting moments and highlights in videos via natural language queries","venue":null,"work_id":"617e507f-0460-4945-9a75-2e9f65fc0026","year":2021},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:0bac9a8c530d17d26e7258222a614f6ff68b1910b3617ec906d6343d8c0baff2","observation_id":"91a23ac5-3689-481b-9ff8-1b88c8632d2f","resolution":{"observed_at":"2026-05-26T13:17:50.061470Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mvbench: A comprehensive multi-modal video understanding benchmark","venue":null,"work_id":"0136fda9-92ba-414e-a9c1-da164a698c63","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:e6d60f9e31e4b8e645bffaf97cbbb869f299fbd234a927ddd171ef89825b25b9","observation_id":"f30fa90e-da7a-4231-a130-10a5f5784162","resolution":{"observed_at":"2026-05-26T13:17:50.039990Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T05:20:43.332327Z","title":"Surgpub-video: A comprehensive surgical video framework for enhanced surgical intelligence in vision-language model","venue":null,"work_id":"f055f872-7a12-44ab-97cd-92e6b8b0681f","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:2fa6fcd88a0e6bdeb0d63a8c03a930d22b38f6bf9cbc44015ffccd5384eb1b39","observation_id":"0a049a47-c8c3-4993-8df4-4fff880a1f8a","resolution":{"observed_at":"2026-05-26T13:17:50.026342Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Kimi K2.5 : Visual agentic intelligence","venue":null,"work_id":"91e71a07-d4d5-4018-b056-a177f11dee83","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:2cf854399af0727cd28e5b0c893ff8e01692718244720bd9d31018d2353b5c92","observation_id":"6d50a534-01eb-4938-8243-4cc1e49f63a1","resolution":{"observed_at":"2026-05-26T13:17:50.036588Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Introducing GPT -5.4","venue":null,"work_id":"3b93e405-1662-4b32-9fe3-ed1f37695748","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:a408c198e42f285a9757c447972820cc9967917544ddca943e8c62992c7d8b2a","observation_id":"90f0868a-1205-4eee-b52b-5f921a9e18c9","resolution":{"observed_at":"2026-05-26T13:17:50.015573Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Introducing GPT-5.4 mini and nano","venue":null,"work_id":"4601919a-2159-4805-942d-328e1bda6e83","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:5761b885f0d65eb3ecfa2fa25dea8237d0ec188066aec8e984cc02bb988dbe45","observation_id":"470f0ce3-485c-403f-b5aa-c26c13b7391a","resolution":{"observed_at":"2026-05-26T13:17:50.033216Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Qwen3.5 : Towards native multimodal agents","venue":null,"work_id":"a033766f-4a1f-404b-b1c4-be1e244e6acd","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:38e75676f7e67f9c6aec8189c2989fbfa15d863944bb19f0de148ee61125a717","observation_id":"9b3d469a-5d8c-4a5c-81b4-d98f790d7574","resolution":{"observed_at":"2026-05-26T13:17:50.050506Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Qwen3.6 : Advancing multimodal intelligence","venue":null,"work_id":"54cb807f-dacf-490d-91e2-fd551122cbe0","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:117d8320416f8cbcc0e9691f0f1c02448901c49ebdf61145bf4be042bc78fe25","observation_id":"b6ae069b-b226-4a0d-844f-2764e7500246","resolution":{"observed_at":"2026-05-26T13:17:50.080383Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Rex, Joseph C","venue":null,"work_id":"186ec27f-c96b-4233-9af9-afc0e602312d","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:41f45755348543149a2b558b8e0321dc43ae8e4e01dbd83070eda60b897c7608","observation_id":"550525d8-3425-4b87-87d4-42c7efe71373","resolution":{"observed_at":"2026-05-26T13:17:50.087126Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"648046bd-02cf-45a4-b18d-9b33d730e39b","year":2022},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:bcf7f0d2d26e1cb8c955edf4c15040b1247afce3bf68fbe74bbc3ed9b5e342e6","observation_id":"acef1349-8ae2-4bd2-8b42-db2413dc9e83","resolution":{"observed_at":"2026-05-26T13:17:50.012116Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Medgrpo: Multi-task reinforcement learning for heterogeneous medical video understanding","venue":null,"work_id":"a78a4bc7-034b-410a-8e10-27356400072b","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:f76c0ee31c2cf8ebc3d9b78a67182b6a439475079e420596e74178f72ee629a0","observation_id":"0577351b-1846-49b2-ad82-dc305ac23921","resolution":{"observed_at":"2026-05-26T13:17:50.023004Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Adaptive keyframe sampling for long video understanding","venue":null,"work_id":"4222cc7b-eede-4721-a3db-883caf034c7e","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:f2a4d8336384a4a4987b968788875dc6d77dbfd309e765d3773ff1b70f276c80","observation_id":"438f8796-466b-4cc1-949e-7d07169cc421","resolution":{"observed_at":"2026-05-26T13:17:50.029670Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mimic-iv-echo-ext-mimicechoqa: A benchmark dataset for echocardiogram-based visual question answering","venue":null,"work_id":"39fdce29-c099-4fdf-b086-c6ec3ba11fa6","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:c534058968987597d73db78de360acdee8c8f80a97d9b449d39e3bbd665a143f","observation_id":"46c8d370-4f9c-4591-8304-160694a87846","resolution":{"observed_at":"2026-05-26T13:17:50.043501Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gemini 3.1 pro: A smarter model for your most complex tasks","venue":null,"work_id":"d9d9fdcc-e2a3-41af-b54b-83ceaf8c25c7","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:0d564c2353bef93acb52ef5ec2d101039afa7ad63b6523b6a5bef06d677fa1d7","observation_id":"bc06ba78-eb71-4e1b-91e9-1367b8f3629b","resolution":{"observed_at":"2026-05-26T13:17:49.987435Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Lvbench: An extreme long video understanding benchmark","venue":null,"work_id":"55f0d6b2-fdc9-417a-80d4-00252682e539","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:083d4860e733314d46da7ac86abd41cfb0cab1c83832ad27650d207f24838113","observation_id":"4cf592f2-3312-4997-9792-e54622d60381","resolution":{"observed_at":"2026-05-26T13:17:49.990767Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Autolaparo: A new dataset of integrated multi-tasks for image-guided surgical automation in laparoscopic hysterectomy","venue":null,"work_id":"e9f8bc62-9e21-4ebf-b15c-da3fc61ca370","year":2022},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:21a50a1ba89403d483877dd9d9983ce91401886ca8cd2f244b43ad5d0366a190","observation_id":"ef52dbea-d4d8-4cf1-b07a-cad1968a9b09","resolution":{"observed_at":"2026-05-26T13:17:50.002139Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Longvideobench: A benchmark for long-context interleaved video-language understanding","venue":null,"work_id":"785f0959-d7b4-4290-bd97-6368ad758917","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:f625c39d6dac86869d9d3f5b8a4c3baed08bd4456d61ed4e538e805299589ec2","observation_id":"9c7a4c55-4aec-4aa3-8b1a-2aff74e4e4d8","resolution":{"observed_at":"2026-05-26T13:17:49.966155Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.01800","last_updated":"2024-08-03T15:02:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-08-03T15:02:21Z","title":"MiniCPM-V: A GPT-4V Level MLLM on Your Phone","version":1},"cited_work":{"arxiv_id":"2408.01800","doi":null,"metadata_source":"pith","pith_arxiv_id":"2408.01800","snapshot_observed_at":"2026-07-10T11:37:03.161139Z","title":"MiniCPM-V: A GPT-4V Level MLLM on Your Phone","venue":"cs.CV","work_id":"0f06e436-0c76-4e3c-be5e-6168f6bc4336","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2408.01800","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:e3feb80f171cc05f5308ecc865d390c071b8b2307b7e6df14c395504d0bd8a5b","observation_id":"b3ecea8b-ce60-4310-b7ab-6efa33b039b2","resolution":{"observed_at":"2026-05-11T19:16:07.486569Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T07:40:47.902764Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":"c7684330-0e20-4ced-aeea-ac696ea54900","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:56f71b85f776bef910d5e329a94990c1ee5619aa7b5b04a4618e9babd08b43f2","observation_id":"f62cf355-6bef-4d82-8735-97fd6c8a9661","resolution":{"observed_at":"2026-05-26T13:17:49.986798Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.18265","last_updated":"2025-08-27T14:39:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-08-25T17:58:17Z","title":"InternVL3.5: Advancing Open-Source Multimodal Models in Versatility, Reasoning, and Efficiency","version":2},"cited_work":{"arxiv_id":"2508.18265","doi":"10.48550/arxiv.2508.18265","metadata_source":"pith","pith_arxiv_id":"2508.18265","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"InternVL3.5: Advancing Open-Source Multimodal Models in Versatility, Reasoning, and Efficiency","venue":"cs.CV","work_id":"b8f5e260-fff5-444e-bcf5-2c42cfefd83d","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2508.18265","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:7f7f81a2be2c6c9c270c5a9a4278d8b5fe57ab8e8b6a9ffb82b7cf08689b18cb","observation_id":"019fdc59-210b-42f5-9ab9-a5c1eb01f229","resolution":{"observed_at":"2026-05-11T19:16:07.513468Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.10479","last_updated":"2025-04-19T03:47:21Z","snapshot_observed_at":"2026-08-10T07:07:56.707005Z","submitted_at":"2025-04-14T17:59:25Z","title":"InternVL3: Exploring Advanced Training and Test-Time Recipes for Open-Source Multimodal Models","version":3},"cited_work":{"arxiv_id":"2504.10479","doi":"10.48550/arxiv.2504.10479","metadata_source":"pith","pith_arxiv_id":"2504.10479","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"InternVL3: Exploring Advanced Training and Test-Time Recipes for Open-Source Multimodal Models","venue":"cs.CV","work_id":"fe8637aa-12bc-4434-8d36-9f57b5eebcbe","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2504.10479","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:c0d04e7d314097b6609d9f033d014edc07a025e5a8ffe2bd8c2c2511d4841c7d","observation_id":"0071f784-2caf-4bb8-a85a-b9944c00c5a5","resolution":{"observed_at":"2026-05-11T19:16:07.424343Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-05-20T07:54:09.017512+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T07:54:09.017512+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12191","last_updated":"2024-10-03T15:54:49Z","snapshot_observed_at":"2026-08-06T05:35:29.109022Z","submitted_at":"2024-09-18T17:59:32Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","version":2},"cited_work":{"arxiv_id":"2409.12191","doi":"10.48550/arxiv.2409.12191","metadata_source":"pith","pith_arxiv_id":"2409.12191","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","venue":"cs.CV","work_id":"8abcfe4f-e0fb-44b7-9123-448fac95f90a","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2409.12191","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:60eb11f950c8d152cd14088e2b8e266d4b30d79a1f7ec8312b3e2e14b1dedbef","observation_id":"e21d934f-0528-4846-8e6a-91bda0f4ea25","resolution":{"observed_at":"2026-05-11T19:16:07.450668Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-11T02:19:33.884263+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T02:19:33.884263+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":"2502.13923","doi":"10.48550/arxiv.2502.13923","metadata_source":"pith","pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen2.5-VL Technical Report","venue":"cs.CV","work_id":"69dffacb-bfe8-442d-be86-48624c60426f","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:f7b6182fd3b057382a5bd3540f11481097f8416cd8d1bcd67bf83aaa675b8dfb","observation_id":"b5a39afa-532c-46f0-b469-d4724efd14d3","resolution":{"observed_at":"2026-05-11T19:16:07.416124Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T16:08:16.864468+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T16:08:16.864468+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.21631","last_updated":"2025-11-27T12:16:54Z","snapshot_observed_at":"2026-07-06T22:37:03.716474Z","submitted_at":"2025-11-26T17:59:08Z","title":"Qwen3-VL Technical Report","version":2},"cited_work":{"arxiv_id":"2511.21631","doi":"10.1016/j.neunet.2025.107777","metadata_source":"pith","pith_arxiv_id":"2511.21631","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen3-VL Technical Report","venue":"cs.CV","work_id":"1fe243aa-e3c0-4da6-b391-4cbcfc88d5c0","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2511.21631","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:25cdfcd5c544783b5bee58226207e2c743db9f6627786acbf698f3b1697db8da","observation_id":"5c15d653-c77a-4879-878e-1d6a1ed7f8dd","resolution":{"observed_at":"2026-05-11T19:16:07.508619Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05201","last_updated":"2026-04-06T19:50:20Z","snapshot_observed_at":"2026-08-07T23:53:34.814339Z","submitted_at":"2025-07-07T17:01:44Z","title":"MedGemma Technical Report","version":4},"cited_work":{"arxiv_id":"2507.05201","doi":"10.48550/arxiv.2507.05201","metadata_source":"pith","pith_arxiv_id":"2507.05201","snapshot_observed_at":"2026-08-05T02:49:54.815029Z","title":"MedGemma Technical Report","venue":"cs.AI","work_id":"3d3f25c0-31e8-4859-bb3d-0d719b47a63d","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2507.05201","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:d348213f1e94226afe086c9465a59ad009f78dff4dc885b87fa51b649544694b","observation_id":"a0a72f57-7d51-48a1-8bb7-ed1803b01720","resolution":{"observed_at":"2026-05-11T19:16:07.492013Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.07044","last_updated":"2025-06-13T04:22:02Z","snapshot_observed_at":"2026-07-06T21:38:34.998414Z","submitted_at":"2025-06-08T08:47:30Z","title":"Lingshu: A Generalist Foundation Model for Unified Multimodal Medical Understanding and Reasoning","version":4},"cited_work":{"arxiv_id":"2506.07044","doi":"10.48550/arxiv.2506.07044","metadata_source":"pith","pith_arxiv_id":"2506.07044","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lingshu: A Generalist Foundation Model for Unified Multimodal Medical Understanding and Reasoning","venue":"cs.CL","work_id":"63908fd8-1967-4f5a-ae33-9d555390043d","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2506.07044","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:82087835eff14f5fd7d34f41ebb444246d60c702cdb371a0efd201d8900cd32b","observation_id":"2ed08202-a445-401a-910a-269f2ec7c7cd","resolution":{"observed_at":"2026-05-15T11:10:59.931236Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.19280","last_updated":"2024-09-30T06:45:16Z","snapshot_observed_at":"2026-08-06T18:57:51.817317Z","submitted_at":"2024-06-27T15:50:41Z","title":"HuatuoGPT-Vision, Towards Injecting Medical Visual Knowledge into Multimodal LLMs at Scale","version":4},"cited_work":{"arxiv_id":"2406.19280","doi":"10.48550/arxiv.2406.19280","metadata_source":"pith","pith_arxiv_id":"2406.19280","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2406.19280 , year=","venue":"cs.CV","work_id":"dd32b8a1-4ad4-4155-b031-c317b565c6e7","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2406.19280","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:0a12642fbc675c9a65ab48cd19d4e54c53a5516dac3addc1737c7e26e38c502a","observation_id":"2eb0ed3f-43a7-4ef2-a9cb-ba6d34e78ddb","resolution":{"observed_at":"2026-05-11T19:16:07.533998Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-05-25T22:53:37.302988+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-25T22:53:37.302988+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"e2f0371f-21e0-4de7-ba58-e60e71bf9b29","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:895a6fdd7cce64ba20f4e66f14b0a0fb62a85e04eeeb3454cc5b08f461e13e61","observation_id":"69fa3833-5c51-42a5-b987-3410237bacfe","resolution":{"observed_at":"2026-05-26T13:17:50.004660Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":"3789cb2c-65d3-43fc-9f85-184fa7e30b44","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4e362f61210a393489b707a31e2ece7a90a35941a7560256399277f7be60d9ed","observation_id":"19a12006-3b84-4e7e-b071-11e7cb984348","resolution":{"observed_at":"2026-05-26T13:17:49.965773Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.00512","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2603.00512 , year=","venue":null,"work_id":"8fda1f20-a92e-4b09-b880-4f106cff8973","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:56ce42d695c492c2804cf434b7c4868f812edefa31485cb9d894cc358b7e2876","observation_id":"e0155bb6-715e-4429-a952-b2da743e96ed","resolution":{"observed_at":"2026-05-11T19:16:07.429415Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"American Journal of Gastroenterology , volume=","venue":null,"work_id":"9c1652bb-073e-48e9-bd15-fe000a937e5d","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:e67203af6ccd84558da69eb8784f4aab55c3bd1cfc661e47725bd78f12efab4f","observation_id":"f1a9b58c-666c-44f1-b94f-1603ec438d54","resolution":{"observed_at":"2026-05-26T13:17:49.972052Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"New England Journal of Medicine , volume=","venue":null,"work_id":"752e1d2a-0013-444a-9227-5dc53968cf94","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:917a6e77b9874b5e73c6856db165182e2b3390c388e69a48fcd071c9732ec1e4","observation_id":"2e138424-c925-461f-80fb-dd668b94af6f","resolution":{"observed_at":"2026-05-26T13:17:49.977297Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Clinical Gastroenterology and Hepatology , volume=","venue":null,"work_id":"42f0b279-01d3-44ac-9311-3c44b8f7dbfa","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:27cbcbbea1922bc603910fcd4dc22b7b61495ff05e81dc7b951442e8741965da","observation_id":"66c1444c-2725-4d8e-99c2-302ea202736b","resolution":{"observed_at":"2026-05-26T13:17:50.001531Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2604.21814","last_updated":"2026-04-23T16:07:51Z","snapshot_observed_at":"2026-08-10T04:37:05.159589Z","submitted_at":"2026-04-23T16:07:51Z","title":"Divide-then-Diagnose: Weaving Clinician-Inspired Contexts for Ultra-Long Capsule Endoscopy Videos","version":1},"cited_work":{"arxiv_id":"2604.21814","doi":null,"metadata_source":"pith","pith_arxiv_id":"2604.21814","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Divide-then-Diagnose: Weaving Clinician-Inspired Contexts for Ultra-Long Capsule Endoscopy Videos","venue":"cs.CV","work_id":"51307048-7ab1-43b6-b5cb-65a0eb219f51","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2604.21814","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:8bf95c70e6d63d0f27da86f4db3ce9d494ad9747a25d5bc0034f4ea3bd964cfe","observation_id":"5e5cfe94-d8c6-4be6-98bd-85fd6cad65fd","resolution":{"observed_at":"2026-05-11T19:16:07.475502Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Biomedical Optics Express , volume=","venue":null,"work_id":"15cf71db-a581-4bf5-8d2d-f525f8c582a2","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:574083d9702237ac1bed78db8a3065a062a30229c85d429ed6d90981938fc6a4","observation_id":"b30d5025-9271-4960-8ff4-1b633bc269d0","resolution":{"observed_at":"2026-05-26T13:17:50.057919Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"6d9934b3-8cca-49e6-94c8-b02771385606","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:6c688826d7edb3376f43613a103980ad4502872d4359753fc87d9eb8c511147f","observation_id":"aef39177-fa86-475b-8d55-11d277d32df9","resolution":{"observed_at":"2026-05-26T13:17:49.941288Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Journal of the American Society of Echocardiography , volume=","venue":null,"work_id":"77716c18-d2d2-4716-b244-dfe5ac99e474","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:8d7335529d02a670e2363accbce1352f4f890bc1092c96084e29cb288b438617","observation_id":"2911fadb-fbe2-40a9-8a8e-8f4b862fc0ef","resolution":{"observed_at":"2026-05-26T13:17:49.938600Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Acta Radiologica , volume=","venue":null,"work_id":"794785a9-53df-41b3-b9c1-c74e2e7add15","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4dd301b53ece5e90932cf7c28a231387bdb826ebb44885f9b834394a4aa1f4b0","observation_id":"5fed9393-475c-4c68-bfe2-5ee5341ebc0b","resolution":{"observed_at":"2026-05-26T13:17:49.951537Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"IEEE Journal of Biomedical and Health Informatics , volume=","venue":null,"work_id":"2be417d6-cf08-46a8-adbc-ca6c3cdf51f6","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:e57813334cf8dfb3d108ebb9da0708027f987b0b5ddbdd0ac7da1367bcae3ec7","observation_id":"e9eeea78-7ff2-4b8e-9826-9581febb9da6","resolution":{"observed_at":"2026-05-26T13:17:50.005172Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"npj Digital Medicine , volume=","venue":null,"work_id":"dc4d17a6-5d74-47ac-8f04-546e9a5aeb74","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:985979463ff581eb4ca1f66cc5e2ea1b7724a38893da4d09797a72c057924152","observation_id":"8b786b6f-8085-42ea-9336-50f1551eb565","resolution":{"observed_at":"2026-05-26T13:17:50.054022Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":"2bba4d0b-b6f2-474d-ac7b-d2e970b8671f","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:9a76abb94c0c6e7546d99a8f18a52805be43b495c55a5b308b1ecf85b0a80c28","observation_id":"cc5d5c52-dfbc-4c47-8f58-ed3167bd43fa","resolution":{"observed_at":"2026-05-26T13:17:49.913560Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":"78162a81-0fff-46ee-a2ff-6adf67e3c6fb","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:6ba84854383f3b31c16d7e9c243b1459ac4cab1f6c92ee225650eccbbfc482a1","observation_id":"91028cd8-6cd9-4428-936a-47fddf4ba213","resolution":{"observed_at":"2026-05-26T13:17:49.916710Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Advances in Neural Information Processing Systems , pages=","venue":null,"work_id":"af1894ff-23d6-4b40-9d24-6d30b07653c1","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:87fbc085cbcefb9810e58de5c686905f139865d6e7bce94fc9b7307afb8dad8f","observation_id":"d4565e1a-02cc-4dd8-aa04-58045645717e","resolution":{"observed_at":"2026-05-26T13:17:50.076680Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Medical Image Computing and Computer Assisted Intervention -- MICCAI 2024 , pages=","venue":null,"work_id":"161d403a-de89-4a88-bf4a-71756e62dee6","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:0fa9982a976a92b1bf0372e93b130d4bbe0b8e80d6088547f42f26460143dcca","observation_id":"38cf701d-9669-4617-ad25-14d23916e5f6","resolution":{"observed_at":"2026-05-26T13:17:49.919692Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Computer Vision -- ECCV 2024 , pages=","venue":null,"work_id":"1cf72466-5118-4195-97da-943fa6086c5e","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:e07412ace07446017b60e7575d9261262292cf75c91e0df603d14ad90ff5c406","observation_id":"bcae6083-f207-4b3b-b375-3781da64a295","resolution":{"observed_at":"2026-05-26T13:17:49.915427Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.14391","last_updated":"2025-05-22T00:02:09Z","snapshot_observed_at":"2026-08-07T16:00:53.159295Z","submitted_at":"2025-04-19T19:32:15Z","title":"How Well Can General Vision-Language Models Learn Medicine By Watching Public Educational Videos?","version":2},"cited_work":{"arxiv_id":"2504.14391","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.14391","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2504.14391 , year=","venue":null,"work_id":"244ea0be-8d09-4c16-996a-b15b63b9bfc9","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2504.14391","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:c2eb167c9666b7d892bca6619cf2b7c4ca50198cff477e140494abc462ee9ae4","observation_id":"187c4c83-cee0-4c58-adf3-e85d4228e7f3","resolution":{"observed_at":"2026-05-11T19:16:07.526906Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.06570","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T20:10:07.469860Z","title":"arXiv preprint arXiv:2603.06570 , year=","venue":null,"work_id":"9b0e7001-3f8c-4a25-b531-ff1feaff8e9e","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:0e515186b5267f4a4c1ebc7a25cb546676baaa07296f01abd4bf5674b4cd3bb0","observation_id":"0554f377-4198-43b6-9270-2c061069cd4a","resolution":{"observed_at":"2026-05-11T19:16:07.444794Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":"b270aee9-a606-44d9-b44e-c4038d516bd8","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:46bd2f6a9ff15f4f0072c3a9f5aefe6720016bac4c8da759b80f987891519483","observation_id":"5a646395-0c8b-4d61-98ce-239c52630c6a","resolution":{"observed_at":"2026-05-26T13:17:49.925229Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2512.06581","last_updated":"2026-04-08T16:04:11Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-12-06T22:27:59Z","title":"MedGRPO: Multi-Task Reinforcement Learning for Heterogeneous Medical Video Understanding","version":4},"cited_work":{"arxiv_id":"2512.06581","doi":null,"metadata_source":"pith","pith_arxiv_id":"2512.06581","snapshot_observed_at":"2026-06-29T23:24:01.724884Z","title":"MedGRPO: Multi-Task Reinforcement Learning for Heterogeneous Medical Video Understanding","venue":"cs.CV","work_id":"369e9872-9603-4303-9d0e-579ce2204120","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2512.06581","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:8130038075b85ee047f2af7b0dd8af93c2f69bda2b7333496086af5fac883e76","observation_id":"c9370133-c00c-4d3a-8ec3-40a856bb46bf","resolution":{"observed_at":"2026-05-11T19:16:07.568594Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"PhysioNet , year=","venue":null,"work_id":"fdce3bb6-d810-43d4-9bfb-118217f64f90","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:28145489f61b436328920e24ad49223557f50db5e20dc32afb75f7b502371cdf","observation_id":"84fb6661-0a53-4bb8-9b00-4d0a6db3d4c6","resolution":{"observed_at":"2026-05-26T13:17:49.947580Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"MiniCPM-V: A","venue":null,"work_id":"98e9780f-bf4c-4322-b5ca-4635ffe1c27a","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:cda9d3f69e7a0b7367529bf3889484a110351ae9b78ba8001ee86e243fff13ff","observation_id":"cb1e5f02-3f42-4bb0-8dcd-e7c51e226f38","resolution":{"observed_at":"2026-05-26T13:17:49.980048Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02713","last_updated":"2025-08-01T16:40:14Z","snapshot_observed_at":"2026-08-02T12:24:31.329178Z","submitted_at":"2024-10-03T17:36:49Z","title":"LLaVA-Video: Video Instruction Tuning With Synthetic Data","version":3},"cited_work":{"arxiv_id":"2410.02713","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.02713","snapshot_observed_at":"2026-07-09T21:36:34.348434Z","title":"LLaVA-Video: Video Instruction Tuning With Synthetic Data","venue":"cs.CV","work_id":"e598f516-d992-449a-ab6d-6c788b3a1d7b","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2410.02713","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:f285a608556ff9f1ecd2950ceeb89bfcb708756e3cc54c4d12750c6bf1ab7a4f","observation_id":"a1207483-aae8-4b3e-a819-60c5bf9f5a17","resolution":{"observed_at":"2026-05-11T19:16:07.411617Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2026 , howpublished=","venue":null,"work_id":"53237011-62d0-4c49-8b3c-e085cc6d8df7","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:bf1f65dd1d2ecfc26f7f2f0738d55bb03590f79e78c9e85eb0aef8e5d63b44f9","observation_id":"ea2ba43d-e694-4a2c-bacc-163bef11ee43","resolution":{"observed_at":"2026-05-26T13:17:49.882550Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2026 , howpublished=","venue":null,"work_id":"3dbc0335-fc96-48cd-8762-9ac324f62d0f","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:603949b79ef8024c732e999361689deebbba2a5e28dda131fd01fbc3a7e76fcd","observation_id":"91fadc65-a633-4257-b8f8-563cae1d340c","resolution":{"observed_at":"2026-05-26T13:17:50.064538Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"We're Expanding Our","venue":null,"work_id":"8d9d9c99-3df4-47f8-940c-6dfe7d26adbe","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:2afc18441ff21ba8e5a2d5eaacf2295d18463d91573a6440aaff907572573e04","observation_id":"d3e82b2f-7208-42c1-a92e-db6769e119c1","resolution":{"observed_at":"2026-05-26T13:17:49.896640Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T13:42:38.156435Z","title":null,"venue":null,"work_id":"9aa422ef-371a-4d66-910b-f90d626e3a37","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:368ee5c1681bcd7a87652a42a901f41ec87b0138710c31fa1685cff01f5892e8","observation_id":"73fa5ec1-1c36-472f-a236-02bf9f1b6674","resolution":{"observed_at":"2026-05-26T13:17:49.876811Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T02:37:54.519855Z","title":null,"venue":null,"work_id":"72647878-5eae-4741-b770-5602b7561189","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4d1893ee78abe6a1eee9f557b92ca9518222b1c306f5f961b1ed01bbcf0cbee8","observation_id":"82b5c1a8-d259-4ac2-88de-e2a36a45c683","resolution":{"observed_at":"2026-05-26T13:17:49.879550Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.08668","doi":"10.48550/arxiv.2510.08668","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jiang, Y","venue":"ArXiv.org","work_id":"7a8fdc22-20a2-4741-865e-cb4dfc77234d","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:5ceb594197a13e046fe1bf9aee5b8ede4db051ae475dd5979e057866662099d4","observation_id":"4a176ee4-a216-495c-ba16-976d85638ee2","resolution":{"observed_at":"2026-05-11T19:16:07.461953Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13106","last_updated":"2025-06-03T03:33:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T18:59:46Z","title":"VideoLLaMA 3: Frontier Multimodal Foundation Models for Image and Video Understanding","version":4},"cited_work":{"arxiv_id":"2501.13106","doi":"10.48550/arxiv.2501.13106","metadata_source":"pith","pith_arxiv_id":"2501.13106","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"VideoLLaMA 3: Frontier Multimodal Foundation Models for Image and Video Understanding","venue":"cs.CV","work_id":"38f52461-37fd-4266-bc46-9dea31be2824","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2501.13106","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:16ace2a1163a0f2ba928b0b4474ed052e41037c20d8890e0d778c61cd0c34c28","observation_id":"82a5614e-764b-46d7-bfb3-09c97d2a20a4","resolution":{"observed_at":"2026-05-11T19:16:07.439260Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.16852","last_updated":"2024-07-01T02:59:29Z","snapshot_observed_at":"2026-08-07T09:52:45.942315Z","submitted_at":"2024-06-24T17:58:06Z","title":"Long Context Transfer from Language to Vision","version":2},"cited_work":{"arxiv_id":"2406.16852","doi":"10.48550/arxiv.2406.16852","metadata_source":"pith","pith_arxiv_id":"2406.16852","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Long Context Transfer from Language to Vision","venue":"cs.CV","work_id":"52f1b946-568f-4819-9d8a-a87296f8852d","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2406.16852","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:fb4130cd66458d30cee736717429bf77e9e8bd803df1b5d75326b03ae97eb1a4","observation_id":"ace738c8-8441-410f-949c-d37a5b0e0ab4","resolution":{"observed_at":"2026-05-12T07:08:36.726593Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.00574","last_updated":"2025-07-13T16:21:16Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-31T18:01:23Z","title":"VideoChat-Flash: Hierarchical Compression for Long-Context Video Modeling","version":4},"cited_work":{"arxiv_id":"2501.00574","doi":null,"metadata_source":"pith","pith_arxiv_id":"2501.00574","snapshot_observed_at":"2026-07-04T16:49:57.265026Z","title":"VideoChat-Flash: Hierarchical Compression for Long-Context Video Modeling","venue":"cs.CV","work_id":"52ec7cb6-1ef6-4366-a43f-d4c13fce23ad","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":97,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2501.00574","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:b467be4d703b459477ae673c7b301ba6376271fc4605e4474f689010cc5615c3","observation_id":"52e12638-e132-44a5-8b65-f8bfc00e0044","resolution":{"observed_at":"2026-05-18T04:02:43.915683Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"ae755293-c2d8-4387-9ff4-a49ab10b498a","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":98,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:570e8c651e7f2d19f7f42e2f5e744dffa2a33ed2a0d0dd01ad7a4084270fdee3","observation_id":"dc519c26-77e9-4426-a602-ef0533594b9a","resolution":{"observed_at":"2026-05-26T13:17:49.874038Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02438","last_updated":"2025-09-10T04:22:46Z","snapshot_observed_at":"2026-08-07T16:12:02.851576Z","submitted_at":"2025-04-03T09:55:09Z","title":"Scaling Video-Language Models to 10K Frames via Hierarchical Differential Distillation","version":5},"cited_work":{"arxiv_id":"2504.02438","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.02438","snapshot_observed_at":"2026-07-02T23:57:29.022097Z","title":"Scaling video-language models to 10k frames via hierarchical differential distillation","venue":null,"work_id":"0a6b0011-908b-4498-b4f3-1ba634df498f","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":99,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2504.02438","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:bc00cd12d8dc5a66d97ebbd852d6552d2ead7561be3b39db38b3a9bc1c45dab5","observation_id":"9b09393b-1b9e-482f-8bcb-fdd8bcb23fa0","resolution":{"observed_at":"2026-05-11T19:16:07.499176Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the IEEE/CVF International Conference on Computer Vision , pages=","venue":null,"work_id":"283c7846-9d25-4316-821f-7568e4dbd110","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":100,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:0e3ce783003a2f3b2303ca11fe7d2188bccd4f31f4913cf2958a902ea33e4253","observation_id":"684cd3d3-6f79-425a-b97b-7290acecb609","resolution":{"observed_at":"2026-05-26T13:17:49.918449Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2022 , eprint=","venue":null,"work_id":"e6419e28-c886-48fa-87b3-1e1f67a30715","year":2022},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":101,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:c69245f32beef59f0026d78966a75a03c0856d2d8dfdacf12398ba5ce26c94e5","observation_id":"ccc84196-5b42-48de-b067-6ba9ee596362","resolution":{"observed_at":"2026-05-26T13:17:49.892926Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Scientific Data , volume=","venue":null,"work_id":"5a373abd-34e6-4563-8600-afd9a7c98836","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":102,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4c5aec0174ace852dd9e28e908c5ae03b375c77ef68635522de33a4290e9af29","observation_id":"25fbe252-25f3-4be1-b5f0-851631c9ade1","resolution":{"observed_at":"2026-05-26T13:17:49.931659Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Medical Image Analysis , pages=","venue":null,"work_id":"c0536f93-142a-43f9-9365-e06847fc1c05","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":103,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:dcfc81c1969eb5c57ecea3577228432825e2c23ebbdd0af47db634609e5e31cc","observation_id":"d10158d3-0f78-441c-bb92-87b5b6673533","resolution":{"observed_at":"2026-05-26T13:17:50.011762Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"International journal of computer assisted radiology and surgery , volume=","venue":null,"work_id":"6d286faa-1c0c-48bc-957e-7db833a9b9ea","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":104,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:d6a37a21d894ef0befd0212a4e997cdde2352bcf21422aa5be95891f65451d80","observation_id":"da4a1ef2-d343-438b-a423-87c73cb583bd","resolution":{"observed_at":"2026-05-26T13:17:50.047466Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Scientific Data , volume=","venue":null,"work_id":"5aa9d441-f7b3-4f0f-bad6-65554b696e8c","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":105,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:7b2f44cd3be4db28772208355b305f31231de0c48cb02191b4f93faaf0e443f0","observation_id":"74e3f639-bc76-438e-a768-410f60ace084","resolution":{"observed_at":"2026-05-26T13:17:49.983377Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Medical Image Analysis , pages=","venue":null,"work_id":"68dcc263-45d1-4067-a5fa-4484368970aa","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":106,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:c60995f5e81fb118ace0bbe03469724b22ffb51bb17158ba345eddb851569a69","observation_id":"bf9a7a3c-8f97-4116-ad64-f06ffb2ee507","resolution":{"observed_at":"2026-05-26T13:17:49.849405Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"American Journal of Obstetrics and Gynecology , volume=","venue":null,"work_id":"b4e42c5a-a026-4b1c-82f1-7079cb214926","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":107,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:09920cbe2570d7629c3c0b129dde6f33fff97fcf747644c627776f6373370769","observation_id":"0585806c-8307-4504-88b9-f0f3ae88af75","resolution":{"observed_at":"2026-05-26T13:17:50.019074Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild"},"reference_resolution":{"displayed":89,"state_counts":{"malformed_identifier":0,"metadata_mismatch":15,"parse_uncertain":2,"unresolved":3,"verified_exact":5,"verified_fuzzy":64},"total_outbound_references":89},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 89 of 89 outbound references and 3 inbound Pith citation observations for arXiv:2605.06537."}