{"as_of":"2026-08-19T16:35:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c9674b2f5d15cb9b99bcc240be12e97dc977fea3be9d20923a922dab283d9dd3","coverage":[{"denominator":60,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":60,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T10:46:09.996574Z","state":"measured"},{"denominator":62,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":62,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T19:34:04.544567Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-06-28T19:22:34.657932Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.03800","snapshot_observed_at":"2026-08-03T19:34:04.544567Z","title":"Med- vista3d: Vision-language modeling for reducing diagnostic errors in 3d ct disease detection, un- derstanding and reporting.arXiv preprint arXiv:2509.03800, 2025b","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2512.00239","last_updated":"2026-06-07T22:52:48Z","snapshot_observed_at":"2026-08-17T23:17:39.627065Z","submitted_at":"2025-11-28T22:53:31Z","title":"Self-Supervised Dynamical System Representations for Physiological Time-Series","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-03T19:34:04.544567Z"},"links":{"cited_paper":"/paper/2509.03800","citing_paper":"/paper/2512.00239"},"observation_digest":"sha256:49121d81784f328c5930cff847bfd26a441f808f1ed14f1e50ef19c2fec200ae","observation_id":"0cd03453-c97c-4b33-9b7c-1d3daeea6f74","resolution":{"observed_at":"2026-08-03T19:34:04.544567Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"cited_work":{"arxiv_id":"2509.03800","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.03800","snapshot_observed_at":"2026-06-28T19:22:34.657932Z","title":"Medvista3d: Vision-language modeling for reducing diagnostic errors in 3d ct disease detection, understanding and reporting,","venue":null,"work_id":"a5511f07-ba18-4865-b330-19daca5fc28e","year":2025},"citing_paper":{"arxiv_id":"2606.00602","last_updated":"2026-05-30T07:59:21Z","snapshot_observed_at":"2026-08-14T18:26:22.873680Z","submitted_at":"2026-05-30T07:59:21Z","title":"ASAP: Advancing Medical Volumetric Representation Learning with Anatomy-aware Semantically-adaptive Pre-training","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-06-28T19:16:42.139096Z"},"links":{"cited_paper":"/paper/2509.03800","citing_paper":"/paper/2606.00602"},"observation_digest":"sha256:5e98aee3bcccb753805bda826826723d23959426943c6f7f79917466afd3e195","observation_id":"e8ecbf8a-69eb-48e9-9d43-16ea5263acd2","resolution":{"observed_at":"2026-06-28T19:22:34.659378Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2509.03800/citation-record","integrity":"/paper/2509.03800/integrity","json":"/paper/2509.03800/citation-record.json","paper":"/paper/2509.03800"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:04.243549Z","title":"Merlin: A vision language foundation model for 3d computed tomography","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:04.243549Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:2455174a5bdc51d462c8ecfcb2d72bc4330b08f46b91a06ee411bf5327e5ea60","observation_id":"a6fa80d1-f5d6-4f4a-8c11-f7590af040bf","resolution":{"observed_at":"2026-08-05T10:46:04.243549Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.953946Z","title":"A vision–language foundation model for the generation of realistic chest x-ray images","venue":null,"work_id":"8cf289c1-4b0b-493a-bf6b-56c20d411669","year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:04.330865Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:d85ec207920c297f6aee8807135a2188a4a4b2b0478b3d4449e14b237a0fb445","observation_id":"b33a7ab9-9f9f-45d3-8445-97b33ff2b16b","resolution":{"observed_at":"2026-08-05T10:46:13.958572Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:04.426003Z","title":"Making the most of text semantics to improve biomedical vision–language processing","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:04.426003Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:174915b43604b3a8652da8a1bc60274086db1c56016d7cf00460c7c530e8ab1a","observation_id":"cb37fff5-0668-4885-9707-e0ec24f3015e","resolution":{"observed_at":"2026-08-05T10:46:04.426003Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.928235Z","title":"Understanding and confronting our mis- takes: the epidemiology of error in radiology and strategies for error reduction","venue":null,"work_id":"584728aa-0d65-4a74-bd77-8eba9d20f6ca","year":2015},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:04.516487Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:5ed8b1bf09984adf9e24407637b9806153c4ca448a9525d5a3341bd1d1757b54","observation_id":"d404663f-9670-4008-8297-2b334af143b6","resolution":{"observed_at":"2026-08-05T10:46:13.932829Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.913157Z","title":"Joint modeling of chest radiographs and radiology reports for pulmonary edema assessment","venue":null,"work_id":"baa2c859-63e5-4a43-a813-e46720f487b5","year":2020},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:04.585735Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:889086db401ab7bb48109caf8a66439606ef46c76b88e0680dcc6546100e4d36","observation_id":"c143d303-e11e-4e10-bd16-89cf43f4e210","resolution":{"observed_at":"2026-08-05T10:46:13.917747Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02746","last_updated":"2025-02-19T05:59:59Z","snapshot_observed_at":"2026-08-18T15:05:50.204557Z","submitted_at":"2024-10-03T17:56:09Z","title":"Contrastive Localized Language-Image Pre-Training","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.02746","snapshot_observed_at":"2026-08-05T10:46:04.661105Z","title":"Contrastive localized language-image pre-training","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:04.661105Z"},"links":{"cited_paper":"/paper/2410.02746","citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:c4035a3be9571d47485795f637418e43976834337db07ed0d5e10e2a603be390","observation_id":"41ed5110-7327-4e3d-bafa-1a1622ef2ee5","resolution":{"observed_at":"2026-08-05T10:46:04.661105Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.895688Z","title":"A review of medical image data augmentation techniques for deep learning applications","venue":null,"work_id":"59f8838b-2c39-4eca-9c72-80287a5e5307","year":2021},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:04.776379Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:40b10943e22850d687728d0b02522e2d7249f82c10674b298942582c3ecbaf65","observation_id":"3f111173-6cf8-4d07-8f39-b1d2c4ddf619","resolution":{"observed_at":"2026-08-05T10:46:13.900852Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.879928Z","title":"Machine-learning-based multiple abnormality prediction with large-scale chest computed tomography volumes","venue":null,"work_id":"2744172b-bf54-4156-84db-79b78d5a2cec","year":2021},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:04.842478Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:0cb1830cc19d4f7f58e5f1e3e6204470b805e7fce471e6eff1163178f54ca2c5","observation_id":"f555adc7-97f4-414c-8a7b-5d9ba7d8b2d8","resolution":{"observed_at":"2026-08-05T10:46:13.884634Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18119","last_updated":"2025-03-27T17:39:55Z","snapshot_observed_at":"2026-08-18T20:17:02.937577Z","submitted_at":"2024-09-26T17:56:59Z","title":"Multi-View and Multi-Scale Alignment for Contrastive Language-Image Pre-training in Mammography","version":2},"cited_work":{"arxiv_id":"2409.18119","doi":null,"metadata_source":"pith","pith_arxiv_id":"2409.18119","snapshot_observed_at":"2026-08-05T10:46:10.496135Z","title":"Multi-View and Multi-Scale Alignment for Contrastive Language-Image Pre-training in Mammography","venue":"cs.CV","work_id":"7036e351-8912-414c-b2e1-a1749a2788ed","year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:04.962958Z"},"links":{"cited_paper":"/paper/2409.18119","citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:bef4599f2e7e576d4b5d330f74d7464c26f3dce763d84c5805252705eb7283e6","observation_id":"85df312c-28df-4027-8a32-0cf6486e38a7","resolution":{"observed_at":"2026-08-05T10:46:10.536952Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-05T10:46:05.092590Z","title":"The llama 3 herd of models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:05.092590Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:01ef52513d8ae509a9ef4e721c2b7cc43950fa0034ac2972ebbb5786483746b9","observation_id":"4feab371-5592-4901-b37e-221704c40765","resolution":{"observed_at":"2026-08-05T10:46:05.092590Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:05.191575Z","title":"Devel- oping generalist foundation models from a multimodal dataset for 3d computed tomography","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:05.191575Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:ceaec1c88b772eb2fa5f62890e062060a27e3fc182e9543aaf87176b138573b7","observation_id":"13ac3f6d-a019-4f0b-825c-d36d7414fa43","resolution":{"observed_at":"2026-08-05T10:46:05.191575Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.09685","last_updated":"2021-10-16T18:40:34Z","snapshot_observed_at":"2026-08-17T18:04:53.578114Z","submitted_at":"2021-06-17T17:37:18Z","title":"LoRA: Low-Rank Adaptation of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.09685","snapshot_observed_at":"2026-08-05T10:46:05.296884Z","title":"Lora: Low-rank adaptation of large language models","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:05.296884Z"},"links":{"cited_paper":"/paper/2106.09685","citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:f29b527ba92ea80d65b115a14bca1d637f1af3141461d6edc25b9184a815a9bc","observation_id":"ac57e9c6-b566-4905-a756-dd5818e81d92","resolution":{"observed_at":"2026-08-05T10:46:05.296884Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:05.360284Z","title":"Gloria: A multimodal global-local representation learning framework for label-efficient medical image recognition","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:05.360284Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:8080cacd9ca13925bc11c5a64640fb2c924acbec8d9257f61ff884b8ba85deac","observation_id":"aa273c70-1ad5-4176-a12e-25ee2a518a90","resolution":{"observed_at":"2026-08-05T10:46:05.360284Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.854850Z","title":"Enhancing representation in medical vision-language foun- dation models via multi-scale information extraction techniques","venue":null,"work_id":"6c5183a8-1834-492e-b0fa-22df6ac53d04","year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:05.458401Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:6bdcd9a7b6ebaf1d3588ee57ba5a8a84970746a28ccc92f6ea673fedf5747052","observation_id":"c7bf766a-ec65-41da-a482-156f75f1b8da","resolution":{"observed_at":"2026-08-05T10:46:13.859199Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.06716","last_updated":"2023-04-13T17:59:13Z","snapshot_observed_at":"2026-08-17T06:57:45.675090Z","submitted_at":"2023-04-13T17:59:13Z","title":"STU-Net: Scalable and Transferable Medical Image Segmentation Models Empowered by Large-Scale Supervised Pre-training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.06716","snapshot_observed_at":"2026-08-05T10:46:05.562057Z","title":"Stu-net: Scalable and transferable medical image segmentation models empowered by large-scale supervised pre-training","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:05.562057Z"},"links":{"cited_paper":"/paper/2304.06716","citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:f8c99e3e3a03c67544e821200453a3125c50968c3e20a965bae3c28cc5a00ddf","observation_id":"da68f205-0616-4a33-8ce2-23eb68858145","resolution":{"observed_at":"2026-08-05T10:46:05.562057Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:05.633725Z","title":"nnu-net: a self-configuring method for deep learning-based biomedical image segmentation","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:05.633725Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:291abffc33aac13d8066df582d93da9ad9bb83cb7d1995ef375306fbc1c54a29","observation_id":"4dbc9100-63b5-4b79-b258-fb0395cb7c62","resolution":{"observed_at":"2026-08-05T10:46:05.633725Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.829867Z","title":"Fool me twice: delayed diagnoses in radiology with emphasis on perpetuated errors","venue":null,"work_id":"5ac2cad2-9f58-4daf-b0d0-96acdc72e4db","year":2014},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:05.740237Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:c12bb0ebfc135b2a41353c7f4572631594a1fd8e01176aad44275741cc79c557","observation_id":"0d20a01c-58e5-4de6-8e64-28d7416906b8","resolution":{"observed_at":"2026-08-05T10:46:13.834498Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.814701Z","title":"Generating synthetic data for medical imaging","venue":null,"work_id":"6a6a4511-5c64-4ea1-8fd3-313176e39656","year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:05.837515Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:5eec848839925caecae282053a21a14a1a6f4d3f152402867fd1b9e4e958cea9","observation_id":"0ef9eb6e-3c40-4f58-8234-d3ac1161d8f4","resolution":{"observed_at":"2026-08-05T10:46:13.819212Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.798105Z","title":"Cxr-llava: a multimodal large language model for interpreting chest x-ray images","venue":null,"work_id":"7baf51ef-361b-40f1-b2b3-1554bcc5c872","year":2025},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:05.952126Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:cd61db5708bb82bac12a6246fac46097da4141a521bade5d9bc93f1c252cd29d","observation_id":"35d0a13e-01ef-4672-bf63-cc177925cee6","resolution":{"observed_at":"2026-08-05T10:46:13.803136Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:06.026153Z","title":"Llava-med: Training a large language-and-vision assistant for biomedicine in one day","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:06.026153Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:96b50139d4c6e4607fbb07c6fb33a8ce0a6922e684743086d509d582686013ca","observation_id":"71cb91cd-edae-46b0-9c62-195ae4964bf7","resolution":{"observed_at":"2026-08-05T10:46:06.026153Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.770165Z","title":"Artificial general intelligence for medical imaging analysis","venue":null,"work_id":"d74d2bbf-6de1-4036-9b54-a1d565dfb15f","year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:06.144188Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:c1d305d9305922290af4fbb682f1ef00b9d639455bcc6311ef1bc04b2a64cc36","observation_id":"6b547427-1974-4c16-a177-2ff59895bdc3","resolution":{"observed_at":"2026-08-05T10:46:13.775179Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:06.214147Z","title":"Ct-glip: 3d grounded language-image pretraining with ct scans and radiology reports for full-body scenarios","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:06.214147Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:98ad9d564b782c2b3b8d3e32a6e920691a9fbfe73510fbd3f263125a109f3737","observation_id":"9920dd3d-25af-4526-9efc-147946bfd878","resolution":{"observed_at":"2026-08-05T10:46:06.214147Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:06.347751Z","title":"Pmc-clip: Contrastive language-image pre-training using biomedical documents","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:06.347751Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:2f6dd13866154b024ffe61cf3205fc5e2ac58e85b2d2c9b3b50e6b19575ab0f9","observation_id":"045e8575-14d9-458f-90dc-c8791db2b0af","resolution":{"observed_at":"2026-08-05T10:46:06.347751Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.13523","last_updated":"2025-02-25T06:19:03Z","snapshot_observed_at":"2026-08-19T08:44:53.472470Z","submitted_at":"2024-10-17T13:11:07Z","title":"Can Medical Vision-Language Pre-training Succeed with Purely Synthetic Data?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.13523","snapshot_observed_at":"2026-08-05T10:46:06.490159Z","title":"Can medical vision-language pre-training succeed with purely synthetic data? arXiv preprint arXiv:2410.13523, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:06.490159Z"},"links":{"cited_paper":"/paper/2410.13523","citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:d096806b43ad7e99b06f0386f3155d50d1e244b0fbf63a12cfb750ad484c65f2","observation_id":"56045d27-ec37-4ba5-a6bd-9f2a0988974a","resolution":{"observed_at":"2026-08-05T10:46:06.490159Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:06.654963Z","title":"Improved baselines with visual instruction tuning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:06.654963Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:1aef12c9d945f2c3cc917fbc9d64ee4c6d5d8b64569bc76cf40a75c56efc6bad","observation_id":"7cfb17f5-68a2-4784-ae44-4289145e1d0d","resolution":{"observed_at":"2026-08-05T10:46:06.654963Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1807.03748","last_updated":"2019-01-22T18:47:12Z","snapshot_observed_at":"2026-08-14T18:53:38.574749Z","submitted_at":"2018-07-10T16:52:11Z","title":"Representation Learning with Contrastive Predictive Coding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1807.03748","snapshot_observed_at":"2026-08-05T10:46:06.818244Z","title":"Representation learning with contrastive predictive coding","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:06.818244Z"},"links":{"cited_paper":"/paper/1807.03748","citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:5984422f7075d2c28c4b1f733fb4ade8b3c46bf91f60ee30483307338f0c87c5","observation_id":"08c8a7a7-1b04-4673-b65c-8d991d660f66","resolution":{"observed_at":"2026-08-05T10:46:06.818244Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.730957Z","title":"Unsupervised medical image translation with adversarial diffusion models","venue":null,"work_id":"123926c1-9146-48d6-b139-1d7a3cf20497","year":2023},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:06.953711Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:53d69596d88f8223bc2516e1dbede0445bb3ecb531f006622db68fc8a6315f4f","observation_id":"051d5200-2466-433f-948e-fa46ba9c4c30","resolution":{"observed_at":"2026-08-05T10:46:13.736423Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.716722Z","title":"On variational bounds of mutual information","venue":null,"work_id":"8834debc-f6c2-45c4-8cca-ce6e64db09ff","year":2019},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:07.058611Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:e50665e99e4061a02f327b9e8752c51bf5a1d9790784e45de12385581bb6e698","observation_id":"5a6d5801-2d1a-43c9-bdba-d88c5914c4df","resolution":{"observed_at":"2026-08-05T10:46:13.720929Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:07.172235Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:07.172235Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:7b7805519b7f95b6af1e94ee551402df3ae35d02f1a159805ddfc9c550d11f62","observation_id":"04d8ab2e-3c09-4969-8f54-f0c86ead5096","resolution":{"observed_at":"2026-08-05T10:46:07.172235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.691468Z","title":"Study of thoracic ct in covid-19: the stoic project","venue":null,"work_id":"cb21fb89-8126-4342-a0c3-c8256e22858f","year":2021},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:07.275215Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:ba1564c06cdb0f3b56b2b57049e8ca213c9ba5e1fd4125c8da9d9bed3dbf4c21","observation_id":"904762dd-f001-4ef0-96f7-33945a9c2c06","resolution":{"observed_at":"2026-08-05T10:46:13.695888Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.675692Z","title":"Deep learning in medical image analysis","venue":null,"work_id":"5d484872-8dad-4199-8db7-026d9f2ed2b2","year":2017},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:07.401048Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:14b0b15fd3a81a74f3160840446afe38772fae3fdc4e9df0c655d3b1b389e107","observation_id":"7526b28c-0615-4fd0-8f25-a789a2a3a0d4","resolution":{"observed_at":"2026-08-05T10:46:13.680527Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.661220Z","title":"Large-scale and fine-grained vision- language pre-training for enhanced ct image understanding","venue":null,"work_id":"6610ba8f-28c9-4a11-b8a7-7119f26a0fbb","year":2025},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:07.524178Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:e0323d678607ce2639b727930f52bb1a712e0e6ef1023af055e1182add30222c","observation_id":"8cea4486-36cd-46ca-a3e0-800c3835fb65","resolution":{"observed_at":"2026-08-05T10:46:13.666394Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.647076Z","title":"Bioclip: A vision foundation model for the tree of life","venue":null,"work_id":"828b54eb-3359-4986-8e2c-b4b060118710","year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:07.697454Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:59dcfca15a2382a2af854406a8d180499e161bfc6a14caba379980accad48a05","observation_id":"d6fa8c35-2d72-432f-8f29-319783f54832","resolution":{"observed_at":"2026-08-05T10:46:13.651773Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.07971","last_updated":"2025-05-07T14:26:09Z","snapshot_observed_at":"2026-08-18T18:51:11.113709Z","submitted_at":"2023-06-13T17:59:59Z","title":"XrayGPT: Chest Radiographs Summarization using Medical Vision-Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.07971","snapshot_observed_at":"2026-08-05T10:46:07.826433Z","title":"Xraygpt: Chest radiographs summarization using medical vision-language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:07.826433Z"},"links":{"cited_paper":"/paper/2306.07971","citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:a5bc48163334fbdc6d4b93f467b8931427e41e6ac82d9f78218842087e60a5ff","observation_id":"196f8bc6-9cd3-4d50-ba35-9d041cf1c193","resolution":{"observed_at":"2026-08-05T10:46:07.826433Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.631316Z","title":"Communication errors in radiology–pitfalls and how to avoid them","venue":null,"work_id":"a180d8d2-93c9-40bd-a4bc-23fff4c39710","year":2018},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:07.970509Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:6878eafb6f5e6bcc5b19d48872fd5c2ff4579d7596c3222c414b4469f42c95fa","observation_id":"fe0aa725-db38-4b9a-9523-c2a648ecd696","resolution":{"observed_at":"2026-08-05T10:46:13.637015Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.617887Z","title":"Multi- granularity cross-modal alignment for generalized medical visual representation learning","venue":null,"work_id":"63201f56-a3b1-4e87-8ac7-429fad89e6fd","year":2022},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:08.098509Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:f08adc327314b569daebfb647bd9bb90af9e9bdeae038e046228ea9e8b242aaa","observation_id":"75a176bc-5a08-4d29-9b3e-e14cd42ca5df","resolution":{"observed_at":"2026-08-05T10:46:13.621985Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.603866Z","title":"Totalseg- mentator: robust segmentation of 104 anatomic structures in ct images","venue":null,"work_id":"bb0ffd0e-8f53-43e4-8830-dd8a6a994b01","year":2023},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:08.279430Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:2babcfd505a592ffa6c171b791e12ef846e94cdd128dfdf0d6c30acd1ba7d126","observation_id":"01ab2a9b-31b2-4afe-9136-fa8d0a2ff7f7","resolution":{"observed_at":"2026-08-05T10:46:13.608638Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.586379Z","title":"Medklip: Medical knowledge enhanced language-image pre-training for x-ray diagnosis","venue":null,"work_id":"3e050f7d-e109-429f-8c43-d1a0863884fd","year":2023},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:08.403517Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:8fde6069a25eecd2dca6eeea82faa5c5261d9b52e492a271b1a259c7d11aca6e","observation_id":"06f050c3-69ee-4cae-a1f7-a1026382da7c","resolution":{"observed_at":"2026-08-05T10:46:13.591474Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.572570Z","title":"Unimiss: Universal medical self-supervised learning via breaking dimensionality barrier","venue":null,"work_id":"faf4c2d9-3170-4213-b52e-611c5fc6e1de","year":2022},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:08.565766Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:dafa940a1809c3b18c54c2de458c08c16b4df0b07137c36f4736147fa64e69c0","observation_id":"99fd17d9-2eb1-4594-bf2f-41f070a5863a","resolution":{"observed_at":"2026-08-05T10:46:13.577184Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.16671","last_updated":"2025-11-23T00:34:43Z","snapshot_observed_at":"2026-08-17T12:17:37.741844Z","submitted_at":"2023-09-28T17:59:56Z","title":"Demystifying CLIP Data","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.16671","snapshot_observed_at":"2026-08-05T10:46:08.732381Z","title":"Demystifying clip data","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:08.732381Z"},"links":{"cited_paper":"/paper/2309.16671","citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:c7d4ef576bb88be79205f9190c890fafffbdda6abb1b7bdf0c70cb81a932e06a","observation_id":"eb497786-556b-4405-a3d4-fa736c636040","resolution":{"observed_at":"2026-08-05T10:46:08.732381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.556266Z","title":"Glipv2: unifying local- ization and vl understanding","venue":null,"work_id":"554e94c8-7617-45a4-867c-00f79bbe0523","year":2022},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:08.830792Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:86c6f0b2cd7506c287b68dbe75792e431a0b37baf9ca88333bd47ca26d01f547","observation_id":"57eb387b-0cc9-4b4b-947e-0e076d79af67","resolution":{"observed_at":"2026-08-05T10:46:13.561088Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.542047Z","title":"Biomedgpt: A unified and generalist biomedical generative pre-trained transformer for vision, language, and multimodal tasks","venue":null,"work_id":"2a0f71f7-7700-4ccb-8ccc-4b03a75be8d3","year":2023},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:08.923378Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:7fe1551a1d911282636a08a915f3bc3381d65e403f5e90bc15a5f9c4e672688b","observation_id":"0871f1a1-4930-4509-b976-2ae4035b12ae","resolution":{"observed_at":"2026-08-05T10:46:13.546935Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.16754","last_updated":"2024-04-25T17:11:37Z","snapshot_observed_at":"2026-08-17T16:25:56.941516Z","submitted_at":"2024-04-25T17:11:37Z","title":"RadGenome-Chest CT: A Grounded Vision-Language Dataset for Chest CT Analysis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.16754","snapshot_observed_at":"2026-08-05T10:46:09.008709Z","title":"Radgenome-chest ct: A grounded vision-language dataset for chest ct analysis","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.008709Z"},"links":{"cited_paper":"/paper/2404.16754","citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:90b46cb37d332b22d999ec57d1bdf34478f5f03eef841319250260a237c7a4b4","observation_id":"fd585d6f-5915-4473-b7d0-29ad8ac20d8c","resolution":{"observed_at":"2026-08-05T10:46:09.008709Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.527786Z","title":"Development of a large-scale medical visual question-answering dataset","venue":null,"work_id":"6ec14da6-fb01-41e0-b473-7e4216a47950","year":2024},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.069233Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:2641d08db78603d4d40bef9dd7b4aae78c775ab9101872dc3c4814435d6b8f72","observation_id":"5278f4c0-7029-4709-814e-c6f077d70839","resolution":{"observed_at":"2026-08-05T10:46:13.532120Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.384847Z","title":"Each of these claims is supported by theoretical analysis, ablation studies, and experimental results","venue":null,"work_id":"a77d69ee-2d30-44b7-a194-dc96df160458","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.130493Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:2e0b37d17316d99f5b853477b130b0a346b932a2ad751b84cf81e1fdc87f959b","observation_id":"52da4f73-68ac-4a46-ae3b-f2ce4f16ed22","resolution":{"observed_at":"2026-08-05T10:46:13.447833Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:13.070214Z","title":"Limitations","venue":null,"work_id":"b386c23d-a164-45d5-9c14-719f7ed8f7bb","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.166745Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:e0cbab11eabda597a4bacf6af19e0a55144661ec07d091e5bbea20b48852d07d","observation_id":"2e310e4c-9ace-41a2-81bd-58b7deb97277","resolution":{"observed_at":"2026-08-05T10:46:13.216943Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:12.693492Z","title":"Guidelines: • The answer NA means that the paper does not include theoretical results","venue":null,"work_id":"768d4dce-2d06-406b-9e06-543ffd83c124","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.222371Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:84d94f8120d61fb596a9e8cf0a32fa0cb9848720b06d69aa71f7716deeb9b533","observation_id":"655d84c8-671a-4e46-8840-a2ce1a9282b3","resolution":{"observed_at":"2026-08-05T10:46:12.864349Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:12.407620Z","title":"Guidelines: • The answer NA means that the paper does not include experiments","venue":null,"work_id":"f590abe0-655b-4c16-8c6d-f79a04a91886","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.304182Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:e6c354110b2eba412634751b46a6a7015a3b3a4242079e8139e76c5ab1d1adcd","observation_id":"4babdb4f-17c1-4e46-af10-1ab8db572560","resolution":{"observed_at":"2026-08-05T10:46:12.564751Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:12.083058Z","title":"Guidelines: • The answer NA means that paper does not include experiments requiring code","venue":null,"work_id":"e974ff14-b685-4dc6-a0aa-f92fd016602c","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.363957Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:8624a5e12deffd918ab3a4da37c0efdd71d53a8d3c6aae4b6db5c6efba9eebec","observation_id":"42aa9b52-0cf8-4efa-b0ce-3849436a662a","resolution":{"observed_at":"2026-08-05T10:46:12.259108Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:11.938282Z","title":"Guidelines: • The answer NA means that the paper does not include experiments","venue":null,"work_id":"2312a503-2ada-45ba-a54e-07423e054bd0","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.437212Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:e486a84eca6da71412009137d878e7dce7c025b99b609f59fb2e6e9c3170558a","observation_id":"81e41ff1-e8e3-43ed-9355-b8b9868615ca","resolution":{"observed_at":"2026-08-05T10:46:12.017340Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:11.875726Z","title":"Guidelines: • The answer NA means that the paper does not include experiments","venue":null,"work_id":"d63cc696-aa1c-4dd1-a9ff-72d542786618","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.500264Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:559643f2a0a4edd841e5b2039b1f3e97575927d64212b3cae0a2a0297a40e31f","observation_id":"e1fad113-1c30-4f52-88ea-f48e99eed218","resolution":{"observed_at":"2026-08-05T10:46:11.928213Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:11.747564Z","title":"Guidelines: • The answer NA means that the paper does not include experiments","venue":null,"work_id":"38c44f8b-17df-4cb3-b5f9-50ea84dde77e","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.569808Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:44bf99cc92cdaf8541d0aedf6cc6e1c65c687d7bcb06ce8ff647010866a8ba2c","observation_id":"60d98338-22c4-40ca-957a-87b29c194562","resolution":{"observed_at":"2026-08-05T10:46:11.810015Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:11.637665Z","title":"All data used in this study are from publicly available, de- identified medical datasets, and no personally identifiable information (PII) was accessed or used","venue":null,"work_id":"9494fa45-26a5-4f9d-8de2-f096a03a6b13","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.648785Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:7de3a3a9924a8dfd4eafe27fd2f1e8278ce3241f4d2cd887487819b610e9a975","observation_id":"fe82fdc4-fa2e-4d71-bcbe-0c308e46cdd5","resolution":{"observed_at":"2026-08-05T10:46:11.712125Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:11.481895Z","title":"Guidelines: 17 • The answer NA means that there is no societal impact of the work performed","venue":null,"work_id":"10e6da7a-d95e-49eb-ad96-0461c54fd63b","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.694965Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:06717bdc9c508261a1f3ead8028711144293b673309c12f953bb5b3fb486e012","observation_id":"4cbe450f-4293-472e-80de-8a150b92989b","resolution":{"observed_at":"2026-08-05T10:46:11.553592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:11.362891Z","title":null,"venue":null,"work_id":"b1f59d8c-c4b5-46c1-885c-eaab66627f1a","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.756872Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:4b58335ca21753c145d57cd2e545b320f7f5aaa2ba371e5efc5b229b3dd226b3","observation_id":"91091c49-eb63-4fbe-b291-6edfceb106fd","resolution":{"observed_at":"2026-08-05T10:46:11.419872Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:11.192746Z","title":"Guidelines: • The answer NA means that the paper does not use existing assets","venue":null,"work_id":"e3232f63-ec54-41f8-b5c4-1f07944541cb","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.813237Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:d0daa7d4e6469859a8a1738909e829ab48ed1feef922d94e3b20e976c2be33fe","observation_id":"68c2e891-feb6-46ed-b4b6-2720a5c2d289","resolution":{"observed_at":"2026-08-05T10:46:11.269223Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:11.052859Z","title":"These assets will be released with accompanying documentation upon paper acceptance","venue":null,"work_id":"174fb173-1c0d-4f80-9e90-30c6b5ab832b","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.853132Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:01e3f5a45af25994b8905fb86f85dd4047e96302fc8d7940a1b8c0f07b3f369e","observation_id":"065dba93-e8cc-466b-ba0b-2abbd75001ec","resolution":{"observed_at":"2026-08-05T10:46:11.105435Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:10.909396Z","title":"All data used are from publicly available, de-identified medical datasets with appropriate licenses and do not involve any direct interaction with individuals","venue":null,"work_id":"42424386-ceb6-4c3e-ae3d-79cfad7a5f5a","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.939998Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:a8f41299ae80b3ff4fce6b23ae6e7cd787736c7f0d3dfcd28b1d7b1df013802b","observation_id":"5cf22d1e-5be8-4e8f-9c0d-983ca3baaf39","resolution":{"observed_at":"2026-08-05T10:46:10.989629Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:10.742870Z","title":"Therefore, IRB approval was not required","venue":null,"work_id":"0ad134fd-5474-40d8-979d-1ac3458b84dc","year":null},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.978316Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:29808dca354339215200bab0cba2ef25afa64efd4f8e080016f1514b7581ce35","observation_id":"304aea77-f964-4354-8c2a-855d2a6a01d3","resolution":{"observed_at":"2026-08-05T10:46:10.820337Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T10:46:10.616630Z","title":"Answer: [Yes] Justification: Large language models such as GPT-4o and Qwen2.5, were used to rewrite radiology reports for improving semantic clarity during pretraining","venue":null,"work_id":"7a5a616c-b517-4d6f-81a9-c249099b79a0","year":2025},"citing_paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-05T10:46:09.996574Z"},"links":{"citing_paper":"/paper/2509.03800"},"observation_digest":"sha256:44984d5b60f55d480929253ae8c7fcde9cf712e0ebc8c157621aeee345fc03a4","observation_id":"2905636d-5717-4d58-9d1d-4f6aeb97aac9","resolution":{"observed_at":"2026-08-05T10:46:10.677975Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2509.03800","last_updated":"2025-09-04T01:28:44Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-17T12:18:24.366486Z","submitted_at":"2025-09-04T01:28:44Z","title":"MedVista3D: Vision-Language Modeling for Reducing Diagnostic Errors in 3D CT Disease Detection, Understanding and Reporting"},"reference_resolution":{"displayed":60,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":20,"verified_exact":1,"verified_fuzzy":39},"total_outbound_references":60},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 60 of 60 outbound references and 2 inbound Pith citation observations for arXiv:2509.03800."}