{"as_of":"2026-08-22T14:11:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6679a2ecfc665bd76a5e041ebed29c25892913b1f7d7cb4166b7521bc41285cf","coverage":[{"denominator":60,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":60,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T04:50:45.382132Z","state":"measured"},{"denominator":61,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":61,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-22T06:32:14.747728+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T04:50:45.050116Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-06T04:50:45.808704Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"cited_work":{"arxiv_id":"2508.02951","doi":null,"metadata_source":"pith","pith_arxiv_id":"2508.02951","snapshot_observed_at":"2026-08-06T04:50:45.808704Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","venue":"cs.AI","work_id":"a3e84c01-cafa-4817-b836-7dc68c8a1ed5","year":2025},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.050116Z"},"links":{"cited_paper":"/paper/2508.02951","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:f6750ef1e74ad034dad1b66e69ccee339b2ee89e7f9e9d141668717cc1e99fd8","observation_id":"47c153a4-45ee-4a12-91f4-2dc59ba0543a","resolution":{"observed_at":"2026-08-06T04:50:45.814234Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2508.02951/citation-record","integrity":"/paper/2508.02951/integrity","json":"/paper/2508.02951/citation-record.json","paper":"/paper/2508.02951"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"cited_work":{"arxiv_id":"2508.02951","doi":null,"metadata_source":"pith","pith_arxiv_id":"2508.02951","snapshot_observed_at":"2026-08-06T04:50:45.808704Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","venue":"cs.AI","work_id":"a3e84c01-cafa-4817-b836-7dc68c8a1ed5","year":2025},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.050116Z"},"links":{"cited_paper":"/paper/2508.02951","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:f6750ef1e74ad034dad1b66e69ccee339b2ee89e7f9e9d141668717cc1e99fd8","observation_id":"47c153a4-45ee-4a12-91f4-2dc59ba0543a","resolution":{"observed_at":"2026-08-06T04:50:45.814234Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.552868Z","title":"These models combine text and image understanding and are typ- ically evaluated using visual question answering (VQA) tasks","venue":null,"work_id":"5939989f-59fa-42eb-8e59-5c86742b092e","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.056467Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:6f7d816ef6de95c272e7d08ba1e6bc2b13940f575c9aca3557010dd5b25d8685","observation_id":"710b94de-7a56-4593-a1ca-35d63f90de9a","resolution":{"observed_at":"2026-08-06T04:50:46.557554Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.521820Z","title":"Perception enables clinicians to ex- tract key visual features before engaging in more complex 2 Figure 2","venue":null,"work_id":"92427927-f619-4631-b96a-dc6cec0d0e0e","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.067950Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:e84a0b05152ae2aca83511116f46c51b12fbf669a5ac14d75aaf9803744a9d9e","observation_id":"6af40122-e672-4103-8449-829ae1cbafde","resolution":{"observed_at":"2026-08-06T04:50:46.527645Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.447521Z","title":null,"venue":null,"work_id":"3f43bac4-d776-4df7-910b-a8d196fc2a04","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.090608Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:b882cd43f5187a4f1257735f9cde37cb122f028a5491d76b9ed2393ff6a52f01","observation_id":"e4d88d05-d7f5-42a3-9f80-f09ff5e32940","resolution":{"observed_at":"2026-08-06T04:50:46.452654Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.504597Z","title":null,"venue":null,"work_id":"9540ef43-5fba-4f2d-85b7-236043ceb559","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.073600Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:27b5d6a40b826211640c2225d2251c5ed9b3190be39412d2b115513d7e8db137","observation_id":"847fb887-7ac6-46d8-bfa2-1993dbbd1d71","resolution":{"observed_at":"2026-08-06T04:50:46.510041Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.487007Z","title":null,"venue":null,"work_id":"9ef2dae0-8738-4557-81fb-fdff30d3f0bb","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.079017Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:fa0f3011b0f31ac314e8bf068adfe1f99a27021661c7cf540ad72925e4cf083a","observation_id":"7f86ac3c-fbac-475f-8c64-101e529991c7","resolution":{"observed_at":"2026-08-06T04:50:46.491579Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.467060Z","title":"We mimic this form of prompting, by leveraging visual cues like points/dots to spatially prompt the models when answering the specified question [43, 50]","venue":null,"work_id":"00e1abd7-7cbd-4abe-9298-5057cf71377f","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.084623Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:1045f6a77db4620d14a0ede75cec6870db970113cce025f9f3ae272a5befb09a","observation_id":"54d22850-cdf2-4baa-ae46-bf76d3454179","resolution":{"observed_at":"2026-08-06T04:50:46.474122Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.429810Z","title":"Current models including leading generalist and domain-specialized systems perform far below human levels on perceptual tasks that clinicians solve effortlessly (best: 65% vs","venue":null,"work_id":"c6f68535-d5bf-4651-86bd-756e0d265928","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.096640Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:6c013501d03c44676af5bdb1aff54d960544b4aa1484d543065bca9a999f93c7","observation_id":"f0ade7de-781a-417f-b728-1541bb262b80","resolution":{"observed_at":"2026-08-06T04:50:46.435678Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.860662Z","title":null,"venue":null,"work_id":"42e73b97-b610-4db8-aecc-93a3a52d281d","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.371172Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:f9aabc40f3012c416c4dc8664c9cb9e47c86de19a90d3de5439fea57d8a6be77","observation_id":"6280ff95-d0ab-431c-bc7b-415cae1e46b7","resolution":{"observed_at":"2026-08-06T04:50:45.865994Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.843318Z","title":null,"venue":null,"work_id":"5f52f24b-c565-47b4-8f6e-5bbeb9029ef1","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.376206Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:88f9b64f393eb098dd0f7678c19959fd798fff1901c184ad06dde1440653a0e0","observation_id":"d1fcc25d-bb44-4297-9949-9cef024da10d","resolution":{"observed_at":"2026-08-06T04:50:45.848416Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.826741Z","title":"Small Specialized models We train small specialized models for some of the tasks with sizeable train sets from the original dataset used to construct the task","venue":null,"work_id":"7e3bf3df-067e-4565-9c22-aa0df3851879","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.382132Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:5b608137601178521b23e69dd547ea0975832b7902e814e6fe88afc44375e7ab","observation_id":"382d428a-6816-4345-9815-1370576706e9","resolution":{"observed_at":"2026-08-06T04:50:45.832023Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.412516Z","title":"Omnimedvqa: A new large- scale comprehensive evaluation benchmark for medical lvlm","venue":null,"work_id":"87ffcca9-c41f-442e-817c-64529ea9993f","year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.101847Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:f0bb176413bc2015836175832452450476b34e9acefd3d8bdac88f5a7b4267e5","observation_id":"29967050-21e7-49b3-b916-a48d40f29a5d","resolution":{"observed_at":"2026-08-06T04:50:46.417870Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21276","last_updated":"2024-10-25T17:43:01Z","snapshot_observed_at":"2026-08-15T14:02:47.366139Z","submitted_at":"2024-10-25T17:43:01Z","title":"GPT-4o System Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21276","snapshot_observed_at":"2026-08-06T04:50:45.107165Z","title":"Gpt-4o system card","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.107165Z"},"links":{"cited_paper":"/paper/2410.21276","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:933bc29a6af030999f4d6ff3b0eb70c1538ad39123a7e281d8219dbd154e5743","observation_id":"3fd89b40-4785-4f9c-9255-30b339a8010b","resolution":{"observed_at":"2026-08-06T04:50:45.107165Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.388840Z","title":"The diagnostic impact of contrast-enhanced computed tomography (cect) in evaluating lymph node involvement in colorectal cancer: a comprehen- sive review","venue":null,"work_id":"839d954e-baff-48e8-8b00-c349d3636f05","year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.113559Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:aeb963bdf76ea56c036b1a17ba4824f29077f31d38aa49569ea5ac5abc6abd20","observation_id":"f3429a48-d7ad-4bf8-9555-7e26f08a7c10","resolution":{"observed_at":"2026-08-06T04:50:46.401814Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.359124Z","title":"Chexpert: A large chest radiograph dataset with uncertainty labels and expert comparison","venue":null,"work_id":"dd1a2fac-f359-4009-9d53-a55da79e4704","year":2019},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.119177Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:57c5ce1da9df69798c316bb12b73a3505db2465e33a020a6b7b4e7b080b6ef56","observation_id":"4fa8b32c-08b4-4b95-8ee6-b032d571f42c","resolution":{"observed_at":"2026-08-06T04:50:46.369421Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.341043Z","title":"A dataset of clinically generated visual questions and answers about radiology images","venue":null,"work_id":"d97d9c4c-7884-4d6a-9f1c-ddb59851b406","year":2018},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.124916Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:42fa8e34ffb03a8623971e2bc482468efc8f02029272e924c91ee3159f47a0f1","observation_id":"4702e604-2bdc-4114-abe0-c795ffa19983","resolution":{"observed_at":"2026-08-06T04:50:46.346514Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.322014Z","title":"Llava-onevision: Easy visual task transfer, 2024","venue":null,"work_id":"b7588736-1298-4b0a-9196-9d04f0910d35","year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.130611Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:7739537a4952b48d8a1a93b39bac15aa137f1190ade45e27e06bddc9e0440778","observation_id":"42ef0246-05b1-41a5-90c1-25d056987509","resolution":{"observed_at":"2026-08-06T04:50:46.328216Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.03326","last_updated":"2024-10-26T16:35:13Z","snapshot_observed_at":"2026-08-22T00:56:08.009523Z","submitted_at":"2024-08-06T17:59:44Z","title":"LLaVA-OneVision: Easy Visual Task Transfer","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.03326","snapshot_observed_at":"2026-08-06T04:50:45.135890Z","title":"Llava-onevision: Easy visual task transfer","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.135890Z"},"links":{"cited_paper":"/paper/2408.03326","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:00322d67942e7c34bfb26bcc6e42c0670256a8f88fed20317dff388d52fe6266","observation_id":"8008b3fc-6a3b-4fff-87cd-a44e45d15971","resolution":{"observed_at":"2026-08-06T04:50:45.135890Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.141848Z","title":"Llava-med: Training a large language- and-vision assistant for biomedicine in one day","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.141848Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:b1c778496e3629ec7a762a96862771437449febdd7ccdfe77bf0401a48d5d9af","observation_id":"80da6ea6-1f71-466c-9550-417b77da3a1f","resolution":{"observed_at":"2026-08-06T04:50:45.141848Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.294902Z","title":"Slake: A semantically-labeled knowledge- enhanced dataset for medical visual question answering","venue":null,"work_id":"9825c6d3-04ad-4451-8bac-3d2e12abe957","year":2021},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.146966Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:ddccefb71e6dc446f9a1299af835a60ed0538c914308b0e5ec1d46641ba422ef","observation_id":"0f1634c8-b201-43e2-94db-5372c0df43d5","resolution":{"observed_at":"2026-08-06T04:50:46.300661Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03744","last_updated":"2024-05-15T19:22:44Z","snapshot_observed_at":"2026-08-15T02:35:59.111911Z","submitted_at":"2023-10-05T17:59:56Z","title":"Improved Baselines with Visual Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03744","snapshot_observed_at":"2026-08-06T04:50:45.152773Z","title":"Improved baselines with visual instruction tuning, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.152773Z"},"links":{"cited_paper":"/paper/2310.03744","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:594a0946b3a26f982274f5466ebed7f7ac671ef284b403d07bb1a55599744ef6","observation_id":"3039d259-df27-4083-9eb5-38244175a009","resolution":{"observed_at":"2026-08-06T04:50:45.152773Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11217","last_updated":"2024-11-29T02:50:45Z","snapshot_observed_at":"2026-08-16T14:17:47.173829Z","submitted_at":"2024-02-17T08:04:23Z","title":"A Spectrum Evaluation Benchmark for Medical Multi-Modal Large Language Models","version":2},"cited_work":{"arxiv_id":"2402.11217","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.11217","snapshot_observed_at":"2026-08-06T04:50:45.731150Z","title":"A Spectrum Evaluation Benchmark for Medical Multi-Modal Large Language Models","venue":"cs.CL","work_id":"2ffda452-88b0-428b-835f-735d7579f115","year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.158675Z"},"links":{"cited_paper":"/paper/2402.11217","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:be47cb0edac08be35407bfcac6a7f1ac0392a7336453e10df78f9413138ea55b","observation_id":"503ad637-3270-4f14-acbd-35c7b1d9855a","resolution":{"observed_at":"2026-08-06T04:50:45.737508Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.276786Z","title":"The coding of roentgen images for computer analysis as applied to lung cancer","venue":null,"work_id":"4c3de253-3a1d-45dc-9643-0bef409d1915","year":1963},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.164214Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:f1bd6e90c762d22d06515218d6a26391b62e142e4e7c0c699b58cd0d53fc99b8","observation_id":"be03e375-8617-4b08-8dc6-5727ce900fe7","resolution":{"observed_at":"2026-08-06T04:50:46.282228Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.258578Z","title":"Lu, Bowen Chen, Drew F","venue":null,"work_id":"0f740f7a-a06b-45a8-9bff-1ccfcc936a48","year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.170000Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:29d5bf6734ab8264b05079799a45d5781d643b9de332fce2c3e982a9d0d04b60","observation_id":"b9d79b9d-06b9-4d12-a317-a8d6f17014b8","resolution":{"observed_at":"2026-08-06T04:50:46.265144Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.241277Z","title":"Lung-rads: pushing the limits","venue":null,"work_id":"307b262f-b4c2-4ca2-af67-f76f41c0bfdb","year":1975},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.175275Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:1731dc863121b8df7c3256b3e2af25ef52c72f2195558abb3bac98ac49465a65","observation_id":"3fa64eb0-1fd8-4a7c-a042-bc6278f9868b","resolution":{"observed_at":"2026-08-06T04:50:46.246274Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.00164","last_updated":"2023-11-30T19:55:51Z","snapshot_observed_at":"2026-08-19T05:53:51.283979Z","submitted_at":"2023-11-30T19:55:51Z","title":"Towards Accurate Differential Diagnosis with Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.00164","snapshot_observed_at":"2026-08-06T04:50:45.181592Z","title":"Towards accurate differential diagnosis with large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.181592Z"},"links":{"cited_paper":"/paper/2312.00164","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:a8be75b03b51e20a2f7846f743dea39db887cd1368b4008aaab42694ca10a6e3","observation_id":"dca4a1ca-1a4c-4c61-89a8-22c858a68100","resolution":{"observed_at":"2026-08-06T04:50:45.181592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.224030Z","title":"Med-flamingo: a multimodal medical few-shot learner","venue":null,"work_id":"c47eb94c-4759-40ba-ba90-2e82a8b31f1c","year":2023},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.188471Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:7c677a4b1f99d1dd836b68c0f2f33a729b6012638cc030dd4f59b333b9f60d22","observation_id":"4614b4aa-68e4-4136-8e44-eb087144fc66","resolution":{"observed_at":"2026-08-06T04:50:46.229278Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.207792Z","title":"Interactions of perceptual and conceptual pro- cessing: Expertise in medical image diagnosis","venue":null,"work_id":"def99191-a69c-40c6-90f3-baf87af7d95d","year":2008},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.194078Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:59647c7433bcb85aac24a88ab63ef0a0bb84b5f88de31555be0590a69863a51b","observation_id":"7a1be3ce-2ba3-40d0-9881-1351627d55e9","resolution":{"observed_at":"2026-08-06T04:50:46.212954Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.08704","last_updated":"2025-04-09T17:42:01Z","snapshot_observed_at":"2026-08-17T12:48:31.714505Z","submitted_at":"2024-08-16T12:32:44Z","title":"Beyond the Hype: A dispassionate look at vision-language models in medical scenario","version":2},"cited_work":{"arxiv_id":"2408.08704","doi":null,"metadata_source":"pith","pith_arxiv_id":"2408.08704","snapshot_observed_at":"2026-08-06T04:50:45.683550Z","title":"Beyond the Hype: A dispassionate look at vision-language models in medical scenario","venue":"cs.CV","work_id":"1e4103b7-cabb-44c0-84d6-4799f990e7d0","year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.199029Z"},"links":{"cited_paper":"/paper/2408.08704","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:1c5502b6a63e2ad3259ce5b2dbdd541323cc4446433fc9f666c53046cfbb807c","observation_id":"0910e13e-2cab-4d99-92ef-3dbdd9f931a7","resolution":{"observed_at":"2026-08-06T04:50:45.691969Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.190405Z","title":"Video-based ai for beat-to-beat assessment of cardiac function","venue":null,"work_id":"9e12af26-6882-464b-be80-4bf4c1330e56","year":2020},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.205271Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:e1e741e1f51eb3602888ddcc9ed09f864cebc4a2c59c04b00b8e47d440b61b73","observation_id":"3dbd8851-1a05-4103-8386-467712ece06c","resolution":{"observed_at":"2026-08-06T04:50:46.196722Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.170979Z","title":"Kvasir: A multi-class image dataset for computer aided gastrointestinal disease detection","venue":null,"work_id":"b61254ed-a84a-4d35-bcff-c070d22eb499","year":2017},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.210922Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:dd06ece4ddf09391995233d3d9f8a977eea8c5884f0a3303bfe08c3f52c17f38","observation_id":"bb599bb3-104d-4f89-b8eb-12d27c251d4f","resolution":{"observed_at":"2026-08-06T04:50:46.177764Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.215891Z","title":"Learning transferable visual models from natural language supervision, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.215891Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:d87124bca8aeb19f64f89062c31821a12acab46681e12470e664f0e0d30d2a8c","observation_id":"4c63d650-f3ba-46a8-9471-4d7e92744df8","resolution":{"observed_at":"2026-08-06T04:50:45.215891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.220710Z","title":"Mul- timedeval: A benchmark and a toolkit for evaluating medical vision-language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.220710Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:7f48e8c347c28186191bbd20a447b48ad5ae3027cb4978df75cee76bcc1b31dc","observation_id":"95b409ea-16bb-41c5-9bbe-3c8adf40784d","resolution":{"observed_at":"2026-08-06T04:50:45.220710Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.134466Z","title":"Hematoxylin and eosin-stained whole slide image dataset annotated for skin tissue segmentation","venue":null,"work_id":"69528176-fc21-4b25-89df-8cfaeaf29938","year":2025},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.225802Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:f9c7dad18806f958c168d638045318910873b49b5efaca2e06f5a9e186578a21","observation_id":"dbf45989-4f88-43f9-bbfb-004d4022d6f9","resolution":{"observed_at":"2026-08-06T04:50:46.142357Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.15477","last_updated":"2025-05-20T23:20:36Z","snapshot_observed_at":"2026-08-16T13:16:07.432594Z","submitted_at":"2024-09-23T18:59:37Z","title":"MediConfusion: Can you trust your AI radiologist? Probing the reliability of multimodal medical foundation models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.15477","snapshot_observed_at":"2026-08-06T04:50:45.231663Z","title":"Mediconfusion: Can you trust your ai radiologist? probing the reliability 10 of multimodal medical foundation models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.231663Z"},"links":{"cited_paper":"/paper/2409.15477","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:26b38eef41dd8eb9bc737a27ea52f267a94fc19975c223095296bb7ad3760fde","observation_id":"cfd50d5a-86c3-4101-adee-4ff98f625024","resolution":{"observed_at":"2026-08-06T04:50:45.231663Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.114892Z","title":"Quilt-llava: Visual instruction tuning by extracting localized narratives from open-source histopathology videos","venue":null,"work_id":"c17fd453-a170-43e7-a05b-a50f3a538612","year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.237379Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:0ebed0d9ce62454a5b09a74fe35d70b34fc51332e20b8c83736d27d6f2bd54ec","observation_id":"46cd4533-fb17-4e19-bb77-767cfc68408f","resolution":{"observed_at":"2026-08-06T04:50:46.121877Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.097130Z","title":"Diagnostic ultrasound imaging: inside out","venue":null,"work_id":"2ceb4a8c-ea55-4952-9156-063779e47277","year":2013},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.242452Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:24af60ca61392d61aa2c1606bb4b7203edf66b1d8ea01b658c5c7bbdfba441a7","observation_id":"72043f7c-db68-416f-b52f-7337a3562546","resolution":{"observed_at":"2026-08-06T04:50:46.102569Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05530","last_updated":"2024-12-16T17:39:39Z","snapshot_observed_at":"2026-08-14T18:15:53.516440Z","submitted_at":"2024-03-08T18:54:20Z","title":"Gemini 1.5: Unlocking multimodal understanding across millions of tokens of context","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.05530","snapshot_observed_at":"2026-08-06T04:50:45.249202Z","title":"Gemini 1.5: Unlocking mul- timodal understanding across millions of tokens of context","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.249202Z"},"links":{"cited_paper":"/paper/2403.05530","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:70708379507eef50140dc334fcb8680bfebe7f43ee7c76be1b1359ba88578913","observation_id":"db8c5afb-1320-454f-9fe9-b77cbb2332d1","resolution":{"observed_at":"2026-08-06T04:50:45.249202Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.078406Z","title":"What clinicians want: contextualizing explainable machine learning for clinical end use","venue":null,"work_id":"ea41f875-7004-409f-b47e-8f1f5b2b2014","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.255259Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:00674c2b3073faaae9498463803f5b425c43bec9c14851620decea8271b91baf","observation_id":"66423372-b8a5-4212-9344-495965152d86","resolution":{"observed_at":"2026-08-06T04:50:46.084094Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-07T12:56:43.323460Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-08-06T04:50:45.261256Z","title":"Llama 2: Open foundation and fine-tuned chat models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.261256Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:5c659524bb6d0b3a24c79db43fb3eab2527c7f5514caaad101cd26cb6406b6ee","observation_id":"92afd454-277d-4b91-8d57-a9cbcf723d8f","resolution":{"observed_at":"2026-08-06T04:50:45.261256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.062111Z","title":"Corrado, Yossi Matias, Karan Singhal, Pete Florence, Alan Karthikesalingam, and Vivek Natarajan","venue":null,"work_id":"2c77c79f-fd7c-47cd-9451-62086ff5c2d3","year":2023},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.268906Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:7c42d31643a257421b907165b82ad80179f2f0e38491fd8de855666b7d06c0a8","observation_id":"f998d0a9-5378-44a1-8740-88c68dbde60a","resolution":{"observed_at":"2026-08-06T04:50:46.067523Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.12671","last_updated":"2025-03-05T08:23:04Z","snapshot_observed_at":"2026-08-19T05:50:08.132110Z","submitted_at":"2025-02-18T09:21:12Z","title":"Baichuan-M1: Pushing the Medical Capability of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.12671","snapshot_observed_at":"2026-08-06T04:50:45.273861Z","title":"Baichuan-m1: Pushing the medical capability of large language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.273861Z"},"links":{"cited_paper":"/paper/2502.12671","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:b8f1a0ffe7cf63eefa50e411cd3d5d54671b3077350009f1b75644d5bc31ae03","observation_id":"00bfc95d-4598-4802-a203-d9aea1a245c1","resolution":{"observed_at":"2026-08-06T04:50:45.273861Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.279006Z","title":"Chestx- ray8: Hospital-scale chest x-ray database and benchmarks on weakly-supervised classification and localization of common thorax diseases","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.279006Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:2ce0635304569cbe838b553816ee9d9e20800c8f68d31ae5bd950cbe56d4d5a5","observation_id":"0a69827b-6308-4eb3-9760-2852c2d11ef6","resolution":{"observed_at":"2026-08-06T04:50:45.279006Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.034344Z","title":"Detection of radiographic abnor- malities in mammograms by means of optical scanning and computer analysis","venue":null,"work_id":"ee638058-66c8-4dc8-967e-6506fa661c43","year":1967},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.283634Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:569f99c22beefb6da618c5fe5e2a0c3826be92c9782341c9a067371bbd27ba1a","observation_id":"fd19d5cf-bcf1-4ea1-a4d5-e45f44bcb0ff","resolution":{"observed_at":"2026-08-06T04:50:46.040026Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.09909","last_updated":"2023-12-04T14:13:35Z","snapshot_observed_at":"2026-08-16T14:51:59.767642Z","submitted_at":"2023-10-15T18:32:27Z","title":"Can GPT-4V(ision) Serve Medical Applications? Case Studies on GPT-4V for Multimodal Medical Diagnosis","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.09909","snapshot_observed_at":"2026-08-06T04:50:45.288838Z","title":"Can gpt-4v (ision) serve medical ap- plications? case studies on gpt-4v for multimodal medical diagnosis","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.288838Z"},"links":{"cited_paper":"/paper/2310.09909","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:cafcdbca97fa5cdcd7ce9b23a26073618093c01d3f8254ed0c4e04a6bb5ee9a6","observation_id":"12078155-60aa-4a78-b8e6-edd367152b85","resolution":{"observed_at":"2026-08-06T04:50:45.288838Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.02463","last_updated":"2023-11-16T12:38:46Z","snapshot_observed_at":"2026-08-16T15:10:43.905696Z","submitted_at":"2023-08-04T17:00:38Z","title":"Towards Generalist Foundation Model for Radiology by Leveraging Web-scale 2D&3D Medical Data","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.02463","snapshot_observed_at":"2026-08-06T04:50:45.295335Z","title":"Towards generalist foundation model for radiology by leveraging web-scale 2d&3d medical data","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.295335Z"},"links":{"cited_paper":"/paper/2308.02463","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:249399eb2f493c8f32f323b009a500b8297dd0788ab36830b8a0e293ab682b6e","observation_id":"360772cb-1472-4e1e-934d-18f1316792ab","resolution":{"observed_at":"2026-08-06T04:50:45.295335Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.538464Z","title":"MediConfusion [ 49] probes failure modes on visually dissimilar image pairs","venue":null,"work_id":"c2f5259e-5d0a-4ba6-9237-49a64232ef34","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.061737Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:dda4ae6eadae17d9c3dd05610af3f3b66f6056512037c43c0fe3832c91a29c4b","observation_id":"a26b6eee-0fca-48a4-a7db-b5e9c62f8bbf","resolution":{"observed_at":"2026-08-06T04:50:46.543033Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06007","last_updated":"2024-11-03T16:54:14Z","snapshot_observed_at":"2026-08-21T23:06:47.098719Z","submitted_at":"2024-06-10T04:07:09Z","title":"CARES: A Comprehensive Benchmark of Trustworthiness in Medical Vision Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06007","snapshot_observed_at":"2026-08-06T04:50:45.300767Z","title":"Cares: A comprehensive benchmark of trustwor- thiness in medical vision language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.300767Z"},"links":{"cited_paper":"/paper/2406.06007","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:0bc672ac3f2b9d34a5979555f800a88dde0544a5fd705e1c5945ce63793cba99","observation_id":"ec15513e-c7ba-42c5-a280-4db773bee185","resolution":{"observed_at":"2026-08-06T04:50:45.300767Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02900","last_updated":"2025-07-10T08:33:52Z","snapshot_observed_at":"2026-08-18T16:54:45.587374Z","submitted_at":"2024-08-06T02:09:35Z","title":"MedTrinity-25M: A Large-scale Multimodal Dataset with Multigranular Annotations for Medicine","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02900","snapshot_observed_at":"2026-08-06T04:50:45.306820Z","title":"Medtrinity-25m: A large-scale multimodal dataset with multigranular annotations for medicine","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.306820Z"},"links":{"cited_paper":"/paper/2408.02900","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:2863b65767f51b5e0c4bc28879c03ea7deb1f057c2ff43f2487b00ac1ac5da3c","observation_id":"42125ddc-4570-45f7-b84c-5dba5e666f68","resolution":{"observed_at":"2026-08-06T04:50:45.306820Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:46.015124Z","title":"Depth anything: Unleashing the power of large-scale unlabeled data","venue":null,"work_id":"c0be3c64-aa63-4c8a-b8f9-0e3bc2f6e8ec","year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.312871Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:346cc96b78705a03a1b2dcc49d52579281d26d7e0e05e5d8a39c8eadc222e188","observation_id":"ea407e41-f8ca-4185-a42f-eee8bff03f30","resolution":{"observed_at":"2026-08-06T04:50:46.022706Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.03162","last_updated":"2024-05-06T04:44:22Z","snapshot_observed_at":"2026-08-16T13:55:11.369835Z","submitted_at":"2024-05-06T04:44:22Z","title":"Advancing Multimodal Medical Capabilities of Gemini","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.03162","snapshot_observed_at":"2026-08-06T04:50:45.317935Z","title":"Advancing multimodal medical capabilities of gemini","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.317935Z"},"links":{"cited_paper":"/paper/2405.03162","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:67004bc022085b281166085233d33984280b0062787ccc74dc2c835abc9fe776","observation_id":"d4185d3c-b5de-49c8-9e27-4923254fee06","resolution":{"observed_at":"2026-08-06T04:50:45.317935Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.996276Z","title":"Gmai-mmbench: A comprehensive multimodal evaluation benchmark towards general medical ai","venue":null,"work_id":"4e836cf6-68cf-4090-a36a-0aa99bb1c1bf","year":2024},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.323714Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:d138837ea39a4c22ab37e4ca9fdb86c81f2af493fa8a066c8e7e2150b06666de","observation_id":"82cbc2ec-e34a-479a-ac03-8f78de3111f6","resolution":{"observed_at":"2026-08-06T04:50:46.002495Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.10415","last_updated":"2024-09-08T01:04:35Z","snapshot_observed_at":"2026-08-21T18:14:57.901749Z","submitted_at":"2023-05-17T17:50:16Z","title":"PMC-VQA: Visual Instruction Tuning for Medical Visual Question Answering","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.10415","snapshot_observed_at":"2026-08-06T04:50:45.328920Z","title":"Pmc-vqa: Visual in- struction tuning for medical visual question answering","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.328920Z"},"links":{"cited_paper":"/paper/2305.10415","citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:87e4c70ef052bd186aa0a4b5f2237c881bccba35143034bdea9f12ca6ad39d80","observation_id":"d7768017-fca9-49a6-b8dd-6a6de030ffa8","resolution":{"observed_at":"2026-08-06T04:50:45.328920Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.978715Z","title":null,"venue":null,"work_id":"be6d3f99-b2c6-4cb1-b17a-6fd508d3bbfc","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.334499Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:145064488fc7c4be49278a0b42f8d25a9abc4a467f50a7b31679a091a17d3b6e","observation_id":"8cfcdeb0-8fa5-4ce3-9f83-fe4581342e22","resolution":{"observed_at":"2026-08-06T04:50:45.983909Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.962860Z","title":null,"venue":null,"work_id":"de347a75-473e-4c0b-a20b-e94c65d6c996","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.339961Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:00ef131d57d1b213a2600ea099720eff904fa0a24abebe98c4113294774b7c98","observation_id":"22370e71-3979-4e4d-8cf4-dbd23db8f2b5","resolution":{"observed_at":"2026-08-06T04:50:45.968069Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.946284Z","title":null,"venue":null,"work_id":"88aae89b-48e9-4400-b6d0-9bd299c24076","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.346086Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:833efa0f05f26a132e3f78ba059a411b982c1e2cf2e174beb816736f04e18397","observation_id":"34fc5d89-f78b-4671-aa47-456e4ddfc969","resolution":{"observed_at":"2026-08-06T04:50:45.951271Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.930437Z","title":null,"venue":null,"work_id":"6cfcebd6-1b09-4f19-b594-92e09756f5a8","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.351161Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:e91c2b254aa26c661989ff9dd22f7a60b4c5de9d4d36f05d13173e58978a5750","observation_id":"04848663-0141-4fc7-b3ed-9f935c04c9e9","resolution":{"observed_at":"2026-08-06T04:50:45.935521Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.912825Z","title":"Here we use two versions as well, the 0.5B parameterized model, and 7B parameter- ized models","venue":null,"work_id":"eb39fd17-9e7f-4471-a437-2e6940f16359","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.356130Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:28ca1ea4465b3358ecb33129d68ab3780b4bed84409312b639e95e1efe2237f2","observation_id":"710a4900-c0a2-41d9-bbfd-5c4fa6eb15e1","resolution":{"observed_at":"2026-08-06T04:50:45.918541Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.895111Z","title":"Specifically, we prompt it with three questions and answers from PMC-VQA benchmark","venue":null,"work_id":"abba655a-8709-4953-819b-91c460004aa2","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.361108Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:cbaed9b4f6569afdd38025ea99fa07886034de437261ca1cfc70f08f964f73fb","observation_id":"b2e27d68-1e17-4fcc-bf01-fa012d90f9a5","resolution":{"observed_at":"2026-08-06T04:50:45.900362Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:50:45.877901Z","title":null,"venue":null,"work_id":"1d8f4471-b0f3-45b3-a2da-c8eed910a936","year":null},"citing_paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-06T04:50:45.366141Z"},"links":{"citing_paper":"/paper/2508.02951"},"observation_digest":"sha256:51ea54acfd09df63ff9e42f96b1e332e5588d6760c1f4fb908f4ed41feba667d","observation_id":"dbc38cef-16cb-40c4-93ca-4cd9ceffcb02","resolution":{"observed_at":"2026-08-06T04:50:45.882897Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2508.02951","last_updated":"2025-08-04T23:19:18Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-15T00:29:03.219436Z","submitted_at":"2025-08-04T23:19:18Z","title":"MedBLINK: Probing Basic Perception in Multimodal Language Models for Medicine"},"reference_resolution":{"displayed":60,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":4,"unresolved":23,"verified_exact":3,"verified_fuzzy":29},"total_outbound_references":60},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"thesis":"As of 22 August 2026, this Paper Citation Record lists 60 of 60 outbound references and 1 inbound Pith citation observation for arXiv:2508.02951."}