{"as_of":"2026-08-16T03:02:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7102a925631648d21cd001445199fb9b2401dc5762fec09d1db4357d1b89146e","coverage":[{"denominator":52,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":52,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T15:28:49.334512Z","state":"measured"},{"denominator":52,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":52,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-15T06:32:42.880941+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2507.15717/citation-record","integrity":"/paper/2507.15717/integrity","json":"/paper/2507.15717/citation-record.json","paper":"/paper/2507.15717"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:57.227525Z","title":"Attention is All you Need [Internet]","venue":null,"work_id":"c0b30817-2a27-4ec9-9b75-fea23df0c8c2","year":2017},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.035082Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:08908c167cacc9f7fb93c22bcd244a533a6859b7143668b09e95b0de7faa230f","observation_id":"41b878ad-7ebf-4384-aebd-e6bc462a0da6","resolution":{"observed_at":"2026-08-06T15:28:57.276819Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.09617","last_updated":"2023-05-16T17:11:29Z","snapshot_observed_at":"2026-08-03T08:52:05.725675Z","submitted_at":"2023-05-16T17:11:29Z","title":"Towards Expert-Level Medical Question Answering with Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.09617","snapshot_observed_at":"2026-08-06T15:28:45.191838Z","title":"Towards Expert-Level Medical Question Answering with Large Language Models [Internet]","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.191838Z"},"links":{"cited_paper":"/paper/2305.09617","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:3189beb7940012960ebe1958097d14237e8586e2e8370846923a34aa6ac63e82","observation_id":"b7a59abd-7264-4026-868e-40f46daa170e","resolution":{"observed_at":"2026-08-06T15:28:45.191838Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.05654","last_updated":"2024-01-11T04:25:06Z","snapshot_observed_at":"2026-08-13T04:44:53.257770Z","submitted_at":"2024-01-11T04:25:06Z","title":"Towards Conversational Diagnostic AI","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.05654","snapshot_observed_at":"2026-08-06T15:28:45.332823Z","title":"Towards Conversational Diagnostic AI [Internet]","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.332823Z"},"links":{"cited_paper":"/paper/2401.05654","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:db9d302328a2edb8e3a4b9950d342f695b332d510ebce4b555f19f120932570a","observation_id":"6968ced4-5dd0-41b2-8e79-eae323c24c12","resolution":{"observed_at":"2026-08-06T15:28:45.332823Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:57.068753Z","title":"Opportunities and challenges for ChatGPT and large language models in biomedicine and health","venue":null,"work_id":"42cf9b75-35d1-4f53-97ff-fe6fd49a9cbc","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.436050Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:a13831990aa7ace124c418bf71da90d851e0d355021f53c966142b090babddf8","observation_id":"2974194f-aad3-4e9a-aa91-77ced40dba2e","resolution":{"observed_at":"2026-08-06T15:28:57.159541Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:56.971021Z","title":"Benchmarking large language models for biomedical natural language processing applications and recommendations","venue":null,"work_id":"89542e2d-4966-426b-8507-6dee989fd004","year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.529386Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:fe5dcb2260fb42cff2bfdba27bee05d43e9831687baed3c7dd61b521baeb9868","observation_id":"ae2c89fd-8c14-49e8-8fe9-893bb4abe982","resolution":{"observed_at":"2026-08-06T15:28:57.005286Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06419","last_updated":"2023-08-29T17:44:28Z","snapshot_observed_at":"2026-08-13T10:24:47.146510Z","submitted_at":"2023-08-29T17:44:28Z","title":"Radiology-Llama2: Best-in-Class Large Language Model for Radiology","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06419","snapshot_observed_at":"2026-08-06T15:28:45.636946Z","title":"Radiology-Llama2: Best-in-Class Large Language Model for Radiology [Internet]","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.636946Z"},"links":{"cited_paper":"/paper/2309.06419","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:74c93eeabff60c4af2fab068c02dc43228c1dacb5190a6681549bbdec0b8b9da","observation_id":"f16aec01-0270-4d23-b243-633889ed059a","resolution":{"observed_at":"2026-08-06T15:28:45.636946Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:56.852278Z","title":"Towards a general-purpose foundation model for computational pathology","venue":null,"work_id":"e47cec6f-cfd9-4c41-be07-aafefd2aa1d6","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.692928Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:8661a88577827ad187d5c7f85e9e8123234cc7228f9b3e558b73ca654380e92c","observation_id":"b435dda8-be36-4702-80eb-6a71439ee956","resolution":{"observed_at":"2026-08-06T15:28:56.893597Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2410.03740","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:51.008718Z","title":"Language Enhanced Model for Eye (LEME): An Open-Source Ophthalmology-Specific Large Language Model [Internet]","venue":null,"work_id":"75ae564f-77ce-46bc-bc49-cb1442f78c72","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.752216Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:62206c2ef0979eb574812ae2cd075054ffdfdb37092119a13f5bda4afac2d101","observation_id":"c4e2200e-a08d-424e-b627-ee91fbd05561","resolution":{"observed_at":"2026-08-06T15:28:51.092097Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:56.714490Z","title":"Utilizing Large Language Models to Simplify Radiology Reports: a comparative analysis of ChatGPT3.5, ChatGPT4.0, Google Bard, and Microsoft Bing","venue":null,"work_id":"8b302a45-0915-43a1-8fe8-7c3c85052c55","year":2023},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.820969Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:a50791937cbfda733b9862c84343449c467f94e139ce562f37cdc89674b0cdb8","observation_id":"d0b17bca-a972-4431-becc-bd81ef17c053","resolution":{"observed_at":"2026-08-06T15:28:56.791303Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:56.570552Z","title":"Can OpenAI’s New o1 Model Outperform Its Predecessors in Common Eye Care Queries? Ophthalmology Science [Internet] 2025 [cited 2025 Apr 26];5(4)","venue":null,"work_id":"d7f2810c-d2eb-417d-8a6a-f98fe282ecc7","year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.911562Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:bb47e7c3efb9d4c1ec2574454e6e8c4e58278afb8c3d362970c1f40379763529","observation_id":"dd25c89c-50a6-45f5-9800-8396f2691f32","resolution":{"observed_at":"2026-08-06T15:28:56.626722Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.00840","last_updated":"2024-02-29T09:35:41Z","snapshot_observed_at":"2026-08-15T04:07:47.600345Z","submitted_at":"2024-02-29T09:35:41Z","title":"EyeGPT: Ophthalmic Assistant with Large Language Models","version":1},"cited_work":{"arxiv_id":"2403.00840","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.00840","snapshot_observed_at":"2026-08-06T15:28:50.667856Z","title":"EyeGPT: Ophthalmic Assistant with Large Language Models","venue":"cs.CL","work_id":"4ad2dc10-d30d-436c-846e-74401474679d","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:45.969695Z"},"links":{"cited_paper":"/paper/2403.00840","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:dd4f72079116c8c143e59197feefbe2fe436ae4bc7e6e05f8bd1ca96d8b2db09","observation_id":"4fbba0a1-01a9-43d9-acd1-4396a2de224a","resolution":{"observed_at":"2026-08-06T15:28:50.835793Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:56.407909Z","title":"Can OpenAI o1’s Enhanced Reasoning Capabilities Extend to Ophthalmology? A Benchmark Study Across Large Language Models and Text Generation Metrics [Internet]","venue":null,"work_id":"fdc485b0-5ba1-48d2-ac85-bcea944834e0","year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.099803Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:0eb86766d9698fe08a206bd95e048bf517d124829a8e3ecc60c21a3cf9848ddb","observation_id":"dc9b8830-4855-42d4-8a87-377e0f522fe7","resolution":{"observed_at":"2026-08-06T15:28:56.497357Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:56.235272Z","title":"ChatGPT-Generated Differential Diagnosis Lists for Complex Case–Derived Clinical Vignettes: Diagnostic Accuracy Evaluation","venue":null,"work_id":"af865f37-1319-4716-86e6-e13f3800e885","year":2023},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.227312Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:f89fb1e063692d4325eec37da614d7c26cab168c45b1cf1d19eade29f080e839","observation_id":"eb89e847-6645-4f25-a93c-09548d54c317","resolution":{"observed_at":"2026-08-06T15:28:56.339531Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:56.114423Z","title":"Performance of ChatGPT in Ophthalmic Registration and Clinical Diagnosis: Cross-Sectional Study","venue":null,"work_id":"3813a00e-0156-478a-8609-e021feebe758","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.311462Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:73ec9d028b134901a15e1fbeca6f90e4c70c4d289baf2803dfedcec10735c26b","observation_id":"77f3be25-38e3-4b8d-af47-cf236a59ef85","resolution":{"observed_at":"2026-08-06T15:28:56.164737Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:55.999094Z","title":"Can off-the-shelf visual large language models detect and diagnose ocular diseases from retinal photographs? BMJ Open Ophth [Internet] 2025 [cited 2025 May 15];10(1)","venue":null,"work_id":"fea55f19-fe5e-4a84-9c5c-08184a17ce11","year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.394293Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:43be6b704a6f8046ca33d7d58853f819c52a07060ad5e546c584c210bf581897","observation_id":"0612924a-dd29-47d7-a6e0-0b232010179e","resolution":{"observed_at":"2026-08-06T15:28:56.057025Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:55.762986Z","title":"ChatGPT: the future of discharge summaries? The Lancet Digital Health 2023;5(3):e107–8","venue":null,"work_id":"2a91cc04-870e-412d-8472-6b3ee78bb93d","year":2023},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.434939Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:4aee304c66063b4d415c831a463c357783c930474911acb77ea622f72954781a","observation_id":"5a512d6d-13b5-485e-b7eb-0988035bb7fb","resolution":{"observed_at":"2026-08-06T15:28:55.890375Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:55.551739Z","title":"Large Language Models Seem Miraculous, but Science Abhors Miracles","venue":null,"work_id":"77824fa1-6476-43d5-88f6-b561efcb737f","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.505475Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:4e82987d39cb39d85e41b0a55e91553e16f3b745fc177b478ae443d0f6549c7c","observation_id":"9523c4c7-4528-4ebe-ba60-3136efa9f739","resolution":{"observed_at":"2026-08-06T15:28:55.584987Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:55.416827Z","title":"Medical Ethics of Large Language Models in Medicine","venue":null,"work_id":"e69cd8a5-3eec-4983-900d-36329b945586","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.581662Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:2f09401402ecfc897ffeff4c319182b69673bd5a39f310464a062fa74906af17","observation_id":"ae5f0ee8-f75f-4af2-9fa2-603115a10bf7","resolution":{"observed_at":"2026-08-06T15:28:55.484975Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:55.288722Z","title":"Google for Developers","venue":null,"work_id":"0c631fdd-1328-46ea-9476-97fb4f381270","year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.668047Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:b68227e64a63d0b23ea6d4419af386208746d3fc1970e51cefd1d28701ba61a9","observation_id":"64f78800-3061-4837-8c3b-de7ea129d70a","resolution":{"observed_at":"2026-08-06T15:28:55.354156Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:55.143289Z","title":"Testing and Evaluation of Health Care Applications of Large Language Models: A Systematic Review","venue":null,"work_id":"3eba9581-a4fb-4599-a9de-c54983be4657","year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.736571Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:0ab5f7da59063e8f665b19040a439aea87c2f8a8cf1035eaf3f272d432e6b9ca","observation_id":"82a1fe2c-2b66-4048-a245-fd54a38c2915","resolution":{"observed_at":"2026-08-06T15:28:55.225606Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2009.03300","last_updated":"2021-01-12T18:57:11Z","snapshot_observed_at":"2026-08-13T20:44:28.824685Z","submitted_at":"2020-09-07T17:59:25Z","title":"Measuring Massive Multitask Language Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.03300","snapshot_observed_at":"2026-08-06T15:28:46.831673Z","title":"Measuring Massive Multitask Language Understanding [Internet]","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.831673Z"},"links":{"cited_paper":"/paper/2009.03300","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:2373233c37aaf07c5e8419505b6dffc92186709d7e93a26436e67e1d42b05a59","observation_id":"ab00acb6-03f0-49cb-918d-09ea297d8ad5","resolution":{"observed_at":"2026-08-06T15:28:46.831673Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:54.982803Z","title":"HealthBench: Evaluating Large Language Models Towards Improved Human Health","venue":null,"work_id":"764d1291-c87c-471a-bea5-0d4b54d081d3","year":null},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.891283Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:ae5c93c5bd974ba53625b52501d7851ac5a7fa4b2abd85f244177d7786d15870","observation_id":"d8afb170-bfd4-4613-84da-88c6a6f8b200","resolution":{"observed_at":"2026-08-06T15:28:55.030073Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:54.848163Z","title":"PathVQA: 30000+ Questions for Medical Visual Question Answering","venue":null,"work_id":"e8bd5638-6c24-4eca-ab57-2a61f8d042b2","year":2020},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:46.949387Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:66493808172ece153f801edcce00bab8d01fe7ec78bb512252846185efdc2e40","observation_id":"ceb9921a-cccb-40af-90b1-b0da72ed42a0","resolution":{"observed_at":"2026-08-06T15:28:54.911878Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:54.716795Z","title":"MIMIC-CXR, a de-identified publicly available database of chest radiographs with free-text reports","venue":null,"work_id":"30589656-cfdc-4762-8f91-fbebe474ee62","year":2019},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.030758Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:d763bda6dba0f4c66d90fe8e266f7001f57bc0fcb42e6f42fd36afd752d2753c","observation_id":"6c297985-bf6d-4bb5-9c5a-ad39439d5b36","resolution":{"observed_at":"2026-08-06T15:28:54.766328Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2009.13081","last_updated":"2020-09-28T05:07:51Z","snapshot_observed_at":"2026-08-15T11:16:02.015660Z","submitted_at":"2020-09-28T05:07:51Z","title":"What Disease does this Patient Have? A Large-scale Open Domain Question Answering Dataset from Medical Exams","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.13081","snapshot_observed_at":"2026-08-06T15:28:47.108780Z","title":"What Disease does this Patient Have? A Large-scale Open Domain Question Answering Dataset from Medical Exams [Internet]","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.108780Z"},"links":{"cited_paper":"/paper/2009.13081","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:79771629d462afef3698785fee6d4ddbceadbcdef833f29932fcf44faa5c96e5","observation_id":"d8c35cad-0cd8-4378-9ce3-dd0d82383e72","resolution":{"observed_at":"2026-08-06T15:28:47.108780Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.13650","last_updated":"2025-04-18T12:09:15Z","snapshot_observed_at":"2026-08-14T07:59:52.157444Z","submitted_at":"2025-04-18T12:09:15Z","title":"EyecareGPT: Boosting Comprehensive Ophthalmology Understanding with Tailored Dataset, Benchmark and Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.13650","snapshot_observed_at":"2026-08-06T15:28:47.162607Z","title":"EyecareGPT: Boosting Comprehensive Ophthalmology Understanding with Tailored Dataset, Benchmark and Model [Internet]","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.162607Z"},"links":{"cited_paper":"/paper/2504.13650","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:181c4bdb4682599c05a2d1bb243763658048ebb42e2b79a2ebcc9a9e7e6264b3","observation_id":"620c27bb-70d6-4918-8fb5-5eef302943b4","resolution":{"observed_at":"2026-08-06T15:28:47.162607Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14304","last_updated":"2024-12-18T20:18:03Z","snapshot_observed_at":"2026-08-14T14:39:37.909655Z","submitted_at":"2024-12-18T20:18:03Z","title":"Multi-OphthaLingua: A Multilingual Benchmark for Assessing and Debiasing LLM Ophthalmological QA in LMICs","version":1},"cited_work":{"arxiv_id":"2412.14304","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.14304","snapshot_observed_at":"2026-08-06T15:28:50.433641Z","title":"Multi-OphthaLingua: A Multilingual Benchmark for Assessing and Debiasing LLM Ophthalmological QA in LMICs","venue":"cs.CL","work_id":"84474afb-4c08-4b6d-8f3f-a01455efdd23","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.251903Z"},"links":{"cited_paper":"/paper/2412.14304","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:89fad121235fd00267144f8a62ef8add79e52b7529ddecc5f9ccc4c91314107b","observation_id":"793f4a7b-ebd2-484d-b992-0539ba088dbf","resolution":{"observed_at":"2026-08-06T15:28:50.529350Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:54.598772Z","title":"Unveiling the clinical incapabilities: a benchmarking study of GPT-4V(ision) for ophthalmic multimodal image analysis","venue":null,"work_id":"7f9b23f7-6e76-4681-aeb6-2bccceed5f48","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.339856Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:7589e78768d3c11c4e277c1000f2080642c07ecad4737e2d10303c01bbb9a8b3","observation_id":"a8106e3a-06e1-46d6-bf98-58024df53215","resolution":{"observed_at":"2026-08-06T15:28:54.660686Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2023.10032","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:50.254977Z","title":"Evaluating the Performance of ChatGPT in Ophthalmology: An Analysis of Its Successes and Shortcomings","venue":null,"work_id":"bf7f7f92-09eb-4a0f-b5af-550926b3b759","year":2023},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.412379Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:2f689d51ab260cad6da33442cbd88569d7ce2f91ddba65632744a6c7367b85eb","observation_id":"7dab4c86-1fba-44e8-97ed-0aa0aeaf3785","resolution":{"observed_at":"2026-08-06T15:28:50.372241Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.04933","last_updated":"2023-11-07T16:19:45Z","snapshot_observed_at":"2026-08-13T05:30:29.319591Z","submitted_at":"2023-11-07T16:19:45Z","title":"Evaluating Large Language Models in Ophthalmology","version":1},"cited_work":{"arxiv_id":"2311.04933","doi":null,"metadata_source":"pith","pith_arxiv_id":"2311.04933","snapshot_observed_at":"2026-08-06T15:28:49.970298Z","title":"Evaluating Large Language Models in Ophthalmology","venue":"cs.CL","work_id":"9f6da05f-feba-469d-961e-0deda93b2157","year":2023},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.476669Z"},"links":{"cited_paper":"/paper/2311.04933","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:aae2ff89921801b96777a804d3dd14eea4dffc4990fd9113b944eb26dd3e4301","observation_id":"7cc5c6d7-369d-43b7-81c6-7381732dce9c","resolution":{"observed_at":"2026-08-06T15:28:50.050972Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01243","last_updated":"2025-02-03T11:04:51Z","snapshot_observed_at":"2026-08-13T15:01:44.141448Z","submitted_at":"2025-02-03T11:04:51Z","title":"OphthBench: A Comprehensive Benchmark for Evaluating Large Language Models in Chinese Ophthalmology","version":1},"cited_work":{"arxiv_id":"2502.01243","doi":null,"metadata_source":"pith","pith_arxiv_id":"2502.01243","snapshot_observed_at":"2026-08-06T15:28:49.752882Z","title":"OphthBench: A Comprehensive Benchmark for Evaluating Large Language Models in Chinese Ophthalmology","venue":"cs.CL","work_id":"0dcbc693-5039-44bd-abcf-d5d394440e21","year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.559357Z"},"links":{"cited_paper":"/paper/2502.01243","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:46427cb5b2d3abfd7dcdf90eec2d89483fcf9af02b31ae7a88c2e4ef72be71c6","observation_id":"71bd979c-52ef-432c-96cc-8e3fa3c7fbc4","resolution":{"observed_at":"2026-08-06T15:28:49.862916Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:54.476737Z","title":"EYE-Llama, an in-domain large language model for ophthalmology","venue":null,"work_id":"d414e7f1-fe9f-41a7-a29e-bdb836a30d87","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.652656Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:3d121a92a2e0dd64213dcf5377d56334a9a1465a575f52d2b611eb77ff54b4f2","observation_id":"7a40c24a-e822-4dee-85cc-54da13ee7c41","resolution":{"observed_at":"2026-08-06T15:28:54.534407Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:54.356492Z","title":"BioASQ-QA: A manually curated corpus for Biomedical Question Answering","venue":null,"work_id":"7b4acff5-89e9-49bf-a1fb-651d6090a43b","year":2023},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.743157Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:b6472c887e31c4d1b1c347dfefd1d75040fe8a12e81330136fb40a1d45f3be27","observation_id":"ac32d7b4-1998-4ac3-9e4d-29eb344278e7","resolution":{"observed_at":"2026-08-06T15:28:54.420785Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:54.178940Z","title":"Medmcqa: A large-scale multi-subject multi-choice dataset for medical domain question answering","venue":null,"work_id":"1621442d-e4f0-4fae-bca4-8911d8cf9adc","year":2022},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.803904Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:50b55c58a77ab05756234d155c694db0bde4c149ed803b0090c9a711b5076c8b","observation_id":"a14f5cc9-a325-4875-a478-471344818985","resolution":{"observed_at":"2026-08-06T15:28:54.253238Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:54.038507Z","title":"PubMedQA: A Dataset for Biomedical Research Question Answering","venue":null,"work_id":"d19db992-e0ac-4624-bb32-dd8490d4156e","year":2019},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.857163Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:22a79376c81c38662b0bec3554c696eca3453113b4c1abfe15280ea763305366","observation_id":"2feb1c9a-705f-40ea-94fd-ce61c8dc9cdb","resolution":{"observed_at":"2026-08-06T15:28:54.129247Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:53.752414Z","title":"Domain-Specific Language Model Pretraining for Biomedical Natural Language Processing","venue":null,"work_id":"888aa788-8f18-487f-8d5e-355358c94a08","year":2022},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:47.913858Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:d0320e1b127f99679068dd0e3cc71ce1f51e862329c4eab58c8242b6b9763582","observation_id":"d80f5379-265e-4c45-8b36-9a23073f0a87","resolution":{"observed_at":"2026-08-06T15:28:53.878466Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:53.501044Z","title":"How Does ChatGPT Perform on the United States Medical Licensing Examination? The Implications of Large Language Models for Medical Education and Knowledge Assessment","venue":null,"work_id":"ab38069a-e19a-4020-bbc9-8b07d9200564","year":2023},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.008049Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:a56882cb49dfb7777f8dcb005c66ff768c37031925b3fefc87054f993ad2f7e1","observation_id":"89a81b06-6227-4b91-8831-4b97f4ad7987","resolution":{"observed_at":"2026-08-06T15:28:53.578158Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.11186","last_updated":"2025-04-15T13:42:34Z","snapshot_observed_at":"2026-08-07T16:05:31.091720Z","submitted_at":"2025-04-15T13:42:34Z","title":"Benchmarking Next-Generation Reasoning-Focused Large Language Models in Ophthalmology: A Head-to-Head Evaluation on 5,888 Items","version":1},"cited_work":{"arxiv_id":"2504.11186","doi":null,"metadata_source":"pith","pith_arxiv_id":"2504.11186","snapshot_observed_at":"2026-08-06T15:28:49.556000Z","title":"Benchmarking Next-Generation Reasoning-Focused Large Language Models in Ophthalmology: A Head-to-Head Evaluation on 5,888 Items","venue":"cs.CL","work_id":"4fad266a-3a63-464f-9ae5-e0d34b2f552a","year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.085423Z"},"links":{"cited_paper":"/paper/2504.11186","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:c1ee10134f64f4cd78c09e37ea4363a5620382bd4e5e8d1d46a3ad93fa0b73c5","observation_id":"d4321c0a-7ddc-4671-b69d-38d1c30a9009","resolution":{"observed_at":"2026-08-06T15:28:49.676971Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:53.230851Z","title":"ROUGE: A Package for Automatic Evaluation of Summaries [Internet]","venue":null,"work_id":"db642ae3-8188-4073-902a-8e11bc241f2e","year":2004},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.171460Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:4f25fed3aa494de888ca953f569e5e364ebddf4e959e229a0b02edffd5016084","observation_id":"56f9c3ce-b564-4ea3-9fa2-20ff9cf01447","resolution":{"observed_at":"2026-08-06T15:28:53.374598Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1904.09675","last_updated":"2020-02-24T18:59:28Z","snapshot_observed_at":"2026-07-29T15:42:51.774083Z","submitted_at":"2019-04-21T23:08:53Z","title":"BERTScore: Evaluating Text Generation with BERT","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1904.09675","snapshot_observed_at":"2026-08-06T15:28:48.277853Z","title":"BERTScore: Evaluating Text Generation with BERT [Internet]","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.277853Z"},"links":{"cited_paper":"/paper/1904.09675","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:f3ffc6c1c7ffeedba96326e302fe148e9f110d79d3aa4bd93ca149f4c0b781de","observation_id":"2f114fce-6175-467b-9f64-0ea339905859","resolution":{"observed_at":"2026-08-06T15:28:48.277853Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.11520","last_updated":"2021-10-27T11:16:01Z","snapshot_observed_at":"2026-08-13T18:59:25.766093Z","submitted_at":"2021-06-22T03:20:53Z","title":"BARTScore: Evaluating Generated Text as Text Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.11520","snapshot_observed_at":"2026-08-06T15:28:48.347345Z","title":"BARTScore: Evaluating Generated Text as Text Generation [Internet]","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.347345Z"},"links":{"cited_paper":"/paper/2106.11520","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:1b3387c9f455e4b4eb479227cc663c56e96267c160b34402fc81efe98a01adf8","observation_id":"803eb354-c9b5-4c9a-aecf-cf64380105c3","resolution":{"observed_at":"2026-08-06T15:28:48.347345Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.16739","last_updated":"2023-05-26T08:41:59Z","snapshot_observed_at":"2026-08-14T23:19:09.773175Z","submitted_at":"2023-05-26T08:41:59Z","title":"AlignScore: Evaluating Factual Consistency with a Unified Alignment Function","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.16739","snapshot_observed_at":"2026-08-06T15:28:48.415628Z","title":"AlignScore: Evaluating Factual Consistency with a Unified Alignment Function [Internet]","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.415628Z"},"links":{"cited_paper":"/paper/2305.16739","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:0d62e4a8fa5f535f75cbad164717bca867a7b38ebd657593ad390da1f8782775","observation_id":"005638be-e921-4c09-bc34-fda54f712046","resolution":{"observed_at":"2026-08-06T15:28:48.415628Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:52.939485Z","title":"A visual-language foundation model for computational pathology","venue":null,"work_id":"d58e75b5-0e04-4395-9870-935a4b00c8d5","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.535484Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:876c8daafcf69bfbbbbdca3b69753e2189f491242a424119a7ac52a095b6abae","observation_id":"901cef58-08f4-4bde-86c0-d97e2fd9bd10","resolution":{"observed_at":"2026-08-06T15:28:53.077371Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:52.673909Z","title":"Developing and Evaluating Large Language Model– Generated Emergency Medicine Handoff Notes","venue":null,"work_id":"759c891c-5250-4fdc-ba21-48303d389fb3","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.666664Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:3da36e1f14357a753bd0d97e00fbbc2e5d4b8ac7fe0585c41249a9149f0ecec9","observation_id":"a3d77380-73a4-477c-a265-12412b9e81ff","resolution":{"observed_at":"2026-08-06T15:28:52.795191Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:52.462778Z","title":"CPMI-ChatGLM: parameter-efficient fine-tuning ChatGLM with Chinese patent medicine instructions","venue":null,"work_id":"a3c1a294-4124-419d-8fcf-3f0eb658f81e","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.703934Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:00afc199ea79b7cae3bbae446e27ad1195c23c9a649f984176f263018275042d","observation_id":"a8c674eb-fc2d-4269-8c39-fa32849d718d","resolution":{"observed_at":"2026-08-06T15:28:52.523882Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:52.233506Z","title":"Reducing hallucinations of large language models via hierarchical semantic piece","venue":null,"work_id":"6e16733d-643f-4890-b2a2-23a2b58442e8","year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.790885Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:a142664ea90b4a16585df3a89d7e8d6c73043bf4da2555e51a1a975e08b26ffd","observation_id":"c9a1c9a2-34c0-41da-a883-64f5401f8426","resolution":{"observed_at":"2026-08-06T15:28:52.356278Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:51.930292Z","title":"Meteor: an automatic metric for MT evaluation with high levels of correlation with human judgments","venue":null,"work_id":"116c7bb5-f9da-4540-aec6-9558412998d5","year":2007},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.835919Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:d0b852e3c5948b5628b4aeb1ce2ed63528b15106928f482f5a4f2ae163a98675","observation_id":"6c1cbc6e-4805-4862-b218-907ab14c75b3","resolution":{"observed_at":"2026-08-06T15:28:52.100594Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.01620","last_updated":"2025-02-05T18:36:42Z","snapshot_observed_at":"2026-08-12T22:32:28.694469Z","submitted_at":"2024-10-02T14:57:58Z","title":"LMOD: A Large Multimodal Ophthalmology Dataset and Benchmark for Large Vision-Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.01620","snapshot_observed_at":"2026-08-06T15:28:48.930283Z","title":"LMOD: A Large Multimodal Ophthalmology Dataset and Benchmark for Large Vision-Language Models [Internet]","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.930283Z"},"links":{"cited_paper":"/paper/2410.01620","citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:229738c121eea2e77f6e4b32584ec6323e944bb83cbcf159a3a9a56b31fa816c","observation_id":"8ddb41ec-a6c3-471f-b196-f0edbf68b762","resolution":{"observed_at":"2026-08-06T15:28:48.930283Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:51.753376Z","title":"medical QA dataset","venue":null,"work_id":"aba8216e-f3b5-4ec2-b2c2-b2b7389a019e","year":2024},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:48.982187Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:3ffd760e7b13f5e926b6583e765f7ce4c4156c95190ad42132429d3ae15448a9","observation_id":"b030e2a6-9b2d-4496-a218-a08b76dbfc17","resolution":{"observed_at":"2026-08-06T15:28:51.799503Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:51.559534Z","title":"Each question offers four options, out of which only one is correct","venue":null,"work_id":"97f07e1a-6f47-42a8-82a7-59802ee2060e","year":2013},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:49.197532Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:764e4ecd4fb897c129f33eb4bf614868997d4f93e45a9669ec939a8ba7574a19","observation_id":"a5402410-20ca-46c3-9d2d-a7920e64ef30","resolution":{"observed_at":"2026-08-06T15:28:51.668008Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:51.421053Z","title":"glaucoma,","venue":null,"work_id":"360fc59a-9fc6-491d-8ece-e07cd55c265c","year":null},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:49.272071Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:ebac6ab57ae457036ff213871c6be98b5bb7e5815a3b63cbaa9c5e9d5655b5bd","observation_id":"af61e65b-595e-4430-b810-8a7cd2432b2e","resolution":{"observed_at":"2026-08-06T15:28:51.495773Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:28:51.234727Z","title":null,"venue":null,"work_id":"f0331a8a-b94c-4b09-a9e0-77ff15bbe21a","year":null},"citing_paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T15:28:49.334512Z"},"links":{"citing_paper":"/paper/2507.15717"},"observation_digest":"sha256:3e7eb9e4f9a75de78cd899570407a104453ca66849c17eea3acfdf166352efb2","observation_id":"c2becfc3-52ba-434d-9744-9d5071f96789","resolution":{"observed_at":"2026-08-06T15:28:51.338945Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2507.15717","last_updated":"2025-07-21T15:27:32Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-12T16:46:35.278321Z","submitted_at":"2025-07-21T15:27:32Z","title":"BEnchmarking LLMs for Ophthalmology (BELO) for Ophthalmological Knowledge and Reasoning"},"reference_resolution":{"displayed":52,"state_counts":{"malformed_identifier":1,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":10,"verified_exact":6,"verified_fuzzy":34},"total_outbound_references":52},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 52 of 52 outbound references and 0 inbound Pith citation observations for arXiv:2507.15717."}