{"as_of":"2026-08-09T14:16:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:9b8f6ddd590aceec4053f88f6f378b8545054bc60ed82520cc8f28c22a3d1d1f","coverage":[{"denominator":53,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":53,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-31T11:25:33.319256Z","state":"measured"},{"denominator":53,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":53,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.28318/citation-record","integrity":"/paper/2607.28318/integrity","json":"/paper/2607.28318/citation-record.json","paper":"/paper/2607.28318"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.196700Z","title":"Image analysis and machine learning in digital pathology: Challenges and opportunities.Medical image analysis, 33:170–175, 2016","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.196700Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:1f002a05ed5c1232d77e5786b1e7220887086235487ad4b90de23a19f581caef","observation_id":"ee392e00-e0f9-4ed4-b7e2-6dd46911f598","resolution":{"observed_at":"2026-07-31T11:25:33.196700Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.200139Z","title":"Digital pathology and artificial intelligence.The lancet oncology, 20(5):e253–e261, 2019","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.200139Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:f14add174e96bb4face6df01434cc5667751cebdd336f223a2bf439d19881f3f","observation_id":"abd20711-c233-4ff2-8d6c-14cb6d194b49","resolution":{"observed_at":"2026-07-31T11:25:33.200139Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.202870Z","title":"Clinical-grade computational pathology using weakly supervised deep learning on whole slide images.Nature medicine, 25(8):1301–1309, 2019","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.202870Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:44bad09a968dcac6cec98a722833a7b6a35a271e4c4e69a8f92adf810d872064","observation_id":"b797886b-4ad4-482a-a990-b06c72eaf96f","resolution":{"observed_at":"2026-07-31T11:25:33.202870Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.205524Z","title":"Data-efficient and weakly supervised computational pathology on whole-slide images.Nature biomedical engineering, 5(6):555–570, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.205524Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:3707651fd29a30b1383e9c98054e5b893f6359d59f6265c984e25c49fca7fe1e","observation_id":"def72c94-67bf-4879-a0c4-ed1839cb6dc8","resolution":{"observed_at":"2026-07-31T11:25:33.205524Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.208012Z","title":"Feature re-embedding: Towards foundation model-level performance in computational pathology","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.208012Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:2eb721eb65ffc26b47df25b82326d45384a31283f25367344a4d70c47128329d","observation_id":"0d992d6a-5d6b-41a8-a1fb-4764e6569e38","resolution":{"observed_at":"2026-07-31T11:25:33.208012Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.210682Z","title":"M4: Multi-proxy multi-gate mixture of experts network for multiple instance learning in histopathology image analysis.Medical Image Analysis, 103:103561, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.210682Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:3dcd4fcdf8920cb739a0b3d682a5766f667ead129f158288f51d228170fff58c","observation_id":"cb01279d-c002-44ff-946e-e3e25a863ffb","resolution":{"observed_at":"2026-07-31T11:25:33.210682Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.213348Z","title":"Smmile enables accurate spatial quantification in digital pathology using multiple-instance learning.Nature Cancer, pages 1–17, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.213348Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:92f6cbcebe335409d9548e2ee116eae964aafafb08e3de173ebadb9ffb130514","observation_id":"5f8db21e-cd39-46a9-9a43-ef766108e4c6","resolution":{"observed_at":"2026-07-31T11:25:33.213348Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.215504Z","title":"A pathology foundation model for cancer diagnosis and prognosis prediction.Nature, 634(8035):970–978, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.215504Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:24025f91f829003e869b7624b71923143dce7273b67a7cf5702123b08b728e55","observation_id":"591e73b3-3b98-4fd0-884f-f1b9f32a0138","resolution":{"observed_at":"2026-07-31T11:25:33.215504Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.217759Z","title":"A vision–language foundation model for precision oncology.Nature, 638(8051):769–778, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.217759Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:2acf6b5b13f7a1741830e9f09ad7a71860d4729ebac21026861b7addadee4fb9","observation_id":"27ead212-24ce-4c62-8867-84dd61a08ebb","resolution":{"observed_at":"2026-07-31T11:25:33.217759Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.219886Z","title":"A generalizable pathology foundation model using a unified knowledge distillation pretraining framework.Nature Biomedical Engineering, pages 1–20, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.219886Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:f6b3f64ace4fe77a4d99a807c5e3e0ead8544388afbb4708936e205e0d9e276f","observation_id":"1062cbf1-698c-4f4d-8eff-033129701bf4","resolution":{"observed_at":"2026-07-31T11:25:33.219886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.222249Z","title":"Pathasst: A generative foundation ai assistant towards artificial general intelligence of pathology","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.222249Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:c744d3d52c588c7c12e62fa4567ec2b7e933d90f98adf565f3e2d29635f0cd0e","observation_id":"dcd0dd6f-3c97-401b-b9e9-a71408a91b45","resolution":{"observed_at":"2026-07-31T11:25:33.222249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.224458Z","title":"Quilt-llava: Visual instruction tuning by extracting localized narratives from open-source histopathology videos","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.224458Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:10bf2953c76bc0d7402a7ab840c7a7548f2902ed6b0b13e940060731573e601a","observation_id":"ad28c6f7-0b46-48d5-b686-1dc94ed708dd","resolution":{"observed_at":"2026-07-31T11:25:33.224458Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.226707Z","title":"Cpath-omni: A unified multimodal foundation model for patch and whole slide image analysis in computational pathology","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.226707Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:dbe3462b29905f2513743756acae3eb1019dc4fb5db4246612320b52450de527","observation_id":"e4d53fa6-b4dd-41a0-8834-f25d4ca7904f","resolution":{"observed_at":"2026-07-31T11:25:33.226707Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.228859Z","title":"Patho-agenticrag: towards multimodal agentic retrieval-augmented generation for pathology vlms via reinforcement learning","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.228859Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:238b26bad57e4c71fd6d3414d5c3fc6b012d62258534ccf72d2451bfeda5d9e0","observation_id":"f44048d6-6ea2-43d5-a2a6-8186c519d14b","resolution":{"observed_at":"2026-07-31T11:25:33.228859Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.231007Z","title":"Wsicaption: Multiple instance generation of pathology reports for gigapixel whole-slide images","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.231007Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:1c2991d1726b3a1017e48363050d495e21088a0a127259b54f94c8bb675a4b68","observation_id":"426ecd3a-7511-472a-acba-a14f5a7e7f04","resolution":{"observed_at":"2026-07-31T11:25:33.231007Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.233225Z","title":"Histgen: Histopathol- ogy report generation via local-global feature encoding and cross-modal context interaction","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.233225Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:6cb18fe009e455755620e1f6f22310fe681e400c42b5f20d4e241f11b6027ef9","observation_id":"6b787e75-5b01-47fa-a5e6-8784c4d30676","resolution":{"observed_at":"2026-07-31T11:25:33.233225Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.235461Z","title":"Generating der- matopathology reports from gigapixel whole slide images with histogpt.Nature communications, 16(1): 4886, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.235461Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:2761ebdd0eb5595f22ca47b77fd8bbb4a7410b6ba8dc292424e3ca918369f049","observation_id":"c38c5b75-45af-4b48-884c-6e541ac7750f","resolution":{"observed_at":"2026-07-31T11:25:33.235461Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.237554Z","title":"Qcagent: An agentic framework for quality-controllable pathology report generation from whole slide image.arXiv preprint arXiv:2603.01647, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.237554Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:64ba939f8b31f625669956c3fa0f0b4f4e1ed84a7489bbbf61836a73bfd3b440","observation_id":"33e7670d-7bbf-42df-a033-9fbb4cb8bd8d","resolution":{"observed_at":"2026-07-31T11:25:33.237554Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.239802Z","title":"Slidechat: A large vision-language assistant for whole-slide pathology image understanding","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.239802Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:f2f1e5058df1f891e71ae24dd16c75a005d87028d3825740167980011649ed41","observation_id":"3cc93f35-d76c-454b-ba72-3ebe5d1a9836","resolution":{"observed_at":"2026-07-31T11:25:33.239802Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.242012Z","title":"Wsi-llava: A multimodal large language model for whole slide image","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.242012Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:c78a0092ab165e45d450e440087ab05db83da5daa7513bc898887fbbc419f47b","observation_id":"cc9189b8-8317-48ba-b818-56bdcbde2675","resolution":{"observed_at":"2026-07-31T11:25:33.242012Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.244085Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.244085Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:9ff5bca84632b024d710485f9ed3aebea35e47ecad753e77d02f55fa8cbe7db7","observation_id":"1f9b6da5-093f-4381-88c9-d9a99b69de5e","resolution":{"observed_at":"2026-07-31T11:25:33.244085Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.19652","last_updated":"2026-06-11T03:10:46Z","snapshot_observed_at":"2026-08-07T01:30:05.915971Z","submitted_at":"2025-11-24T19:33:56Z","title":"Navigating Gigapixel Pathology Images with Large Multimodal Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.19652","snapshot_observed_at":"2026-07-31T11:25:33.246243Z","title":"Navigating gigapixel pathology images with large multimodal models.arXiv preprint arXiv:2511.19652, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.246243Z"},"links":{"cited_paper":"/paper/2511.19652","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:04cfed0eb3a33670979fbf56cc1c3e91c47dc43fc88c030e2de537140b3d6122","observation_id":"62d7bd0b-69ec-4d96-b79d-5424dd192775","resolution":{"observed_at":"2026-07-31T11:25:33.246243Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.248859Z","title":"Pathology-cot: Learning visual chain-of-thought agent from expert whole slide image diagnosis behavior","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.248859Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:17ea79a85b31c38dc797d75701132b383d371d5043e4f2cb9915c0cf8b6dfdfb","observation_id":"45d3f587-1f71-456a-880a-8ca35e0b9e3b","resolution":{"observed_at":"2026-07-31T11:25:33.248859Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.251062Z","title":"Pathagent: Toward interpretable analysis of whole-slide pathology images via large language model-based agentic reasoning.arXiv preprint arXiv:2511.17052, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.251062Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:a9b501d524e211a6091822a9eb8157e19d03cb11df93ffe6aef16d89e83cb4b2","observation_id":"6c21b42d-db0c-474a-8b7c-dcadd5b1098a","resolution":{"observed_at":"2026-07-31T11:25:33.251062Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.253268Z","title":"Pathfound: An agentic multimodal model activating evidence- seeking pathological diagnosis.arXiv preprint arXiv:2512.23545, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.253268Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:ef07f61f9e1c27bd5435542cb37811568cf8f124cf0ae0f9a99c73293e6eba25","observation_id":"31acb2cd-3bb5-4ab8-8105-a408b70f2a91","resolution":{"observed_at":"2026-07-31T11:25:33.253268Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.255404Z","title":"Patho-r1: A multimodal reinforcement learning-based pathology expert reasoner","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.255404Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:68fc711f52454ddd9a6760347a48f72feb7c8c24408f15a10348e00a81e30f6b","observation_id":"db6563bd-d314-471d-8f18-f08f1617d05c","resolution":{"observed_at":"2026-07-31T11:25:33.255404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.257725Z","title":"Pathreasoner-r1: Instilling structured reasoning into pathology vision-language model via knowledge-guided policy optimization","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.257725Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:7efbf437feab8537ee7db6fd2a31c23304ed63b9464193f6015b87337504bb85","observation_id":"2201a5e4-8a02-480e-a037-51e27e677a7a","resolution":{"observed_at":"2026-07-31T11:25:33.257725Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.14680","last_updated":"2025-07-19T16:11:03Z","snapshot_observed_at":"2026-08-06T15:47:49.716888Z","submitted_at":"2025-07-19T16:11:03Z","title":"WSI-Agents: A Collaborative Multi-Agent System for Multi-Modal Whole Slide Image Analysis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.14680","snapshot_observed_at":"2026-07-31T11:25:33.259941Z","title":"Wsi-agents: A collaborative multi-agent system for multi-modal whole slide image analysis.arXiv preprint arXiv:2507.14680, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.259941Z"},"links":{"cited_paper":"/paper/2507.14680","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:7cae83c47ba299345b2aa94d0d879bfc2fc30e8a2e501ad8c734004d1ba0626b","observation_id":"e823c913-5129-4c3c-a401-9fbaeb3d6419","resolution":{"observed_at":"2026-07-31T11:25:33.259941Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.10286","last_updated":"2020-03-07T17:55:41Z","snapshot_observed_at":"2026-07-06T09:06:42.071993Z","submitted_at":"2020-03-07T17:55:41Z","title":"PathVQA: 30000+ Questions for Medical Visual Question Answering","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.10286","snapshot_observed_at":"2026-07-31T11:25:33.262455Z","title":"Pathvqa: 30000+ questions for medical visual question answering.arXiv preprint arXiv:2003.10286, 2020","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.262455Z"},"links":{"cited_paper":"/paper/2003.10286","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:21b4a3aa090128fff7debccbac60ae83faec340c0fb6e756b8796e3cdd6322e3","observation_id":"ae8864d8-f321-4734-913f-52df9b100b9b","resolution":{"observed_at":"2026-07-31T11:25:33.262455Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.264929Z","title":"Pathmmu: A massive multimodal expert-level benchmark for understanding and reasoning in pathology","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.264929Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:39959a2e958bed88bef1d61b8b92529825ea73c367f5aed35ef309f678f46d27","observation_id":"6feb706e-0aaa-473a-8b41-f8a39b3d2b1d","resolution":{"observed_at":"2026-07-31T11:25:33.264929Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.267148Z","title":"Wsi-vqa: Interpreting whole slide images by generative visual question answering","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.267148Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:4af6df02c5ddeb0c199bbd049519bc6dd4494f154311e4a97c91c5ab6160db48","observation_id":"77d682a6-f533-41b5-9f27-2ea681d07658","resolution":{"observed_at":"2026-07-31T11:25:33.267148Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.269324Z","title":"Micro-bench: A microscopy benchmark for vision-language understanding.Advances in Neural Information Processing Systems, 37:30670–30685, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.269324Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:2cfaba16647ade632b841e5b87a566d7756e179d7d26665fd98cd5c8a6abbdd0","observation_id":"e81cf23a-faea-437b-ab77-37b508914e7d","resolution":{"observed_at":"2026-07-31T11:25:33.269324Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.271409Z","title":"Pathbench: Advancing the benchmark of large multimodal models for pathology image understanding at patch and whole slide level.IEEE Transactions on Medical Imaging, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.271409Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:3850dc64ed03fa0e2ed716361795ccdee61857202f6fc37137f568573680a61f","observation_id":"84d7f53d-e864-4643-874f-9b81da64e2da","resolution":{"observed_at":"2026-07-31T11:25:33.271409Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.273526Z","title":"Pathvg: A new benchmark and dataset for pathology visual grounding","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.273526Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:e150982f6a87ef748f3c5d10295325da1b1da6ec1158376f35c4c94cc4346e6a","observation_id":"ad097bcc-b8a3-4316-b0a2-551e33810608","resolution":{"observed_at":"2026-07-31T11:25:33.273526Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.275912Z","title":"Quilt-1m: One million image-text pairs for histopathology.Advances in neural information processing systems, 36:37995–38017, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.275912Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:bf7295a519b41b263dcf127b93534316853e655cb10c5132142c73c746fa22bb","observation_id":"2d22eb76-75f0-4f08-b154-ac0c0dc1a5fd","resolution":{"observed_at":"2026-07-31T11:25:33.275912Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.278192Z","title":"Pathgen-1.6 m: 1.6 million pathology image-text pairs generation through multi-agent collaboration","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.278192Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:a5ae17ed35344f57aed3d9763edceb140ea015c27eb0b2025386c514fb114fca","observation_id":"ee76702a-bd3b-4255-9545-79aaeaafea3a","resolution":{"observed_at":"2026-07-31T11:25:33.278192Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.280311Z","title":"Mirage the illusion of visual understanding.arXiv preprint arXiv:2603.21687, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.280311Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:3f77b4039240309260927ff63aa1d1aa2544b1699136f70c8ce31a2dd859ded6","observation_id":"c0acdfef-b43a-4892-97be-087a93208172","resolution":{"observed_at":"2026-07-31T11:25:33.280311Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21276","last_updated":"2024-10-25T17:43:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-25T17:43:01Z","title":"GPT-4o System Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21276","snapshot_observed_at":"2026-07-31T11:25:33.282518Z","title":"Gpt-4o system card.arXiv preprint arXiv:2410.21276, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.282518Z"},"links":{"cited_paper":"/paper/2410.21276","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:c17e70fd6a099e5609b645dea300f9327c945ea1c944a484f276099aa96182bd","observation_id":"3ef165bc-27fd-494c-985c-a281cbb48eb2","resolution":{"observed_at":"2026-07-31T11:25:33.282518Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.21631","last_updated":"2025-11-27T12:16:54Z","snapshot_observed_at":"2026-07-06T22:37:03.716474Z","submitted_at":"2025-11-26T17:59:08Z","title":"Qwen3-VL Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.21631","snapshot_observed_at":"2026-07-31T11:25:33.285026Z","title":"Qwen3-vl technical report.arXiv preprint arXiv:2511.21631, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.285026Z"},"links":{"cited_paper":"/paper/2511.21631","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:1734f9d9235eb2b06434637fce7036314d566b3bdc97f561e10a3c3fca314dda","observation_id":"4c76a3d6-ca4d-4139-b23d-11e60f2f86da","resolution":{"observed_at":"2026-07-31T11:25:33.285026Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.06261","last_updated":"2025-12-19T14:25:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-07T17:36:04Z","title":"Gemini 2.5: Pushing the Frontier with Advanced Reasoning, Multimodality, Long Context, and Next Generation Agentic Capabilities","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.06261","snapshot_observed_at":"2026-07-31T11:25:33.287610Z","title":"Gemini 2.5: Pushing the frontier with ad- vanced reasoning, multimodality, long context, and next generation agentic capabilities.arXiv preprint arXiv:2507.06261, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.287610Z"},"links":{"cited_paper":"/paper/2507.06261","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:9f7ac0af12106a736fc301785562ff18a8a265bde8b1c17642d2acfe85cc8ecd","observation_id":"7265a8ce-f0a0-4a68-87b6-95411dc6e419","resolution":{"observed_at":"2026-07-31T11:25:33.287610Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.290327Z","title":"Med-flamingo: a multimodal medical few-shot learner","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.290327Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:4d949daccd6bb119c841a460671661f68f20beedc5c5e0b5e532e7f2cf88050f","observation_id":"557c46fd-dbcb-487e-a7e0-389781926525","resolution":{"observed_at":"2026-07-31T11:25:33.290327Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.292436Z","title":"Llava-med: Training a large language-and-vision assistant for biomedicine in one day.Advances in Neural Information Processing Systems, 36:28541–28564, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.292436Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:4fdf69732c9c15b831516832f3ae933839d2ce16528ade2c89e111df99afc554","observation_id":"a2a51705-7601-4331-a1b4-18a3f3757211","resolution":{"observed_at":"2026-07-31T11:25:33.292436Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.294928Z","title":"Towards generalist biomedical ai.Nejm Ai, 1(3): AIoa2300138, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.294928Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:e19f1c5f34ff25edd6910e5a9e58f94074815322f735f6653d2c3344b0307508","observation_id":"b2f26496-71cb-402c-833d-93131a566e2a","resolution":{"observed_at":"2026-07-31T11:25:33.294928Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.297132Z","title":"Towards injecting medical visual knowledge into multimodal llms at scale","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.297132Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:af1d05c4a25bbe8e7c819b4abe6c20f1581411f95b82db6db2bcdd5decdc9f5f","observation_id":"2cb723e4-f6ed-4059-ba78-08850222c441","resolution":{"observed_at":"2026-07-31T11:25:33.297132Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.07044","last_updated":"2025-06-13T04:22:02Z","snapshot_observed_at":"2026-07-06T21:38:34.998414Z","submitted_at":"2025-06-08T08:47:30Z","title":"Lingshu: A Generalist Foundation Model for Unified Multimodal Medical Understanding and Reasoning","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.07044","snapshot_observed_at":"2026-07-31T11:25:33.299406Z","title":"Lingshu: A generalist foundation model for unified multimodal medical understanding and reasoning.arXiv preprint arXiv:2506.07044, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.299406Z"},"links":{"cited_paper":"/paper/2506.07044","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:400638737ef2d6babb9a2fd3e7048605be1b0de59e5bbfc89368b7634904d4b8","observation_id":"a9080ed6-d833-4a70-b8aa-4ced1550cd2d","resolution":{"observed_at":"2026-07-31T11:25:33.299406Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.302220Z","title":"Deep learning in histopathology: the path to the clinic.Nature medicine, 27(5):775–784, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.302220Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:093ee08e2ea9db01824822ee5a7e236a8336c6b40483984d64634baafe891f1e","observation_id":"bd0b8059-45fe-45f5-803e-cc5aa2d74574","resolution":{"observed_at":"2026-07-31T11:25:33.302220Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2601.03267","last_updated":"2026-05-01T23:55:43Z","snapshot_observed_at":"2026-08-02T10:52:10.211700Z","submitted_at":"2025-12-19T07:05:38Z","title":"OpenAI GPT-5 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2601.03267","snapshot_observed_at":"2026-07-31T11:25:33.304369Z","title":"Openai gpt-5 system card.arXiv preprint arXiv:2601.03267, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.304369Z"},"links":{"cited_paper":"/paper/2601.03267","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:5b68bdecf801c27d681bacd6795fdb772e444fb200ea0d389bb2c334c672bb54","observation_id":"43dce112-8cce-4730-98ce-14ca5469789e","resolution":{"observed_at":"2026-07-31T11:25:33.304369Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.10479","last_updated":"2025-04-19T03:47:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-14T17:59:25Z","title":"InternVL3: Exploring Advanced Training and Test-Time Recipes for Open-Source Multimodal Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.10479","snapshot_observed_at":"2026-07-31T11:25:33.306731Z","title":"Internvl3: Exploring advanced training and test-time recipes for open-source multimodal models.arXiv preprint arXiv:2504.10479, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.306731Z"},"links":{"cited_paper":"/paper/2504.10479","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:9d3e03e5d9eca69d8380e188c8cb6d46c2f2254a920d0c6ccb023cdb88a1d15c","observation_id":"97564667-adfe-4de1-b4fd-ef8a326ae823","resolution":{"observed_at":"2026-07-31T11:25:33.306731Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.18265","last_updated":"2025-08-27T14:39:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-08-25T17:58:17Z","title":"InternVL3.5: Advancing Open-Source Multimodal Models in Versatility, Reasoning, and Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.18265","snapshot_observed_at":"2026-07-31T11:25:33.309300Z","title":"Internvl3","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.309300Z"},"links":{"cited_paper":"/paper/2508.18265","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:54d04412747bde418824b3186521c52454451171d33c3d150a467e1de29643e1","observation_id":"7c41e68e-2bd2-42f4-942a-fcece5efe631","resolution":{"observed_at":"2026-07-31T11:25:33.309300Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.01006","last_updated":"2026-01-01T13:07:25Z","snapshot_observed_at":"2026-08-03T18:50:30.558321Z","submitted_at":"2025-07-01T17:55:04Z","title":"GLM-4.5V and GLM-4.1V-Thinking: Towards Versatile Multimodal Reasoning with Scalable Reinforcement Learning","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.01006","snapshot_observed_at":"2026-07-31T11:25:33.311621Z","title":"Glm-4.5 v and glm-4.1 v-thinking: Towards versatile multimodal reasoning with scalable reinforcement learning.arXiv preprint arXiv:2507.01006, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.311621Z"},"links":{"cited_paper":"/paper/2507.01006","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:39f24e4215ac302390dec342d52de542781c8861067912cc7d18e0af84d3d4ae","observation_id":"d89db2b3-e227-4a50-bd36-fe48366a8526","resolution":{"observed_at":"2026-07-31T11:25:33.311621Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05201","last_updated":"2026-04-06T19:50:20Z","snapshot_observed_at":"2026-08-07T23:53:34.814339Z","submitted_at":"2025-07-07T17:01:44Z","title":"MedGemma Technical Report","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.05201","snapshot_observed_at":"2026-07-31T11:25:33.314380Z","title":"Medgemma technical report","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.314380Z"},"links":{"cited_paper":"/paper/2507.05201","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:d2ae32bb1e635e6625a1a660c09f96f5e79b423d8eb9afe093f3f70d07e11850","observation_id":"9eac6142-4919-48ec-83d6-77f5d935a294","resolution":{"observed_at":"2026-07-31T11:25:33.314380Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2604.05081","last_updated":"2026-05-01T19:02:06Z","snapshot_observed_at":"2026-08-06T09:30:02.681376Z","submitted_at":"2026-04-06T18:35:57Z","title":"MedGemma 1.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2604.05081","snapshot_observed_at":"2026-07-31T11:25:33.316821Z","title":"Medgemma 1.5 technical report.arXiv preprint arXiv:2604.05081, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.316821Z"},"links":{"cited_paper":"/paper/2604.05081","citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:31917ff35946d39c1af3fa60f5ea8960c47ba6e209d853ae2cf138b71eb707ba","observation_id":"df2cc097-0eca-4162-8d90-b8dad3c02646","resolution":{"observed_at":"2026-07-31T11:25:33.316821Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:25:33.319256Z","title":"Swift: a scalable lightweight infrastructure for fine-tuning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-07-31T11:25:33.319256Z"},"links":{"citing_paper":"/paper/2607.28318"},"observation_digest":"sha256:a04af693f804f0d2c6798071019ac6500612f243964bf72157c1a021c2126031","observation_id":"c6f8719f-9d93-4c49-8dd5-6f54af641860","resolution":{"observed_at":"2026-07-31T11:25:33.319256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.28318","last_updated":"2026-07-30T14:53:51Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-08T09:23:17.196800Z","submitted_at":"2026-07-30T14:53:51Z","title":"PathView-Bench: Can Multimodal Large Language Models Achieve Fine-grained Multiscale Understanding of Pathology Images?"},"reference_resolution":{"displayed":53,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":53,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":53},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 53 of 53 outbound references and 0 inbound Pith citation observations for arXiv:2607.28318."}