{"as_of":"2026-08-10T06:26:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:37ef557b477047b1e7cf09b2e7cfb46f9446aa39bf8315fb488577590fa1e302","coverage":[{"denominator":86,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":86,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:41:45.954845Z","state":"measured"},{"denominator":86,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":86,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.01738/citation-record","integrity":"/paper/2506.01738/integrity","json":"/paper/2506.01738/citation-record.json","paper":"/paper/2506.01738"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:44.721620Z","title":"ITU-R Rec","venue":null,"work_id":null,"year":2000},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:44.721620Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:96d4dc09fe8061cbaacc9c227233a1c4f6079f4542d1ed3e6e180be1ae9592a0","observation_id":"48efcff6-b4d3-4560-a46e-732fc4eec187","resolution":{"observed_at":"2026-08-07T11:41:44.721620Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-07T11:41:45.183775Z","title":"GPT-4 technical report","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.183775Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:1929117fbcda37899a959806a59f9d9405da2f5b78505b23cf188b5c8826c105","observation_id":"64375bc7-713e-4ba1-b12d-4fead49fa700","resolution":{"observed_at":"2026-08-07T11:41:45.183775Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.16609","last_updated":"2023-09-28T17:07:49Z","snapshot_observed_at":"2026-08-09T21:25:20.369782Z","submitted_at":"2023-09-28T17:07:49Z","title":"Qwen Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.16609","snapshot_observed_at":"2026-08-07T11:41:45.190449Z","title":"Qwen technical report","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.190449Z"},"links":{"cited_paper":"/paper/2309.16609","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:c3947e892f8f0389ae8f1d9d87fe9c28e3143a17a4cda550078506ab583f5612","observation_id":"c26dd91f-91ef-4bec-b177-8a110403c479","resolution":{"observed_at":"2026-08-07T11:41:45.190449Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.12966","last_updated":"2023-10-13T02:41:28Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-24T17:59:17Z","title":"Qwen-VL: A Versatile Vision-Language Model for Understanding, Localization, Text Reading, and Beyond","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.12966","snapshot_observed_at":"2026-08-07T11:41:45.199828Z","title":"Qwen-vl: A frontier large vision-language model with versatile abilities","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.199828Z"},"links":{"cited_paper":"/paper/2308.12966","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:925bfe057b19e84a97f7e151ebd2d29f109168d06d9b1b0ceec1afaa51ed3029","observation_id":"5bd38879-55f5-4d1b-9dd6-d3c06e26c427","resolution":{"observed_at":"2026-08-07T11:41:45.199828Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:45.208211Z","title":"Face recognition and retrieval using cross-age reference coding with cross-age celebrity dataset","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.208211Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:a8398b8784f9067dcf99f6ef02f8189bc7dd59e82890bf014966c1b128748a53","observation_id":"32637ee5-b018-43cc-b171-1167dcf8e333","resolution":{"observed_at":"2026-08-07T11:41:45.208211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:45.217117Z","title":"Using ranking-CNN for age estimation","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.217117Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:4ad71b5ec4ceff767498f4071c591dfcad51f1e160db8cf86464f19d6fae3256","observation_id":"15592c2a-2529-48ab-a4e7-2c699772bb1f","resolution":{"observed_at":"2026-08-07T11:41:45.217117Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2209.06794","last_updated":"2023-06-05T17:55:12Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-09-14T17:24:07Z","title":"PaLI: A Jointly-Scaled Multilingual Language-Image Model","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.06794","snapshot_observed_at":"2026-08-07T11:41:45.227076Z","title":"PaLI: A jointly-scaled multilingual language-image model","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.227076Z"},"links":{"cited_paper":"/paper/2209.06794","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:710b0a58a8c9676a5bac8bfed4dc7c85f15c17306b3c256da5573f5988b26c33","observation_id":"1a5a414c-3f3f-4cc6-a5e0-a0dea53eb3e4","resolution":{"observed_at":"2026-08-07T11:41:45.227076Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.12307","last_updated":"2024-03-08T15:26:38Z","snapshot_observed_at":"2026-07-06T16:21:57.649550Z","submitted_at":"2023-09-21T17:59:11Z","title":"LongLoRA: Efficient Fine-tuning of Long-Context Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.12307","snapshot_observed_at":"2026-08-07T11:41:45.233393Z","title":"LongLoRA: Efficient fine-tuning of long-context large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.233393Z"},"links":{"cited_paper":"/paper/2309.12307","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:f8d968a52979590d77722d29942b25f98b314f84a7b41224bab56951aa0b554e","observation_id":"aa3f6c67-8b3f-4a8c-b7b3-37fdfa9c163b","resolution":{"observed_at":"2026-08-07T11:41:45.233393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:45.245613Z","title":"Vicuna: An open-source chatbot impressing GPT-4 with 90%* ChatGPT quality","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.245613Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:9f8785140ec111b083ee400a679ac981da5099c01d48695aea31f73d47a6e757","observation_id":"3ab120c8-ff92-4fa3-a548-c34f7f2adbb0","resolution":{"observed_at":"2026-08-07T11:41:45.245613Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:45.251589Z","title":"Soft labels for ordinal regression","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.251589Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:b08b6d32969ac38a7711b61bd0bc2b59b68a9400532bff5680cbd47765c8ac4b","observation_id":"db9b99a9-37fc-4434-a30d-644d91f55354","resolution":{"observed_at":"2026-08-07T11:41:45.251589Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-10T01:12:16.468283Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-07T11:41:45.258396Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.258396Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:3a928b9b7869f85ad13889ed6fc93f99238a2645dad25af2c01fa5aa4e399d41","observation_id":"da309fbb-8f42-4730-9f76-6cfe4d68660d","resolution":{"observed_at":"2026-08-07T11:41:45.258396Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:45.275974Z","title":"Teach CLIP to develop a number sense for ordinal regression","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.275974Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:adbd24054cd0f25dce002d5c21bb4e347bda7ef446af32100938c2302e6f9fc2","observation_id":"ea5582a3-46c1-4ca1-9f97-1a1d62e78854","resolution":{"observed_at":"2026-08-07T11:41:45.275974Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.432710Z","title":"Diabetic retinopathy detection (2015)","venue":null,"work_id":"fec5874f-595f-491f-afae-24597c8ce5bb","year":2015},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.289801Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:33014ef613706c24b85c33bd3b204c03d0bad3410091522c9f2721a821cc3bf8","observation_id":"8c226f87-0456-4b1d-9e68-395cb1ef0a5c","resolution":{"observed_at":"2026-08-07T11:42:08.440159Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.409752Z","title":"Perceptual quality assessment of smartphone photography","venue":null,"work_id":"7b2e0fd3-78b0-4070-bae4-707fc9390e76","year":2020},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.301808Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:d3d27051cb3cc890a15be133a98e9b78e570fcb70526b703407cc8d21ce68ce3","observation_id":"8bf600ab-8776-4b02-821b-84905882de5c","resolution":{"observed_at":"2026-08-07T11:42:08.416446Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.375003Z","title":"A simple approach to ordinal classification","venue":null,"work_id":"35984eca-dd21-4351-b8e7-61fc8d46135b","year":2001},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.307548Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:2def5540945edf54318bb27a4f9a5631bf7f8a8f3b420aebfd43d51b3b5a8e99","observation_id":"f78a2149-28fc-4545-9c07-66993fe318e5","resolution":{"observed_at":"2026-08-07T11:42:08.382346Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.351855Z","title":"Deep ordinal regression network for monocular depth estimation","venue":null,"work_id":"c931634c-fa39-4add-be04-e334e92214e9","year":2002},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.313409Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:552fe6bdb084c608c8e14d9636c31a3825fbd08c292b47bd7766f9cc0a959d6a","observation_id":"24e83b33-aa71-4fa0-8ffc-e42d6f7a8ee9","resolution":{"observed_at":"2026-08-07T11:42:08.359669Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.327266Z","title":"Facial age estimation by learning from label distribu- tions","venue":null,"work_id":"b8f41193-c08b-41be-a8b6-03434dcd13e5","year":2013},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.319080Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:c6a32b6c2760188d6356fb191f21993e3d1d529efc2b32228e413099a239880d","observation_id":"9c5a68cd-2626-4884-a9d0-6a05f6d20b59","resolution":{"observed_at":"2026-08-07T11:42:08.334343Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.307723Z","title":"Massive online crowdsourced study of subjective and objective picture quality","venue":null,"work_id":"4f5eec62-ba4f-427e-9605-09970de599a3","year":2015},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.325828Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:e80d91f61d4f6463fda26afff7ad0646e92dc9ad5c31016918130c4a0034c0a9","observation_id":"649e96d6-3324-450f-b7b5-b1f4cea2561b","resolution":{"observed_at":"2026-08-07T11:42:08.313631Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.288684Z","title":"A V A: A video dataset of spatio-temporally localized atomic visual actions","venue":null,"work_id":"7cd9df25-f381-4c4a-85ab-0641f41d4abe","year":2018},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.333105Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:f7ef335f0b5ed3ffd1e1efcbd8705740cda1c079e68de4e6da80c23c3d718c39","observation_id":"6a65e19f-7b0d-44cc-aa1e-24025548fbb1","resolution":{"observed_at":"2026-08-07T11:42:08.295027Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.262471Z","title":"Rethinking image aesthetics assessment: Models, datasets and benchmarks","venue":null,"work_id":"350df8a3-66c9-4fc3-b136-a6afd9a97ef3","year":2022},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.340197Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:d73473264b9230708397a2925275a33cf21749d0db184baa59a74be396fc2b50","observation_id":"8c9b9920-cef3-432d-967c-ab24b197084d","resolution":{"observed_at":"2026-08-07T11:42:08.271058Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.234498Z","title":"KonIQ-10k: An ecologically valid database for deep learning of blind image quality assessment","venue":null,"work_id":"e40395d1-c497-4803-a3e5-3975cb2431fb","year":2020},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.349147Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:036e4e3727a85feb4f8417cf11c7226ef3d74bb6cce161a3d922c93ef3b53194","observation_id":"1da7bc38-b928-46e8-a733-7b34303b89da","resolution":{"observed_at":"2026-08-07T11:42:08.245441Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.215938Z","title":"LoRA: Low-rank adaptation of large language models","venue":null,"work_id":"306072e5-e917-4d0c-9615-7f6bfdefb292","year":2022},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.355864Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:f4d2b5c956eca61c61a2d625864414c16cd8c33d2417d69210f87aace5318333","observation_id":"39c9ac2c-97d3-4230-953d-1473c0e0e4c4","resolution":{"observed_at":"2026-08-07T11:42:08.221921Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.195062Z","title":"and Tamirat Tesafaye","venue":null,"work_id":"3f34438f-47a4-4e6e-8007-4e0436b79568","year":2006},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.361364Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:6b8e4432d460a6ceb44df2cbfccea274d0bd19d5f2ba173039436a784579c421","observation_id":"985d1775-92d3-4c97-bd3e-c530f280ea3a","resolution":{"observed_at":"2026-08-07T11:42:08.201986Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.162661Z","title":"APTOS 2019 blindness detection","venue":null,"work_id":"06be0c99-57e3-42dc-939b-2eb23f2f1a8f","year":2019},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.368840Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:9104c24e0ea4391e0f087ee58714d52169caf0932a1dce530d296c607a6af0d1","observation_id":"59b23d50-0a81-49ec-9d74-b6ea1e1e2429","resolution":{"observed_at":"2026-08-07T11:42:08.169669Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.136563Z","title":"Generating images with multimodal language models","venue":null,"work_id":"7164ab34-1868-46d5-bb04-0eb7333a9fa9","year":2024},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.375581Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:5b39f2ac403cffe838c2a30d4a74c3645c816b677289b580539f0365266eaec9","observation_id":"8c7517d0-cd52-4539-b581-0f1c95718f29","resolution":{"observed_at":"2026-08-07T11:42:08.145067Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.00692","last_updated":"2024-05-01T05:10:13Z","snapshot_observed_at":"2026-08-06T07:43:56.889679Z","submitted_at":"2023-08-01T17:50:17Z","title":"LISA: Reasoning Segmentation via Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.00692","snapshot_observed_at":"2026-08-07T11:41:45.382441Z","title":"LISA: Reasoning segmentation via large language model","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.382441Z"},"links":{"cited_paper":"/paper/2308.00692","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:ae2ed61bef6f71d467d41e829e52d4070af994189ded5e1b5356e9544a2ed562","observation_id":"11113140-16c8-4c23-abf7-a76ebf07b908","resolution":{"observed_at":"2026-08-07T11:41:45.382441Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.106043Z","title":"Deep repulsive clustering of ordered data based on order- identity decomposition","venue":null,"work_id":"ef6c8e18-2d59-4d08-8b5d-f4f02ae3db37","year":2020},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.388773Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:079bab7792d8456d92fe475e5215d45f238e47852e96d4064da92d049b8b37e2","observation_id":"61b40e72-401c-46d6-9a4b-816daf181223","resolution":{"observed_at":"2026-08-07T11:42:08.114275Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.079607Z","title":"Age and gender classification using convolutional neural networks","venue":null,"work_id":"9d2eab95-e764-49ea-bf31-998903af5245","year":2015},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.396693Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:9fe9c0df091f3f05da8533b41b84f3deed390439975953930e481a09c85c78b4","observation_id":"113dcc0b-fad4-46ac-8cf9-748009a0180e","resolution":{"observed_at":"2026-08-07T11:42:08.088154Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.12597","last_updated":"2023-06-15T07:57:29Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-01-30T00:56:51Z","title":"BLIP-2: Bootstrapping Language-Image Pre-training with Frozen Image Encoders and Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.12597","snapshot_observed_at":"2026-08-07T11:41:45.411264Z","title":"BLIP-2: Bootstrapping language- image pre-training with frozen image encoders and large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.411264Z"},"links":{"cited_paper":"/paper/2301.12597","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:f69f9eb079bf90433c0194f846a0e54550edd09436df07a8633aa5ea4ba1fdfb","observation_id":"a836edaf-6aff-4960-be37-2024856d5f8e","resolution":{"observed_at":"2026-08-07T11:41:45.411264Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.054011Z","title":"BLIP: Bootstrapping language- image pre-training for unified vision-language understanding and generation","venue":null,"work_id":"d4d08150-34b1-47cc-9898-3d3440fe9e68","year":2022},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.416381Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:482c9ba5665a35da04386a9f2d245ec48814f7ba51f135778d768a0fda2224bc","observation_id":"f3fc91f8-d063-44f5-a936-c406e46da75c","resolution":{"observed_at":"2026-08-07T11:42:08.061376Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.032693Z","title":"Ordinal regression by extended binary classification","venue":null,"work_id":"39bfa1dc-36c3-4e7a-9fc0-1e45828e7230","year":2006},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.420421Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:e5a5346d3918a38ac274cd2f6c0372579ddb5d226578d14a3c90f4890684d0f8","observation_id":"6482919f-604b-4732-877d-10ed51c8d8bc","resolution":{"observed_at":"2026-08-07T11:42:08.039739Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:08.006786Z","title":"Learning probabilistic ordi- nal embeddings for uncertainty-aware regression","venue":null,"work_id":"46a160f5-bc18-4c81-a664-f41b362d0b0b","year":2021},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.427262Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:043408508071eda82f70076a6123751b99e66fe124a552dd66ed719df6165e5f","observation_id":"a9fe9de4-a93b-44d3-8c5d-1da25cbe8919","resolution":{"observed_at":"2026-08-07T11:42:08.017074Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.982105Z","title":"OrdinalCLIP: Learning rank prompts for language-guided ordinal regression","venue":null,"work_id":"eaa15ba1-d0d1-4a01-b6ed-467bab2a8f49","year":2022},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.432178Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:02b4754d14acd97c4446c697de100f1c2feda4f5c85f8f86b01595681eca72ca","observation_id":"b2d9d653-4984-43c2-9b5e-d940b56c2178","resolution":{"observed_at":"2026-08-07T11:42:07.989243Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.17043","last_updated":"2023-11-28T18:53:43Z","snapshot_observed_at":"2026-08-06T05:08:17.428238Z","submitted_at":"2023-11-28T18:53:43Z","title":"LLaMA-VID: An Image is Worth 2 Tokens in Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.17043","snapshot_observed_at":"2026-08-07T11:41:45.444716Z","title":"LLaMA-VID: An image is worth 2 tokens in large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.444716Z"},"links":{"cited_paper":"/paper/2311.17043","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:97679773bd8786389da5bbe790c13f8eac7cf161a0bbad8fffc2a9096193c573","observation_id":"fe0280de-c759-4176-aace-871854251637","resolution":{"observed_at":"2026-08-07T11:41:45.444716Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.956712Z","title":"Order learning and its application to age estimation","venue":null,"work_id":"5eed3f88-41d2-4a31-a1bd-d9c06ca31569","year":2019},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.450459Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:379c544429b7f4b578f19ef4c48ede828110745a7a7dd839fd4f7204ade002e5","observation_id":"a8f9fd7c-c5e6-4541-aff0-6861c7c48f53","resolution":{"observed_at":"2026-08-07T11:42:07.964083Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:45.459209Z","title":"Improved baselines with visual instruction tuning, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.459209Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:7f7438bffdb1cb818e0572c0cdfe872d192c6b6b57cd114a17914ee95703162c","observation_id":"99318978-0529-4e63-aa6d-15e2379e510d","resolution":{"observed_at":"2026-08-07T11:41:45.459209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:45.465201Z","title":"Visual instruction tuning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.465201Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:7e83d90a2f69b7846c23850a3be81731138f3d6404f2d62732734199c53ca413","observation_id":"d04adafd-e419-4e3b-8384-63bd96708557","resolution":{"observed_at":"2026-08-07T11:41:45.465201Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.900825Z","title":"DeepDRiD: Diabetic retinopathy—grading and image quality estimation challenge","venue":null,"work_id":"dcc4fa65-8c7e-4d10-9c8b-128b48c70548","year":2022},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.470887Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:b9a6dd861dbe63801a33c5331d3f030df7b0ac495d07f94fb25bf07adcd535f7","observation_id":"f05ad454-8b4e-474d-8b1e-1dd97ba32072","resolution":{"observed_at":"2026-08-07T11:42:07.911017Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.875029Z","title":"Deep ordinal regression based on data relationship for small datasets","venue":null,"work_id":"387a77ee-0f42-47b6-81d6-3c64b4ddbd4c","year":2017},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.476642Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:9ae5951b2cbcfe00037fedfcbaf72ce5ae98eab42583b7e1421fbaefacbf649c","observation_id":"a51759ff-de48-4db9-8acd-b1465fc88e45","resolution":{"observed_at":"2026-08-07T11:42:07.882916Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.849689Z","title":"A constrained deep neural network for ordinal regression","venue":null,"work_id":"af6648fe-3970-49fb-b2c4-3adf18d5fdd1","year":2018},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.482841Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:297e20657fc24b5a460d1950acbd71231a685a90c4e23cc8caaf9577fe344373","observation_id":"12ab238b-f708-479f-96a0-72cbd0e00117","resolution":{"observed_at":"2026-08-07T11:42:07.858694Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.829186Z","title":"Probabilistic deep ordinal regression based on Gaussian processes","venue":null,"work_id":"29ea38fc-6a79-471b-8845-ac7dafc69c9c","year":2019},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.489203Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:b53dad980affae7ecf0b92cfee522e2712a621fcdaa10ffc2b7a3273a03c9057","observation_id":"32379c1d-cf7e-42e1-a5a1-db02b10451d9","resolution":{"observed_at":"2026-08-07T11:42:07.836480Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.808687Z","title":"Cheap and quick: Efficient vision-language instruction tuning for large language models","venue":null,"work_id":"a648b79c-a04e-4c76-98a1-61831bf17543","year":2024},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.496057Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:9c57a5dfc5df77534e1ceeb837e9f250c9ee6f9afb096bacd87bcfaeb8a11c3a","observation_id":"f3d060e5-8a46-4f45-aa6c-994af158da1c","resolution":{"observed_at":"2026-08-07T11:42:07.815390Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.788142Z","title":"Dating historical color images","venue":null,"work_id":"c3f5e9d9-5c5e-446d-b5ac-ca6e547b0669","year":2012},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.503343Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:8cc015407078ab5a9d196c8399ebc1a6ad45e146a6d9ce01506e397d9035e8a1","observation_id":"4049f57e-15af-4d21-b845-67b4af4e8ae6","resolution":{"observed_at":"2026-08-07T11:42:07.795007Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.00750","last_updated":"2023-03-01T18:59:33Z","snapshot_observed_at":"2026-08-09T15:21:12.878321Z","submitted_at":"2023-03-01T18:59:33Z","title":"StraIT: Non-autoregressive Generation with Stratified Image Transformer","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.00750","snapshot_observed_at":"2026-08-07T11:41:45.509078Z","title":"StraIT: Non-autoregressive generation with stratified image Transformer","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.509078Z"},"links":{"cited_paper":"/paper/2303.00750","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:8f9ebd14d3970c4a55d70e02b4ed63b2957a4a2b31dfd0eb47f80e678fa55c4a","observation_id":"dbc733e0-8470-4bf6-bec1-56048caec544","resolution":{"observed_at":"2026-08-07T11:41:45.509078Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:45.513986Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.513986Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:c26b6ec3df53069ec0b7bfc4ee5495573fff2b70d578f719afbc9812f6f38858","observation_id":"7e1af7ad-deaa-41ad-aa0c-d9c05000eadd","resolution":{"observed_at":"2026-08-07T11:41:45.513986Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:45.520478Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.520478Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:fa4d3dc7855a6cb3f300136e1c72d267498457f1c6dcdc395baf61a19cd06b9b","observation_id":"3fe8e758-c5ea-4192-9e92-490f9f1182d8","resolution":{"observed_at":"2026-08-07T11:41:45.520478Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.734656Z","title":"Deep expectation of real and apparent age from a single image without facial landmarks","venue":null,"work_id":"8c924c1e-48f3-4815-b9bd-9885e6367002","year":2018},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.532502Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:58256b8e1c61ecc5f43a4b4146719965d3bcf27f751acd87b67404540d0ca8ba","observation_id":"40b90b9c-719d-4fb3-b11d-c17cbb65454f","resolution":{"observed_at":"2026-08-07T11:42:07.741135Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.715655Z","title":"HuggingGPT: Solving ai tasks with ChatGPT and its friends in hugging face","venue":null,"work_id":"18a4d0e4-8dc3-4123-89d6-40761c01a1f9","year":2024},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.538686Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:79a17743570407e681cbc316c030dbaa8c0ee172e1509a7bf9f098ca8ea5f5f8","observation_id":"a1a457e0-91e4-48f4-a7b8-0e3d694dfd4b","resolution":{"observed_at":"2026-08-07T11:42:07.720846Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.698648Z","title":"Moving window regression: A novel approach to ordinal regression","venue":null,"work_id":"853560be-9a5b-42c1-a0b1-c3122507559a","year":2022},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.544344Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:6871fe0aecff28692b10cc30082f644276f88ab04d2896ff0adde68fef228f78","observation_id":"40ead950-642b-44f2-98de-764ac4f42131","resolution":{"observed_at":"2026-08-07T11:42:07.704797Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-08-07T11:41:45.552758Z","title":"Gemini: A family of highly capable multimodal models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.552758Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:cee40dd98aa96fc078501b38ac088ed2213ed8ac90b7930e1663c3669335f6d7","observation_id":"1d3cd1f5-3aaa-496d-909a-079da2fbccb0","resolution":{"observed_at":"2026-08-07T11:41:45.552758Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-07T11:41:45.557980Z","title":"LLaMA: Open and efficient foundation language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.557980Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:27725e72fc7097b71fb9d9bc28e2057c79feafc4f1181f5c88459a7df1942482","observation_id":"9cbd0fc9-79b0-4e4f-84db-c627044890c4","resolution":{"observed_at":"2026-08-07T11:41:45.557980Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.676902Z","title":"Ord2Seq: Regarding ordinal regression as label sequence prediction","venue":null,"work_id":"bb36bd79-c4e2-4c92-89e0-e197c32f2110","year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.565046Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:7c184bff5c261d8d5f4fe5141d6e4d99c874efbe9640849917e44aa5eb339add","observation_id":"a3f923ab-af1c-4e9c-abf7-01b0eee8fdec","resolution":{"observed_at":"2026-08-07T11:42:07.684603Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.653001Z","title":"Learning-to- rank meets language: Boosting language-driven ordering alignment for ordinal classification","venue":null,"work_id":"f6e6ef7b-7f5c-4f85-96e9-ca398a1f3114","year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.569880Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:06302bf72013ae6ee9190b54b5c4aac60ef369f193bb8aa8ded3959a0af954c9","observation_id":"29f86852-257e-40d0-9f73-bf1aa80748a5","resolution":{"observed_at":"2026-08-07T11:42:07.660806Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.03079","last_updated":"2024-02-04T08:23:04Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-11-06T13:04:39Z","title":"CogVLM: Visual Expert for Pretrained Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.03079","snapshot_observed_at":"2026-08-07T11:41:45.575021Z","title":"CogVLM: Visual expert for pretrained language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.575021Z"},"links":{"cited_paper":"/paper/2311.03079","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:65ab1e33f171a51abebb2e62993b4a594022513bb0fe208749922ea3f10d8c0c","observation_id":"3b8c0aeb-1a16-42fb-97e6-11fe07fa9861","resolution":{"observed_at":"2026-08-07T11:41:45.575021Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.632040Z","title":"VisionLLM: Large language model is also an open-ended decoder for vision-centric tasks","venue":null,"work_id":"073f0683-086e-4434-b1b5-da13cc695736","year":2024},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.580149Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:2bc10f90d9d8e570c6f8480a977601c4928626a08bac9cd5a950963e7b91f48f","observation_id":"f506f633-26e2-4ae6-bae8-859db88fab4f","resolution":{"observed_at":"2026-08-07T11:42:07.638795Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.04671","last_updated":"2023-03-08T15:50:02Z","snapshot_observed_at":"2026-07-06T15:00:13.368355Z","submitted_at":"2023-03-08T15:50:02Z","title":"Visual ChatGPT: Talking, Drawing and Editing with Visual Foundation Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.04671","snapshot_observed_at":"2026-08-07T11:41:45.584826Z","title":"Visual ChatGPT: Talking, drawing and editing with visual foundation models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.584826Z"},"links":{"cited_paper":"/paper/2303.04671","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:c7fbbc6c0032755771e6d79301a222d1dc9235b118b8cc10b5eb29979b350d09","observation_id":"1896a52f-be55-47eb-8670-0135fce85bea","resolution":{"observed_at":"2026-08-07T11:41:45.584826Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.611849Z","title":"Q-Bench: A benchmark for general-purpose foundation models on low-level vision","venue":null,"work_id":"2094c7b1-e717-402b-8d36-a582b1f0a3d6","year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.591105Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:a283394f420bb11249e4773f74103954d5061adde0aa723c7d5dfce287fbb748","observation_id":"e0bbe672-2ec6-4d8e-b133-43e072b6d21b","resolution":{"observed_at":"2026-08-07T11:42:07.617607Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.06783","last_updated":"2023-11-12T09:10:51Z","snapshot_observed_at":"2026-07-06T16:46:10.142093Z","submitted_at":"2023-11-12T09:10:51Z","title":"Q-Instruct: Improving Low-level Visual Abilities for Multi-modality Foundation Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.06783","snapshot_observed_at":"2026-08-07T11:41:45.598200Z","title":"Q-Instruct: Improving low-level visual abilities for multi-modality foundation models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.598200Z"},"links":{"cited_paper":"/paper/2311.06783","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:2254379d3220bbe8bc274a03bedbdd3b620e6d0967f53f0715e22a7c9be5540d","observation_id":"d46953ed-94ab-4b7f-9676-3bdd2604202f","resolution":{"observed_at":"2026-08-07T11:41:45.598200Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.17090","last_updated":"2023-12-28T16:10:25Z","snapshot_observed_at":"2026-08-02T07:14:02.308302Z","submitted_at":"2023-12-28T16:10:25Z","title":"Q-Align: Teaching LMMs for Visual Scoring via Discrete Text-Defined Levels","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.17090","snapshot_observed_at":"2026-08-07T11:41:45.603653Z","title":"Q-Align: Teaching LMMs for visual scoring via discrete text-defined levels","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.603653Z"},"links":{"cited_paper":"/paper/2312.17090","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:a48de330a4ee20bca0d5a6c7f67e2c3681d1670db8d088879060c6831acfe3a6","observation_id":"84741159-a626-41a1-a0a5-87d135d9fb88","resolution":{"observed_at":"2026-08-07T11:41:45.603653Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.11381","last_updated":"2023-03-20T18:31:47Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-20T18:31:47Z","title":"MM-REACT: Prompting ChatGPT for Multimodal Reasoning and Action","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.11381","snapshot_observed_at":"2026-08-07T11:41:45.608949Z","title":"MM-REACT: Prompting ChatGPT for multimodal reasoning and action","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.608949Z"},"links":{"cited_paper":"/paper/2303.11381","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:027670c142c23b7c406d1e651171fa79a1cb335fe1b0e95d001a1f58227e8992","observation_id":"5473ffa8-f4fe-4703-8b8c-60865f78cf0e","resolution":{"observed_at":"2026-08-07T11:41:45.608949Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.02858","last_updated":"2023-10-25T06:23:31Z","snapshot_observed_at":"2026-07-06T15:38:39.712379Z","submitted_at":"2023-06-05T13:17:27Z","title":"Video-LLaMA: An Instruction-tuned Audio-Visual Language Model for Video Understanding","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.02858","snapshot_observed_at":"2026-08-07T11:41:45.614543Z","title":"Video-LLaMA: An instruction-tuned audio-visual language model for video understanding","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.614543Z"},"links":{"cited_paper":"/paper/2306.02858","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:f4e64517a19ef0d03392387ad73a30e30be2de49710c586f8836c9cc46ebe1ba","observation_id":"7aa4ca83-908f-4536-b33d-39453273bcd2","resolution":{"observed_at":"2026-08-07T11:41:45.614543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.03601","last_updated":"2025-06-12T00:15:18Z","snapshot_observed_at":"2026-08-07T18:00:59.339869Z","submitted_at":"2023-07-07T13:43:44Z","title":"GPT4RoI: Instruction Tuning Large Language Model on Region-of-Interest","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.03601","snapshot_observed_at":"2026-08-07T11:41:45.624615Z","title":"GPT4RoI: Instruction tuning large language model on region-of-interest","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.624615Z"},"links":{"cited_paper":"/paper/2307.03601","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:0e0947cae658aa01602344ea25a30edb370d5df5fa1351f44ba82fa9330a8d68","observation_id":"ffe94371-9feb-4a6e-becc-1e14a431a0cb","resolution":{"observed_at":"2026-08-07T11:41:45.624615Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.592536Z","title":"Age progression/regression by conditional adver- sarial autoencoder","venue":null,"work_id":"ddac1cd1-2457-4b5d-8315-a638b2c888ea","year":2017},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.630287Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:186fa7834f04a737d9c9ee8c4bd7ac2e11cadd886e3a51d035bbc67c168cbd63","observation_id":"564fd652-22a4-4708-b7b7-8a337c2951de","resolution":{"observed_at":"2026-08-07T11:42:07.598929Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.10592","last_updated":"2023-10-02T16:38:35Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-20T18:25:35Z","title":"MiniGPT-4: Enhancing Vision-Language Understanding with Advanced Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.10592","snapshot_observed_at":"2026-08-07T11:41:45.636103Z","title":"MiniGPT-4: Enhancing vision-language understanding with advanced large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.636103Z"},"links":{"cited_paper":"/paper/2304.10592","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:9fbc8973f16db601f944dcf1976eaeb908f704506a250cd4691da9c06b2f3bb1","observation_id":"83d50c66-a655-4fc1-80d8-e27d84467c49","resolution":{"observed_at":"2026-08-07T11:41:45.636103Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.13046","last_updated":"2024-10-31T17:39:34Z","snapshot_observed_at":"2026-08-05T02:05:40.291312Z","submitted_at":"2024-04-19T17:59:48Z","title":"MoVA: Adapting Mixture of Vision Experts to Multimodal Context","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.13046","snapshot_observed_at":"2026-08-07T11:41:45.645267Z","title":"MoV A: Adapting mixture of vision experts to multimodal context.arXiv preprint arXiv:2404.13046, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.645267Z"},"links":{"cited_paper":"/paper/2404.13046","citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:463c5a481cf11bffacc345770c64369803e2d4e57f97037e8b46c62870f1b6c2","observation_id":"b081692f-0a49-4923-9733-6e860d28dc38","resolution":{"observed_at":"2026-08-07T11:41:45.645267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.571316Z","title":null,"venue":null,"work_id":"4f050ec4-f19c-41d0-953f-8839d232ffd3","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.658787Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:5b4378955921df52d76f5259a07a66d99e9f31d4140720241e19accddb200053","observation_id":"496e8e50-3964-4823-8868-44b4310db3dd","resolution":{"observed_at":"2026-08-07T11:42:07.578133Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:45.665243Z","title":"(a) Did you state the full set of assumptions of all theoretical results? [NA] (b) Did you include complete proofs of all theoretical results? [NA]","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.665243Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:507f94d4bba2546b132ff47a195a1c5cc4c2e9beddb0219816b79ebf1fc77a5b","observation_id":"1142e179-4943-42af-a3d8-81bc5ea1ca6e","resolution":{"observed_at":"2026-08-07T11:41:45.665243Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.529647Z","title":"for benchmarks)","venue":null,"work_id":"931ee3fd-16c6-4b8a-8bb0-088216ef9585","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.674844Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:e68d5f66cf4bd60831fc971ed80623e6624fe6360eb6cc8fba1635b806391a01","observation_id":"c42ac8c5-e19b-470d-8ad3-c531492b1ae4","resolution":{"observed_at":"2026-08-07T11:42:07.535976Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.504424Z","title":null,"venue":null,"work_id":"012fb426-27df-49fa-b85f-f5ace4262a39","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.684663Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:162a49ffdec8708cc3a0188011292176d1d25b35c1c4c652816e09ce24eb5848","observation_id":"86d22b7b-b7a4-4b70-bc07-c3653a0d7a5f","resolution":{"observed_at":"2026-08-07T11:42:07.510874Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.482281Z","title":null,"venue":null,"work_id":"bb1d52c2-761d-45bb-9dbf-490a9f8c46c0","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.693235Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:72bd7aa74203fdf5572e8edecc0fd96bd094435c6f6e3d0966e1eda85373280c","observation_id":"23e7e30c-9d78-4db0-b27d-5cbf4195a047","resolution":{"observed_at":"2026-08-07T11:42:07.489071Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.438730Z","title":null,"venue":null,"work_id":"d185fb0c-e5f1-4a43-94ee-50498786a10f","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.705215Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:d47329b3f419398c990982f12ac579b8612cff39844dd77d033db6e34e0c5c8e","observation_id":"4604e6fb-c779-4566-9d37-410edb684d02","resolution":{"observed_at":"2026-08-07T11:42:07.446414Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.419103Z","title":null,"venue":null,"work_id":"cf5f5240-d07d-4a18-ae80-2844fcd05bad","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.710360Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:3d60e92e86c0f8b3b467b6a2cc8b0d8c6ec36b06ca9fee7830f9baa4d23d2a00","observation_id":"df9f9520-4d43-4c3d-b6f0-0a8b61caf085","resolution":{"observed_at":"2026-08-07T11:42:07.424456Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.399010Z","title":null,"venue":null,"work_id":"43e7355d-acbe-4244-816e-d47eb14e1519","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.847637Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:b190da278cb586c959d7d63ea9712636fbe1abd6418e7371f978e3fbe84d77bd","observation_id":"4b5387d3-9e16-4442-ab5c-103de4c64e3e","resolution":{"observed_at":"2026-08-07T11:42:07.404663Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.364406Z","title":null,"venue":null,"work_id":"fb0fabcc-e231-43d1-8cb1-0a29c646b65a","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.854530Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:b4a9cf857d107d0db4c867dad5c3586292317381866a43f275966d52eb926ffb","observation_id":"f20801d4-b345-4311-8bf3-c24ea2c6424b","resolution":{"observed_at":"2026-08-07T11:42:07.376035Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.329539Z","title":null,"venue":null,"work_id":"efc2852c-3511-49f9-b918-703a11402e2f","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.871507Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:893fef91f9f963ad3c269564e4d800edaa82c051c03606b97347df942b4e94cb","observation_id":"881dc3e8-d910-46e3-a13b-1c21df372d46","resolution":{"observed_at":"2026-08-07T11:42:07.336940Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.307975Z","title":null,"venue":null,"work_id":"bf790495-d978-4593-a33f-7f1d8a88ea9a","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.876846Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:edec23e8072b3fee78579f8bc966cea5142e740076cfe725d434d6842fe65b53","observation_id":"9d1349cf-50a5-4a9e-81ff-874683d6d11d","resolution":{"observed_at":"2026-08-07T11:42:07.315762Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.287999Z","title":null,"venue":null,"work_id":"62c0d888-72cc-40fc-9bee-5fb1063294e8","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.884500Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:18ffd039e1ebcfe010904fa14f19714e136b1d347239e2774cee32fa4fed9652","observation_id":"dbf26c28-7f69-4017-8d4a-7a32ff5375a5","resolution":{"observed_at":"2026-08-07T11:42:07.295063Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.265911Z","title":null,"venue":null,"work_id":"3c37883d-3d8a-482c-a238-0b46a4759aae","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.898081Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:c174dbae778bbf85d6005c01ce3ec1a176e85fe37ab39da74c819e0d23fdc31e","observation_id":"6551b1a3-cc60-48bf-8925-8c11429ef6a4","resolution":{"observed_at":"2026-08-07T11:42:07.273526Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.245200Z","title":null,"venue":null,"work_id":"fdba72e8-2289-47c1-89ba-da1a2bde93fb","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.905549Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:12e65155234628a48093ffe7e8c16db459285f93d3bf9ae62e6fae7fe3cb8bdb","observation_id":"29592a58-8a6a-4c43-a050-f1e34544656e","resolution":{"observed_at":"2026-08-07T11:42:07.251845Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.219375Z","title":null,"venue":null,"work_id":"85700b6d-fedf-4c67-bc59-8ed4c1a8b4ac","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.912776Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:3d9364d5fb4449e397082ae1f5f2aba6625923cdb07c2c38956f6d30556b0d91","observation_id":"6017a24f-302a-4416-a205-db0d13c5732d","resolution":{"observed_at":"2026-08-07T11:42:07.228201Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.192580Z","title":null,"venue":null,"work_id":"a41afbbd-7cdd-4c84-9eb0-fca4f6a89024","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.926489Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:a406417aa022f4d3e895966dd1cca1cf5684fdfb32817eb398d4db66869679ab","observation_id":"b73a12e4-3897-4c57-8c4e-3be948b6fa48","resolution":{"observed_at":"2026-08-07T11:42:07.200225Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.167001Z","title":null,"venue":null,"work_id":"28b50404-b91c-41a7-8057-f09342e1843c","year":1930},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.933298Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:cb374812eaf19030b7e32de0c9aaf1bb1ac2da5386636be12faf1dffb52f9999","observation_id":"c13abde2-f530-4d1f-8d5f-d0f8876480c6","resolution":{"observed_at":"2026-08-07T11:42:07.173786Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.139588Z","title":null,"venue":null,"work_id":"fe3b244d-5de6-43a8-b719-01b08d272f90","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.939904Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:3bf93a821547867f2f51fd177a29b32ff48a0dbe8cfa9919032148ea3bd51a48","observation_id":"74b7eb04-54a5-44a8-9c13-373ede481376","resolution":{"observed_at":"2026-08-07T11:42:07.146131Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:05.886714Z","title":null,"venue":null,"work_id":"b5397bd5-7028-40af-9224-2bd4dfc8248a","year":1930},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.947388Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:7089cbd6e3c346a740b83fec42ce279168464c92f5034bbf2774ac19ad04b49d","observation_id":"857c7449-5f96-456e-8db7-3c8a2a7d01a0","resolution":{"observed_at":"2026-08-07T11:42:07.119073Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:41:46.522427Z","title":"Answer: [Coarse answer], [Predicted Phase] 11 F Limitations The definitions of labels for different domain tasks are quite diverse","venue":null,"work_id":"161dbe56-2814-4051-8958-081bbd8a9611","year":null},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.954845Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:ed42ecbc10cd60e7537cb25c0443084905ff240ca827f57b8d687945deabf23a","observation_id":"7a549368-dff7-449a-a3c1-ada6bbbb9513","resolution":{"observed_at":"2026-08-07T11:42:02.306459Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:42:07.461085Z","title":"To conserve GPU memory during fine-tuning, we employ FSDP (Full Shard Data Parallel) with ZeRO3-style","venue":null,"work_id":"60f0b6ac-4486-4027-9c15-bfa83a0205e0","year":2014},"citing_paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset","version":1},"reference_index":128,"source":"pdf_text","source_observed_at":"2026-08-07T11:41:45.699257Z"},"links":{"citing_paper":"/paper/2506.01738"},"observation_digest":"sha256:531dc0d566e667bc244a544a5e13fa19fa0852b96cb845b15188222f777ee321","observation_id":"7949144e-9d0a-43b4-b040-0f3824226b57","resolution":{"observed_at":"2026-08-07T11:42:07.468031Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.01738","last_updated":"2025-06-02T14:48:15Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-09T15:21:33.651698Z","submitted_at":"2025-06-02T14:48:15Z","title":"STORM: Benchmarking Visual Rating of MLLMs with a Comprehensive Ordinal Regression Dataset"},"reference_resolution":{"displayed":86,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":49,"verified_exact":0,"verified_fuzzy":37},"total_outbound_references":86},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 86 of 86 outbound references and 0 inbound Pith citation observations for arXiv:2506.01738."}