{"as_of":"2026-08-09T23:16:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6b699475faa35166542ad8aff620ceb1ee14ae31c4b3ea7002080f64c2c7ddf3","coverage":[{"denominator":38,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":38,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:22:39.581294Z","state":"measured"},{"denominator":39,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":39,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:22:39.462293Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-07T11:22:39.638306Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"cited_work":{"arxiv_id":"2506.02708","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.02708","snapshot_observed_at":"2026-08-07T11:22:39.638306Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","venue":"cs.CV","work_id":"d02e2601-daa1-420c-a57e-2ed5b152d80d","year":2025},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.462293Z"},"links":{"cited_paper":"/paper/2506.02708","citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:3bf892c084f21a13ff96f3a656daa83e993a5f6fa2bc06f8317d58f1d45b75e8","observation_id":"2d14ec3e-2093-4db5-8a10-afb076786963","resolution":{"observed_at":"2026-08-07T11:22:39.644055Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.02708/citation-record","integrity":"/paper/2506.02708/integrity","json":"/paper/2506.02708/citation-record.json","paper":"/paper/2506.02708"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"cited_work":{"arxiv_id":"2506.02708","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.02708","snapshot_observed_at":"2026-08-07T11:22:39.638306Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","venue":"cs.CV","work_id":"d02e2601-daa1-420c-a57e-2ed5b152d80d","year":2025},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.462293Z"},"links":{"cited_paper":"/paper/2506.02708","citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:3bf892c084f21a13ff96f3a656daa83e993a5f6fa2bc06f8317d58f1d45b75e8","observation_id":"2d14ec3e-2093-4db5-8a10-afb076786963","resolution":{"observed_at":"2026-08-07T11:22:39.644055Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.989981Z","title":"IAA datasets (e.g., A V A [3] and AADB [4]) consist of photographs scored by human annotators","venue":null,"work_id":"7997a939-45ec-43b4-b63a-41972c18a3d7","year":null},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.466279Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:fa51184a91dfec33e8f4e982c87d1182559c7be0fdbda85f0bc33765edbb9fdb","observation_id":"2ed680fc-593e-4dd6-b113-e3cb9919d7a4","resolution":{"observed_at":"2026-08-07T11:22:39.993197Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.981227Z","title":"chosen” and “rejected","venue":null,"work_id":"d36c29ca-bb0b-4a96-9dec-6e475d727b6b","year":null},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.470068Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:83c6be49ed9ed4aa9f84a394c25665d291c62eb91dd9ae3f54d0331121c02aa3","observation_id":"d0cefc2e-78b5-4ffc-81df-db73e2e4be8b","resolution":{"observed_at":"2026-08-07T11:22:39.984013Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.972080Z","title":"Datasets As IAA datasets, we use A V A [3] and AADB [4]","venue":null,"work_id":"adabfbfe-fc36-4c4c-a758-fb5021ab2547","year":2024},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.473651Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:88117ac308f425875f4a5814bd3cf2f3a147a5db85b684b8317fe518e76c1321","observation_id":"4f956278-0b1b-41a5-a61c-0cb590db9b00","resolution":{"observed_at":"2026-08-07T11:22:39.975082Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.962093Z","title":"Main Results Table 1 presents the results of our method applied to the LLaV A-NeXT-7B model on the A V A dataset, a part of which is also plotted in Fig","venue":null,"work_id":"ce0d7a08-3315-4b8e-97f7-a3f7cf71bf98","year":null},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.477154Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:77968c42cd4993574787f4e51d334d294ddd02b0bc55b5b309d74d9b4ed2f1d7","observation_id":"d154969d-a14e-4f9f-bc31-524d2980b052","resolution":{"observed_at":"2026-08-07T11:22:39.966050Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.952348Z","title":null,"venue":null,"work_id":"63874961-c3dc-4267-97c3-524dbba9e60f","year":null},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.480624Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:8b7b64db0a941f86dca9814218dd1f5b401b68bb5c7021b7096ef351b9eae982","observation_id":"61266e04-615f-43eb-abd5-27118940a0e8","resolution":{"observed_at":"2026-08-07T11:22:39.955560Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2006.11371","last_updated":"2020-06-23T01:48:56Z","snapshot_observed_at":"2026-08-08T09:58:42.167795Z","submitted_at":"2020-06-16T02:58:10Z","title":"Opportunities and Challenges in Explainable Artificial Intelligence (XAI): A Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2006.11371","snapshot_observed_at":"2026-08-07T11:22:39.483907Z","title":"Opportunities and challenges in explainable artificial intelligence (xai): A survey,","venue":null,"work_id":null,"year":2006},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.483907Z"},"links":{"cited_paper":"/paper/2006.11371","citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:b290ebeb78f4007276e4e9e7607ab70d2dcf26f4f8bc573d5a133a13cbf77fc0","observation_id":"1c806957-c1bc-4ce1-a21f-14b2cd3f5338","resolution":{"observed_at":"2026-08-07T11:22:39.483907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.943155Z","title":"Di- rect preference optimization: your language model is secretly a reward model,","venue":null,"work_id":"3dc075eb-ab7a-41f4-8188-b27197bc241a","year":2023},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.487350Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:9fd79471922d8673ddfbfa6c063be7f932cd93a1c72ebdfddd3bb42bf8dbee26","observation_id":"6f17033b-efbc-443f-b088-961e8e222ce9","resolution":{"observed_at":"2026-08-07T11:22:39.946167Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.933869Z","title":"Ava: A large-scale database for aesthetic visual analy- sis,","venue":null,"work_id":"af6842b4-4b1f-44e6-98fb-626303d43d35","year":2012},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.490408Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:4988eacd61fe4223b81036e4cd719ca5093ed883c0a619e0e9ff661ed6887171","observation_id":"82622785-d97b-4c65-b9a4-80d9c7d07411","resolution":{"observed_at":"2026-08-07T11:22:39.936942Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.924639Z","title":"Photo aesthetics ranking network with attributes and content adaptation,","venue":null,"work_id":"bcf5cd0e-90d2-4d1b-bcb0-0d7523d675e9","year":2016},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.493391Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:2475bc3cc2b929edbaf8419b136c3a53ac544bd3d736bd4c383a588b0101d101","observation_id":"2b5a600b-eecc-42df-a7e5-e780e169823b","resolution":{"observed_at":"2026-08-07T11:22:39.927803Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.915957Z","title":"Image aesthetic assessment: An experimental survey,","venue":null,"work_id":"4383bddd-2cbc-46f7-be04-57c8e374ec4e","year":2017},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.496724Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:eb6273e991b1903238d2bb25ec4e8aaa2b892837d63a0f22c8620c44fa4e7767","observation_id":"db186e22-6a14-405b-93b7-a10073d233c0","resolution":{"observed_at":"2026-08-07T11:22:39.919018Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.906524Z","title":"Vila: Learning image aes- thetics from user comments with vision-language pre- training,","venue":null,"work_id":"724c94cd-235f-4b4d-a70e-943143abbc30","year":2023},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.499831Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:a4e4f55115814b966172533c3c2a8a7ea97935720aa4fa265e885057d880c6c2","observation_id":"d2e0ca25-5989-4e1e-ab28-8230dd6166ae","resolution":{"observed_at":"2026-08-07T11:22:39.909522Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.897964Z","title":"Q-align: Teach- ing LMMs for visual scoring via discrete text-defined levels,","venue":null,"work_id":"72f74092-b8cc-4b54-bff5-2a78a4ca5b10","year":2024},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.502993Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:f1c611c6a469a57b71ddf28fd86d452699910f837df57c1f92f25b5eb6dfcdf7","observation_id":"9c5efb31-a617-4d67-a0e8-a8ca548135cb","resolution":{"observed_at":"2026-08-07T11:22:39.900889Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.888884Z","title":"Chain-of-thought prompting elicits reasoning in large language models,","venue":null,"work_id":"51deda1c-1f36-46bb-927c-539a838a9ea2","year":2022},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.506200Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:a959724aa4bab7eaa604f9cdb947932759188dd42975c9d398239f95d4dc1cb4","observation_id":"8f456c4b-6661-45af-be63-7522482f3b63","resolution":{"observed_at":"2026-08-07T11:22:39.892057Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.879902Z","title":"Multimodal explanations: Justifying de- cisions and pointing to the evidence,","venue":null,"work_id":"86bc9370-9760-49f9-ae0d-2e2851cb515c","year":2018},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.509131Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:700c92f49bd56e1b9bff9e949fe9b5147e1b001164aca458e7f433bdd0ed07a1","observation_id":"2029f5e8-406e-46f5-a8c8-d1c05c741f62","resolution":{"observed_at":"2026-08-07T11:22:39.882817Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.14546","last_updated":"2020-04-30T02:20:14Z","snapshot_observed_at":"2026-08-08T18:22:05.258756Z","submitted_at":"2020-04-30T02:20:14Z","title":"WT5?! Training Text-to-Text Models to Explain their Predictions","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.14546","snapshot_observed_at":"2026-08-07T11:22:39.512099Z","title":"Wt5?! training text-to-text models to explain their predictions,","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.512099Z"},"links":{"cited_paper":"/paper/2004.14546","citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:b8a75a4a49193116d5f85c20f2b17047e4ae1768ff1aad822fcecb7164308f0a","observation_id":"5673ca2f-f78d-4879-9147-6e596ce96ab1","resolution":{"observed_at":"2026-08-07T11:22:39.512099Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.870909Z","title":"Measuring association between labels and free-text ra- tionales,","venue":null,"work_id":"3e00569f-c681-410d-9f9c-a72cd773c287","year":2021},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.515418Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:e89841d766e522a20feb7ef522684c6371d942b5d932f1422028a8b8d2774e5f","observation_id":"ad6da841-b9bd-4c49-95c0-ed565071435c","resolution":{"observed_at":"2026-08-07T11:22:39.874019Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.860667Z","title":"Few-shot self-rationalization with natural language prompts,","venue":null,"work_id":"e88de83f-c052-4860-8524-2783f13befb0","year":2022},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.518373Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:62d8f6fb394e47a4523d3ee97a4e856ddc05847a09de7373ef169b979c61d1f9","observation_id":"3d03dce8-5ad1-4684-8103-5460ec01d802","resolution":{"observed_at":"2026-08-07T11:22:39.863981Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.851976Z","title":"Reframing human- ai collaboration for generating free-text explanations,","venue":null,"work_id":"a29fff39-e18d-40be-94c3-b63836de8fdb","year":2022},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.521342Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:2543b41b4e15563cc66120d9dcd2a9f23ed1c46f352331f615590a62faf5ff30","observation_id":"a44cd642-adf4-4651-92fc-d1a6ddababa9","resolution":{"observed_at":"2026-08-07T11:22:39.855095Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.843307Z","title":"Large language models can self-improve,","venue":null,"work_id":"434bfbce-e188-48a6-a913-64359f75c7be","year":2023},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.524875Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:c0d1d0b9a0905aa601dc9503d4196732772f6d5a19ee26e502fa206fe2c52d9b","observation_id":"ad0369cc-d466-4107-a54e-ac317df08b42","resolution":{"observed_at":"2026-08-07T11:22:39.846348Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.834265Z","title":"Self-refine: iterative refinement with self-feedback,","venue":null,"work_id":"ea982088-945a-4eab-b5fb-c44d8c7d2d39","year":2023},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.527683Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:5397c243f0a168f0f39c870d6ecfc5fa93a04b9bb2833bb36309af93debec5d1","observation_id":"07d8dcea-0925-4a37-b2bf-5716232228b3","resolution":{"observed_at":"2026-08-07T11:22:39.837216Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.824318Z","title":"Self-rewarding language models,","venue":null,"work_id":"8deb43b7-bd26-4196-a534-f388321208a5","year":2024},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.530544Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:fbd64688c7835ccfa5dddd9e1047dc335c636f74fde37b4cf551d2243bcda33d","observation_id":"2829e04c-99cb-42c8-a077-7f354de3a44e","resolution":{"observed_at":"2026-08-07T11:22:39.828007Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.814351Z","title":"Iterative reasoning preference optimization,","venue":null,"work_id":"93e0061b-4ecf-43d5-b13e-a0931fec2ad9","year":2024},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.533483Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:be0438cf5f8d4bf1520aa7d45aaeda38659c6263bd27fae737aa5c525bd3a367","observation_id":"bc1c3e5f-368c-4f23-89cd-462e691966c9","resolution":{"observed_at":"2026-08-07T11:22:39.817727Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.804053Z","title":"Enhancing large vision language models with self-training on image comprehension,","venue":null,"work_id":"231d4342-e35a-453c-a769-04ff20f8dab5","year":2024},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.536376Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:fb47bab8f9055eebd4ef5919491127e4cd6bea6c187c93c3b9b331d7aee460d5","observation_id":"299e03f4-a8c4-463d-bd06-8f77d8d84019","resolution":{"observed_at":"2026-08-07T11:22:39.807526Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.793468Z","title":"Ties-merging: resolving in- terference when merging models,","venue":null,"work_id":"fdac057c-29c7-48b6-807b-f0bccf616c26","year":2023},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.539337Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:47c1c79b02155b0ff0af1cf001d2a50ddc173e401c827d1507b3b097acf9cb50","observation_id":"1c97a69f-b9b5-4734-bb92-2f0283b2d956","resolution":{"observed_at":"2026-08-07T11:22:39.797117Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.782490Z","title":"Nima: Neural image assessment,","venue":null,"work_id":"6a890518-4ccf-493b-97a5-7ab3d4724d41","year":2018},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.542387Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:c3880d5d14bff5c2deb02d80f2695a639a126622cd40dd3b292876d6b17f42cb","observation_id":"1488d0b1-0aac-4909-85c7-4f6544fdcbbf","resolution":{"observed_at":"2026-08-07T11:22:39.786075Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.771663Z","title":"Llava- next-interleave: Tackling multi-image, video, and 3d in large multimodal models,","venue":null,"work_id":"e3ea9caa-e867-4c43-8cf7-65d3b5135bcf","year":2024},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.545790Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:051d1ace94b0c73dd825d8ed3a2c04e74c1f425c0af4e6dc2204d985ef8b2add","observation_id":"e9b41c18-baa8-414e-9a52-7767347f148e","resolution":{"observed_at":"2026-08-07T11:22:39.775329Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.760982Z","title":"Internvl: Scaling up vi- sion foundation models and aligning for generic visual- linguistic tasks,","venue":null,"work_id":"54db34ec-5af2-42df-ad0b-6a658ca02ed7","year":2024},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.548993Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:15460b1628d9193e0fc0d1eb627556ecea6ad13a4a586578fbb82b19ae34f408","observation_id":"e8c2f30f-c856-4a62-8df6-cf265bfdc897","resolution":{"observed_at":"2026-08-07T11:22:39.764830Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.749523Z","title":"Improved baselines with visual instruction tun- ing,","venue":null,"work_id":"64e86d65-9092-4e9b-a4a6-07f1e08617d9","year":2024},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.552025Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:1282286c28088eab8de17481c7bdc71c65174da8af71a111d4dd89f25315bcdd","observation_id":"f998c6e0-e86f-4f45-b827-f04f267bfcac","resolution":{"observed_at":"2026-08-07T11:22:39.753499Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.738159Z","title":"Llava- next: Improved reasoning, ocr, and world knowl- edge,","venue":null,"work_id":"14db29ce-28e3-4b20-a9b0-012934e144b7","year":2024},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.555089Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:a6f319ef7d7c48e897d3c12829f679e9c97424710e9f9a0d1d9673eb3880e3c0","observation_id":"c06aba25-5c77-4954-a84e-1d1cdd220dca","resolution":{"observed_at":"2026-08-07T11:22:39.741525Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.727387Z","title":"LoRA: Low-rank adaptation of large language models,","venue":null,"work_id":"df4bd32a-170e-4d1f-84f2-6e3fa8857d37","year":2022},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.558258Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:251ebd3c1b56f9a19f02db65cee0838c3c72d76006d474022d4ef12f6ed2c9fb","observation_id":"9af17de3-0583-4a8f-8e3d-1c9f7708f031","resolution":{"observed_at":"2026-08-07T11:22:39.731063Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.716635Z","title":"Musiq: Multi-scale image quality transformer,","venue":null,"work_id":"041bb95f-bf1f-49a3-8ced-c7629a6f774e","year":2021},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.561156Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:221890b27e49f6c6a8614db11c70a91283a540860309fd5cb02a3e0d8734d39d","observation_id":"6a5de88f-dabe-45e2-ba61-986fdeda7801","resolution":{"observed_at":"2026-08-07T11:22:39.720158Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.705790Z","title":"This could be achieved by using a higher megapixel camera or a camera with better low-light capabilities","venue":null,"work_id":"1738f822-c8f4-4ca3-90eb-d9deaa23aaad","year":null},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.564721Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:c3964f7186034b6fb234720a735fa7c40a63a56802c833ae92783ed02c0603b3","observation_id":"afa70c55-5d1e-46c5-870e-caad3cb39951","resolution":{"observed_at":"2026-08-07T11:22:39.709989Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.694898Z","title":"Natural light is often preferred for food photography, as it can create a more appealing and natural look","venue":null,"work_id":"77640c00-e34f-439f-a308-470e3e1e2a1e","year":null},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.568061Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:600b9458808288e14b6631acf8d9431b85ce1d0e45661d2a41b6a0392ca385eb","observation_id":"85d6f7f8-b259-4e38-ada3-5c5160260a68","resolution":{"observed_at":"2026-08-07T11:22:39.698301Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.684216Z","title":"Consider the rule of thirds and the use of negative space to create a more balanced and visually appealing image","venue":null,"work_id":"d6ece0ec-508b-414b-bad0-b091cb90b5dc","year":null},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.571989Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:a6cd269d471f880b40bb4fdce3b01bf6d60c44a46f3622916c90da5bdd40d03f","observation_id":"1a86cca3-b04c-4bbc-b476-cc50d4e9ae10","resolution":{"observed_at":"2026-08-07T11:22:39.687945Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.672659Z","title":"This can help draw attention to the main subject of the image","venue":null,"work_id":"994f018e-7cd1-40d3-afe0-2068e27ecff2","year":null},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.575232Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:5dbdcb523b236b2190d8ac4dc1244ed604af8e5ea4974ff5400c40031f47b33e","observation_id":"9b7adff3-b65d-4a04-ae04-407561c8797e","resolution":{"observed_at":"2026-08-07T11:22:39.676315Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.661347Z","title":"This can be done using photo editing software","venue":null,"work_id":"86aa2365-31d6-4152-bf29-6094b8b3af10","year":null},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.578134Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:c9e674980812103781b5ef8bedf5081724cd14148edcb1461fc3bf2ad1086652","observation_id":"7f568fa4-573f-43c6-901a-642a4da335f5","resolution":{"observed_at":"2026-08-07T11:22:39.664944Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:22:39.651375Z","title":"This could include adjusting the exposure, contrast, and saturation to make the colors pop and the image more vibrant","venue":null,"work_id":"ad99c614-ff31-42d2-80d4-2c87a52a3d4e","year":null},"citing_paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:39.581294Z"},"links":{"citing_paper":"/paper/2506.02708"},"observation_digest":"sha256:0b4e7c0763ecc3764715a038cde00d4ae9220be96c5f17532001a726163818fd","observation_id":"6d47d51f-f11b-4494-a38d-5880d631303e","resolution":{"observed_at":"2026-08-07T11:22:39.654556Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.02708","last_updated":"2025-06-03T10:04:19Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-07T11:15:25.307392Z","submitted_at":"2025-06-03T10:04:19Z","title":"Iterative Self-Improvement of Vision Language Models for Image Scoring and Self-Explanation"},"reference_resolution":{"displayed":38,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":3,"verified_exact":0,"verified_fuzzy":34},"total_outbound_references":38},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 38 of 38 outbound references and 1 inbound Pith citation observation for arXiv:2506.02708."}