{"as_of":"2026-08-12T07:38:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:60afb25f5f59d0da952a0fe7e8c0b5767516c1a551769e7a7883db48de2f4cf2","coverage":[{"denominator":44,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":44,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T19:36:21.792561Z","state":"measured"},{"denominator":50,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":50,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":6,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":6,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T18:48:44.077282Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-01T09:25:40.548668Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09927","snapshot_observed_at":"2026-08-06T18:48:44.077282Z","title":"Ie-bench: Advancing the measurement of text- driven image editing for human perception alignment","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.07317","last_updated":"2025-07-28T17:21:56Z","snapshot_observed_at":"2026-08-09T23:43:22.969472Z","submitted_at":"2025-07-09T22:29:47Z","title":"ADIEE: Automatic Dataset Creation and Scorer for Instruction-Guided Image Editing Evaluation","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-06T18:48:44.077282Z"},"links":{"cited_paper":"/paper/2501.09927","citing_paper":"/paper/2507.07317"},"observation_digest":"sha256:f73b833316ff5fa89a2184ab6371b984e3044e791274be6d909716c92de53b2a","observation_id":"08b2c90e-c4ba-43fd-98e1-1110031dbea4","resolution":{"observed_at":"2026-08-06T18:48:44.077282Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09927","snapshot_observed_at":"2026-08-06T15:19:57.777876Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.16193","last_updated":"2025-09-07T09:10:06Z","snapshot_observed_at":"2026-08-07T01:00:51.428857Z","submitted_at":"2025-07-22T03:11:07Z","title":"LMM4Edit: Benchmarking and Evaluating Multimodal Image Editing with LMMs","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-06T15:19:57.777876Z"},"links":{"cited_paper":"/paper/2501.09927","citing_paper":"/paper/2507.16193"},"observation_digest":"sha256:69acc447900452a574cf127dc205c9b8574945a3042fdfe984c7387217553ece","observation_id":"6b84952a-e902-471d-ae7b-088eb86ac2c0","resolution":{"observed_at":"2026-08-06T15:19:57.777876Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"cited_work":{"arxiv_id":"2501.09927","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.09927","snapshot_observed_at":"2026-07-01T09:25:40.548668Z","title":"arXiv preprint arXiv:2501.09927 (2025) DSH-Bench: A comprehensive benchmark for Subject-Driven T2I 19","venue":null,"work_id":"9eac7c49-8345-4155-a28d-2123c7d046db","year":2022},"citing_paper":{"arxiv_id":"2603.08090","last_updated":"2026-06-30T03:35:03Z","snapshot_observed_at":"2026-08-11T17:26:06.161529Z","submitted_at":"2026-03-09T08:30:28Z","title":"DSH-Bench: A Difficulty- and Scenario-Aware Benchmark with Hierarchical Subject Taxonomy for Subject-Driven Text-to-Image Generation","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-05-15T15:16:59.801347Z"},"links":{"cited_paper":"/paper/2501.09927","citing_paper":"/paper/2603.08090"},"observation_digest":"sha256:e938a8ccd86ee7c17bc418cc3261eb8981bd6b92aee84b6fdfa3a59e10e250ec","observation_id":"8157b4ed-8e00-4056-a8c3-33f50e314ac8","resolution":{"observed_at":"2026-05-15T15:20:09.134175Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09927","snapshot_observed_at":"2026-07-15T12:51:20.969593Z","title":"arXiv preprint arXiv:2501.09927 (2025) DSH-Bench: A comprehensive benchmark for Subject-Driven T2I 19","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.08090","last_updated":"2026-06-30T03:35:03Z","snapshot_observed_at":"2026-08-11T17:26:06.161529Z","submitted_at":"2026-03-09T08:30:28Z","title":"DSH-Bench: A Difficulty- and Scenario-Aware Benchmark with Hierarchical Subject Taxonomy for Subject-Driven Text-to-Image Generation","version":3},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-07-15T12:51:20.969593Z"},"links":{"cited_paper":"/paper/2501.09927","citing_paper":"/paper/2603.08090"},"observation_digest":"sha256:7d175b2b6e53dbbe7ab03411dc5701e3569b05c7e91a24038577c7e93deb3282","observation_id":"69c85a52-3c01-4786-a059-385d3abe06a5","resolution":{"observed_at":"2026-07-15T12:51:20.969593Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"cited_work":{"arxiv_id":"2501.09927","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.09927","snapshot_observed_at":"2026-07-01T09:25:40.548668Z","title":"arXiv preprint arXiv:2501.09927 (2025) DSH-Bench: A comprehensive benchmark for Subject-Driven T2I 19","venue":null,"work_id":"9eac7c49-8345-4155-a28d-2123c7d046db","year":2022},"citing_paper":{"arxiv_id":"2604.24023","last_updated":"2026-06-19T10:09:08Z","snapshot_observed_at":"2026-07-06T23:10:07.178566Z","submitted_at":"2026-04-27T04:11:06Z","title":"ServImage: An Image Generation and Editing Benchmark from Real-world Commercial Imaging Services","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-08T04:57:27.709861Z"},"links":{"cited_paper":"/paper/2501.09927","citing_paper":"/paper/2604.24023"},"observation_digest":"sha256:a0272f51f8f003f98385d6782d802cf536c4705ea87301877ac500d1510c8ed5","observation_id":"f906f36a-ffb6-4402-ac74-d35917afde2d","resolution":{"observed_at":"2026-05-11T21:36:14.802353Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"cited_work":{"arxiv_id":"2501.09927","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.09927","snapshot_observed_at":"2026-07-01T09:25:40.548668Z","title":"arXiv preprint arXiv:2501.09927 (2025) DSH-Bench: A comprehensive benchmark for Subject-Driven T2I 19","venue":null,"work_id":"9eac7c49-8345-4155-a28d-2123c7d046db","year":2022},"citing_paper":{"arxiv_id":"2604.24023","last_updated":"2026-06-19T10:09:08Z","snapshot_observed_at":"2026-07-06T23:10:07.178566Z","submitted_at":"2026-04-27T04:11:06Z","title":"ServImage: An Image Generation and Editing Benchmark from Real-world Commercial Imaging Services","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-01T09:07:27.666877Z"},"links":{"cited_paper":"/paper/2501.09927","citing_paper":"/paper/2604.24023"},"observation_digest":"sha256:958192cf1d54c83cf7a9205c8c4a47935339d9e5fbaf3b995b2a2e628855e58e","observation_id":"c7aca730-da12-4871-9790-d4eb0c64734f","resolution":{"observed_at":"2026-07-01T09:25:40.551336Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2501.09927/citation-record","integrity":"/paper/2501.09927/integrity","json":"/paper/2501.09927/citation-record.json","paper":"/paper/2501.09927"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.536729Z","title":"In- structpix2pix: Learning to follow image editing instructions","venue":null,"work_id":"7f12f3c3-37cf-4737-b2a1-a65f7bacc9c4","year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.566470Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:d9d15ffdcf6cefa8d2b740319a04ee9bcba412086dcba5cfdb8edd7865256301","observation_id":"502f9cec-b1a6-434e-9f3a-4d0b918cb344","resolution":{"observed_at":"2026-08-10T19:36:22.542055Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.520130Z","title":"Masactrl: Tuning-free mu- tual self-attention control for consistent image synthesis and editing","venue":null,"work_id":"4ba630a9-c72f-464a-8e56-6cda89de6ebe","year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.571552Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:d31bca74361d10bafb117db0c6a29bda96226315d708a5a2e14f1680e5e0fef9","observation_id":"34204ceb-1a72-4e33-bd80-aecdbba03b85","resolution":{"observed_at":"2026-08-10T19:36:22.525600Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:21.577092Z","title":"Diffusion models beat gans on image synthesis","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.577092Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:591107fabfebcfaca03c3673b9ec8987b2793f140d08c71518851c07512236f6","observation_id":"ea9b4dbc-7fc2-4800-8daa-fd108c4c45e5","resolution":{"observed_at":"2026-08-10T19:36:21.577092Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:21.582588Z","title":"Gpt-3: Its nature, scope, limits, and consequences","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.582588Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:091533d060a8014a6dfd7e64e8a1fd587019ca8b61b776936efe444ad775729a","observation_id":"66817db1-f33e-4b3c-9be1-6dee1742a650","resolution":{"observed_at":"2026-08-10T19:36:21.582588Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.483044Z","title":"Gans trained by a two time-scale update rule converge to a local nash equilib- rium","venue":null,"work_id":"71b7b9c0-5b73-4ffe-b930-a38dd772607d","year":2017},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.588075Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:664d1521caa482e09cf020b1439143668bc7e207f3815604e9959ec3df1dc32e","observation_id":"4f52388c-5b2a-4d72-b060-d511495df727","resolution":{"observed_at":"2026-08-10T19:36:22.488181Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:21.593454Z","title":"Denoising dif- fusion probabilistic models","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.593454Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:d51d22784c1253d5cd2371a37eba5698e58535f2f88b7e266d1a60992f19a51c","observation_id":"649b2363-0593-4b23-b93c-3496dfe05590","resolution":{"observed_at":"2026-08-10T19:36:21.593454Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.17525","last_updated":"2025-03-11T02:16:29Z","snapshot_observed_at":"2026-08-11T07:46:04.646035Z","submitted_at":"2024-02-27T14:07:09Z","title":"Diffusion Model-Based Image Editing: A Survey","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.17525","snapshot_observed_at":"2026-08-10T19:36:21.605637Z","title":"Diffusion model-based image editing: A survey","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.605637Z"},"links":{"cited_paper":"/paper/2402.17525","citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:57fa74a24e583d7820df1544199461367aae6da39eac2f90996126fa41e2e4c3","observation_id":"5f988682-1955-4131-bcbf-dc5d45a8f24c","resolution":{"observed_at":"2026-08-10T19:36:21.605637Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.456408Z","title":"Smartedit: Exploring complex instruction-based image editing with multimodal large lan- guage models","venue":null,"work_id":"04d08077-20db-4557-b61d-03ae5d26c732","year":2024},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.610812Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:30830347d199a3a6393ffb4a7685c47e4c999ca767fc6fcbb1bdac2dc8bb544f","observation_id":"1c523f88-23c5-44a4-9ab0-ab2a632a07b0","resolution":{"observed_at":"2026-08-10T19:36:22.461295Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.439979Z","title":"Methodology for the subjective as- sessment of the quality of television pictures itu-r recommen- dation","venue":null,"work_id":"70dac757-9be3-4d63-b9e2-8ca30b260bf1","year":2000},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.615507Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:cd9ce3453b0b04c10f25c4c0335549b8ce00b67a8b9378e9d14d6443c10b33e0","observation_id":"deea53be-d91d-4a58-980b-22c75c29f7c9","resolution":{"observed_at":"2026-08-10T19:36:22.445483Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.421169Z","title":"Convolu- tional neural networks for no-reference image quality assess- ment","venue":null,"work_id":"29d75409-a7ab-46b1-a6d5-3fb178f77d55","year":2014},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.625706Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:135358b99baa0cef9756a12c49c3a4f5cf4c95441bfc5632a6742496372e725f","observation_id":"2a45546d-1e03-4850-b179-ccc13d82ac9a","resolution":{"observed_at":"2026-08-10T19:36:22.427072Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.403145Z","title":"The tum high definition video datasets","venue":null,"work_id":"b2c533d1-d0aa-4001-8cea-cd2177c99b36","year":2012},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.630440Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:d5793b90b2e26c2155d0a69cdaf34f5903594dd08bd1d7ba244eb820ae78dfa4","observation_id":"5eff281a-a5dd-4cde-8340-462e6e916c0a","resolution":{"observed_at":"2026-08-10T19:36:22.408782Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6980","last_updated":"2017-01-30T01:27:54Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2014-12-22T13:54:29Z","title":"Adam: A Method for Stochastic Optimization","version":9},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.6980","snapshot_observed_at":"2026-08-10T19:36:21.634926Z","title":"Adam: A method for stochastic optimization","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.634926Z"},"links":{"cited_paper":"/paper/1412.6980","citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:bd4514554b7f5adfc4951f1d1456ffa6f2d6732de516466c8af634847c9faed9","observation_id":"619decb3-1569-4208-9c25-5f29d2c48b72","resolution":{"observed_at":"2026-08-10T19:36:21.634926Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.386823Z","title":"Pick-a-pic: An open dataset of user preferences for text-to-image generation","venue":null,"work_id":"c6c7ef73-a443-4d74-928c-1423b0e39fdf","year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.639859Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:6dd5ad9730426b617144d470525ccb13bac7aef4c662a9ac062c463b8ed79543","observation_id":"5fb24ac3-f994-474c-ab9d-3884278b86ed","resolution":{"observed_at":"2026-08-10T19:36:22.392184Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.11956","last_updated":"2024-08-07T17:02:00Z","snapshot_observed_at":"2026-07-06T17:46:25.142387Z","submitted_at":"2024-03-18T16:52:49Z","title":"Subjective-Aligned Dataset and Metric for Text-to-Video Quality Assessment","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.11956","snapshot_observed_at":"2026-08-10T19:36:21.644637Z","title":"Subjective-aligned dateset and metric for text-to-video qual- ity assessment","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.644637Z"},"links":{"cited_paper":"/paper/2403.11956","citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:4c4f9e9e876821e929b972be7fdeb967966042514c9b8896cf950e959b6650a6","observation_id":"914a7771-8066-4d7e-9c79-0462a2fc1460","resolution":{"observed_at":"2026-08-10T19:36:21.644637Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.370469Z","title":"Most apparent dis- tortion: full-reference image quality assessment and the role of strategy","venue":null,"work_id":"42fac38a-5dea-4a4f-a026-049c573dcebc","year":2010},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.649740Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:ef6171981628036951b6b24da606e558feb2f0e5dab90352a5c8cc1424deb4de","observation_id":"64035f23-d5fe-4885-b78d-a8e86a20b28d","resolution":{"observed_at":"2026-08-10T19:36:22.375709Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.354101Z","title":"Agiqa-3k: An open database for ai-generated image quality assessment","venue":null,"work_id":"01cd0ef3-1fff-4918-972c-d88a8535683d","year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.654324Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:c9444a5c46d5641efd2e1581bba2a069bfdaa061de9547f47329118b4e440fa2","observation_id":"09969bec-b1bc-4e95-bf4b-d022ae718b07","resolution":{"observed_at":"2026-08-10T19:36:22.359346Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.03407","last_updated":"2024-04-04T12:12:24Z","snapshot_observed_at":"2026-07-06T17:55:36.748672Z","submitted_at":"2024-04-04T12:12:24Z","title":"AIGIQA-20K: A Large Database for AI-Generated Image Quality Assessment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.03407","snapshot_observed_at":"2026-08-10T19:36:21.658974Z","title":"Aigiqa-20k: A large database for ai-generated image quality assessment","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.658974Z"},"links":{"cited_paper":"/paper/2404.03407","citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:b8ee20b54ed1866ecd3c5491f9b4436762f94636bcf122345e3df7a14d2c2f16","observation_id":"4516808b-8054-4d1f-91b8-4b685a3b08b9","resolution":{"observed_at":"2026-08-10T19:36:21.658974Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.337271Z","title":"Norm-in-norm loss with faster convergence and better performance for im- age quality assessment","venue":null,"work_id":"24aec8ce-b634-47aa-8a7f-839f2d07d537","year":2020},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.663812Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:6edd4dde8eb733494cf8997eba1683cae1ef5f85430b73a550ce8839345dabb7","observation_id":"ac781497-b60b-409f-8209-bcd5f08615dd","resolution":{"observed_at":"2026-08-10T19:36:22.342852Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:21.668577Z","title":"Microsoft coco: Common objects in context","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.668577Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:57ddafc7b30b48bc4f82174860351eb1e1a09a3d3f6d7e4a07edce733a97da6a","observation_id":"2a79ab37-51c8-47ac-8884-85b328cabdcc","resolution":{"observed_at":"2026-08-10T19:36:21.668577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.309688Z","title":"Referring image editing: Object-level image editing via referring expressions","venue":null,"work_id":"5faab7e4-10a0-4b7a-bfba-8e258993fcfc","year":2024},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.673014Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:2c22da77d6bd317a6c64ccd6bb074c4295a1dbc032afd0fb1e71f2cf8f61edde","observation_id":"f151ae3f-cc5c-4bac-8fd7-dab40770ce0c","resolution":{"observed_at":"2026-08-10T19:36:22.314930Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.293786Z","title":"Dinov2: Learning robust visual features without super- vision","venue":null,"work_id":"7a76072c-da17-4df0-8f6d-e60635a26b7d","year":2024},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.677586Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:3bf03594199965cbfe0a163ce317ff37bf95dae9327830ca9eacd1e479dead19","observation_id":"e32e6631-46bc-4e1c-809f-d9751c26bb30","resolution":{"observed_at":"2026-08-10T19:36:22.298737Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.18714","last_updated":"2024-05-21T06:57:24Z","snapshot_observed_at":"2026-07-06T17:52:00.714976Z","submitted_at":"2024-03-27T16:02:00Z","title":"Bringing Textual Prompt to AI-Generated Image Quality Assessment","version":2},"cited_work":{"arxiv_id":"2403.18714","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.18714","snapshot_observed_at":"2026-08-10T19:36:21.886929Z","title":"Bringing Textual Prompt to AI-Generated Image Quality Assessment","venue":"cs.CV","work_id":"102bd8a7-c540-4f48-b53b-59a8c6275230","year":2024},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.682423Z"},"links":{"cited_paper":"/paper/2403.18714","citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:f8eb06c273885dfa3737880e35022188cd5591cc1b803542559fcc4f2683d972","observation_id":"6b8ff5a6-3b87-4d7f-b889-1d4b20837920","resolution":{"observed_at":"2026-08-10T19:36:21.894630Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.277016Z","title":"Learning transferable visual models from natural language supervi- sion","venue":null,"work_id":"354d970a-cf34-4603-9a13-6f1fb60760ba","year":2021},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.687317Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:cbc65675f46453841e799dbecec2866eb72c2fb717f839db297a7b383a6eb06d","observation_id":"9c5efc64-3486-4784-a8a6-7adbf9f8616c","resolution":{"observed_at":"2026-08-10T19:36:22.282731Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:21.692017Z","title":"Dreambooth: Fine tuning text-to-image diffusion models for subject-driven generation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.692017Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:40532aeb6e5d8348b6d796f53321f4fea2debcf77a04da6e4cc8a53620ea50a1","observation_id":"968e2f5b-5794-4271-843c-33aa147aa981","resolution":{"observed_at":"2026-08-10T19:36:21.692017Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.248867Z","title":"Methodology for the subjective assessment of the quality of television pictures","venue":null,"work_id":"4d486423-66b0-433d-94e2-d0849a469dbb","year":2002},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.696631Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:b83394acf2724fe20926d570758f6ec76adc4ed3c9f46adae8215b571c0fec88","observation_id":"4140a877-dad5-4f4c-b16e-68d747d49a6e","resolution":{"observed_at":"2026-08-10T19:36:22.254218Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.227524Z","title":"Doubly abductive coun- terfactual inference for text-based image editing","venue":null,"work_id":"57d472b5-6d7a-47ae-a861-9e2ab20084d1","year":2024},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.701289Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:d5a152125a86c8a557dca78e93f5502e17c3ca51f2bbf1532a3ad83ff94b5529","observation_id":"185e7f0f-d713-49ab-ad19-af887855770f","resolution":{"observed_at":"2026-08-10T19:36:22.232998Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.211038Z","title":"Blindly assess image qual- ity in the wild guided by a self-adaptive hyper network","venue":null,"work_id":"5cd8b80b-7687-443a-9c51-9ad46853a70e","year":2020},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.705910Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:14c920ed579be4540c1e011ca76cd0934000d1d42a1185b7ba6cceb3a7fff9b3","observation_id":"d70009f6-405a-409f-b656-c2c70831031d","resolution":{"observed_at":"2026-08-10T19:36:22.216295Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.11481","last_updated":"2024-12-18T09:55:40Z","snapshot_observed_at":"2026-08-09T04:37:01.020759Z","submitted_at":"2024-08-21T09:49:32Z","title":"VE-Bench: Subjective-Aligned Benchmark Suite for Text-Driven Video Editing Quality Assessment","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.11481","snapshot_observed_at":"2026-08-10T19:36:21.710559Z","title":"E-bench: Subjective-aligned benchmark suite for text-driven video editing quality assessment","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.710559Z"},"links":{"cited_paper":"/paper/2408.11481","citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:79ec409dba71fa48d7668579826ea6e2fd4c27ce555165e334550b42a67af0d5","observation_id":"8d13b8be-9317-4c3b-9142-d73a83d1721b","resolution":{"observed_at":"2026-08-10T19:36:21.710559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.193785Z","title":"Improved artgan for conditional synthesis of natural image and artwork","venue":null,"work_id":"f2dc9127-f66e-4b55-9faa-c396fc12472b","year":2019},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.716003Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:035012a12ebe66ab27062181b31884ca387f6f0305151d17cd4cd1b3e88aad77","observation_id":"aeb39a3a-f678-488f-9f33-752f100bfbc1","resolution":{"observed_at":"2026-08-10T19:36:22.199251Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.177585Z","title":"Full-reference image quality as- sessment by combining features in spatial and frequency do- mains","venue":null,"work_id":"c3e0b741-4521-4de3-b6b1-ba37761b32a3","year":null},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.720781Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:a32daaea4019da0e4feabf5e36649fee71ed39016c1467fa19ec942e3e843e01","observation_id":"da4870b4-742b-4c2d-b74a-6f94a7006f1a","resolution":{"observed_at":"2026-08-10T19:36:22.182825Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.162299Z","title":"Plug-and-play diffusion features for text-driven image-to-image translation","venue":null,"work_id":"a4f62cb1-901d-4ce5-bfef-bd380cded8fd","year":1921},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.725878Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:747e49a09391afe3bea01bc89d837904865a0d30731bcb3912945895b54663f7","observation_id":"a3d20650-6897-48ac-a165-7aab62cc6785","resolution":{"observed_at":"2026-08-10T19:36:22.167316Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:21.731329Z","title":"Ex- ploring clip for assessing the look and feel of images","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.731329Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:9c417575b27ede1d20ac8c74cfc110e35fa7581bb9b7de35d5e61912d253f3df","observation_id":"7eb4a0cc-98e8-4931-9b58-c0b26ab8bf7d","resolution":{"observed_at":"2026-08-10T19:36:21.731329Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.135148Z","title":"Image quality assessment: from error visibility to structural similarity","venue":null,"work_id":"ee46475d-6b7b-47f4-9273-cf6f7162e4fc","year":2004},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.736520Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:f2ca75b78bdd6eb8c4b5bcb97f7c7116d25ba2f2ccb9405d075ea48a8ad91384","observation_id":"482f1358-acc7-452a-bcaa-554d90976311","resolution":{"observed_at":"2026-08-10T19:36:22.140283Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.09341","last_updated":"2023-09-25T08:19:23Z","snapshot_observed_at":"2026-08-11T00:30:13.593181Z","submitted_at":"2023-06-15T17:59:31Z","title":"Human Preference Score v2: A Solid Benchmark for Evaluating Human Preferences of Text-to-Image Synthesis","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.09341","snapshot_observed_at":"2026-08-10T19:36:21.741256Z","title":"Human preference score v2: A solid benchmark for evaluating human preferences of text-to-image synthesis","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.741256Z"},"links":{"cited_paper":"/paper/2306.09341","citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:19572d96bc965f19292d400613d8e43fd3bda67eb491a225827379b1fb818bea","observation_id":"d5371387-204c-48f5-9b20-b71d5a430727","resolution":{"observed_at":"2026-08-10T19:36:21.741256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.118432Z","title":"Human preference score: Better aligning text- to-image models with human preference","venue":null,"work_id":"09af78bf-92a2-420b-b662-035f1207ec66","year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.746869Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:fc214e98c5d77baf8b071d8fdf2e4401e27d8354bc11a7bfe3816f2f4c74131d","observation_id":"eab88633-98f6-4330-9859-8c1225c2b654","resolution":{"observed_at":"2026-08-10T19:36:22.123688Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.100777Z","title":"Imagere- ward: Learning and evaluating human preferences for text- to-image generation, 2023","venue":null,"work_id":"7915017b-e863-4986-bf9a-a1293b24cd6b","year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.752019Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:095dbaf0fbd09f3db4ac271b19db3ca1f767f9bd99138fbb4f6f30e350bb4806","observation_id":"3538caf0-39c0-4a0a-b215-f5e812b1c2e1","resolution":{"observed_at":"2026-08-10T19:36:22.106402Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.04965","last_updated":"2023-12-07T18:58:27Z","snapshot_observed_at":"2026-07-06T16:58:47.790985Z","submitted_at":"2023-12-07T18:58:27Z","title":"Inversion-Free Image Editing with Natural Language","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.04965","snapshot_observed_at":"2026-08-10T19:36:21.757097Z","title":"Inversion-free image editing with natural language","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.757097Z"},"links":{"cited_paper":"/paper/2312.04965","citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:91d0be05cb070f7c7ba30bef440fbdb1de4710ba8a671e4c5a181022311d1247","observation_id":"bbcde75c-40f1-42ba-a3ec-daace9c75c35","resolution":{"observed_at":"2026-08-10T19:36:21.757097Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.084398Z","title":"Magicbrush: A manually annotated dataset for instruction- guided image editing","venue":null,"work_id":"828a4ecf-064b-4683-bb35-f20dd872dd75","year":2024},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.762300Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:10d2693c28c257817752046a6613928f02401975233da31c75a56fac524a71e8","observation_id":"8b78195b-df6f-45d3-a220-f970ad20feae","resolution":{"observed_at":"2026-08-10T19:36:22.089582Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.066960Z","title":"A comprehensive evaluation of full reference image quality as- sessment algorithms","venue":null,"work_id":"f34e1c7c-4c28-43f6-9c0b-bb9a932cab50","year":2012},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.767134Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:b8a415be74b1eaecb8f19dc10cde047637c17d15767225a2f24acff5ac540a43","observation_id":"7b311c75-7141-4899-b99d-02ba7d0f9fc9","resolution":{"observed_at":"2026-08-10T19:36:22.072681Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.049082Z","title":"The unreasonable effectiveness of deep features as a perceptual metric","venue":null,"work_id":"4858ce64-ef65-448a-adb1-09c8316fd62a","year":2018},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.772049Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:f805dd7f44f2b9c544dfb2139c8e7bb988d2a232d19519a78071427ac9f10506","observation_id":"f3c4e0f1-852f-4a0b-a19b-4cf2b0200726","resolution":{"observed_at":"2026-08-10T19:36:22.054492Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.032265Z","title":"Blind image quality assessment using a deep bilinear convolutional neural network","venue":null,"work_id":"fd6e76e8-c02a-4727-a3e9-93f3b0c6a386","year":2018},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.776759Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:83075fcda50013e2208c30f46ca84219695fe5802535ffa744ddba3cf1e9474c","observation_id":"c89e77c8-f7cb-4add-abeb-03cb2620c113","resolution":{"observed_at":"2026-08-10T19:36:22.037653Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:22.014904Z","title":"Blind image quality assessment via vision- language correspondence: A multitask learning perspective","venue":null,"work_id":"cacddba6-fae0-4a4e-a9b9-449dc90993f9","year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.781577Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:636189048060ef6667e1bc7afd68c1d081eb8c508485abc9d57a1d6ed7356c90","observation_id":"0743fe08-9b87-4615-b588-1649bb01eb09","resolution":{"observed_at":"2026-08-10T19:36:22.020914Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:21.997706Z","title":"Sine: Single image editing with text- to-image diffusion models","venue":null,"work_id":"d8e0cea7-e016-4193-8b77-cead206a5b74","year":2023},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.786976Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:2d6818c35d8bf435f3156d2de4e7a4f68c632bacc015adf369a4970301e1c531","observation_id":"94f42c52-be8c-4cfd-8b3b-016e79814100","resolution":{"observed_at":"2026-08-10T19:36:22.002868Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T19:36:21.981534Z","title":"Scene parsing through ade20k dataset","venue":null,"work_id":"331f85aa-722d-4fa8-b4ce-69095b605c11","year":2017},"citing_paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-10T19:36:21.792561Z"},"links":{"citing_paper":"/paper/2501.09927"},"observation_digest":"sha256:ed1b187966a99f4a6d6f1d52ec5802260223b31f08f754dc126223160d6ae0ea","observation_id":"e5eb6ef7-be34-4383-9f5f-9f3e7fff7d18","resolution":{"observed_at":"2026-08-10T19:36:21.986592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2501.09927","last_updated":"2025-01-17T02:47:25Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-11T22:10:47.897345Z","submitted_at":"2025-01-17T02:47:25Z","title":"IE-Bench: Advancing the Measurement of Text-Driven Image Editing for Human Perception Alignment"},"reference_resolution":{"displayed":44,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":13,"verified_exact":1,"verified_fuzzy":30},"total_outbound_references":44},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 44 of 44 outbound references and 6 inbound Pith citation observations for arXiv:2501.09927."}