{"as_of":"2026-08-18T18:56:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:23b81e68478336a7a271e13848cb0ddf7b97a2435c97cbf5a3525b3ca9b169e1","coverage":[{"denominator":39,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":39,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T20:56:57.572238Z","state":"measured"},{"denominator":47,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":47,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":8,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":8,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T02:13:11.616524Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-09T21:16:34.356874Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"cited_work":{"arxiv_id":"2505.11493","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.11493","snapshot_observed_at":"2026-07-09T21:16:34.356874Z","title":"Gie-bench: Towards grounded evaluation for text-guided image editing.arXiv preprint","venue":"cs.CV","work_id":"489714b6-180f-4824-aef6-be392bf3adda","year":2025},"citing_paper":{"arxiv_id":"2602.23622","last_updated":"2026-05-19T03:38:52Z","snapshot_observed_at":"2026-08-15T15:16:14.894154Z","submitted_at":"2026-02-27T02:59:34Z","title":"DLEBench: Evaluating Small-scale Object Editing Ability for Instruction-based Image Editing Model","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-21T12:04:08.487678Z"},"links":{"cited_paper":"/paper/2505.11493","citing_paper":"/paper/2602.23622"},"observation_digest":"sha256:b6724c2a42a411daee13f4023178d57976595c4a27e65320e5e0fdc807d75aa3","observation_id":"0a1e7150-bde0-49f0-92bd-4beac5722fe4","resolution":{"observed_at":"2026-05-21T12:05:04.951868Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"cited_work":{"arxiv_id":"2505.11493","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.11493","snapshot_observed_at":"2026-07-09T21:16:34.356874Z","title":"Gie-bench: Towards grounded evaluation for text-guided image editing.arXiv preprint","venue":"cs.CV","work_id":"489714b6-180f-4824-aef6-be392bf3adda","year":2025},"citing_paper":{"arxiv_id":"2605.08163","last_updated":"2026-05-04T16:21:25Z","snapshot_observed_at":"2026-08-15T03:16:28.209720Z","submitted_at":"2026-05-04T16:21:25Z","title":"MULTITEXTEDIT: Benchmarking Cross-Lingual Degradation in Text-in-Image Editing","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-05-12T01:11:35.480399Z"},"links":{"cited_paper":"/paper/2505.11493","citing_paper":"/paper/2605.08163"},"observation_digest":"sha256:e80ee5faffd6fd480946af0aeb02208f74275a149d55890daefe05c726ceb080","observation_id":"4f0abb18-d53c-4ad5-b0ee-428d954cb699","resolution":{"observed_at":"2026-05-12T08:26:23.441082Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"cited_work":{"arxiv_id":"2505.11493","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.11493","snapshot_observed_at":"2026-07-09T21:16:34.356874Z","title":"Gie-bench: Towards grounded evaluation for text-guided image editing.arXiv preprint","venue":"cs.CV","work_id":"489714b6-180f-4824-aef6-be392bf3adda","year":2025},"citing_paper":{"arxiv_id":"2605.14842","last_updated":"2026-05-14T13:46:24Z","snapshot_observed_at":"2026-08-12T22:57:10.943424Z","submitted_at":"2026-05-14T13:46:24Z","title":"Editor's Choice: Evaluating Abstract Intent in Image Editing through Atomic Entity Analysis","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-06-30T21:41:10.852265Z"},"links":{"cited_paper":"/paper/2505.11493","citing_paper":"/paper/2605.14842"},"observation_digest":"sha256:1d04c9847daa278ff6187c4ee39788d65a587b6587c87a5ca8cb583b0726b2e0","observation_id":"847f0ef5-9046-4ddd-afc7-dac030866610","resolution":{"observed_at":"2026-07-01T14:25:45.982484Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"cited_work":{"arxiv_id":"2505.11493","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.11493","snapshot_observed_at":"2026-07-09T21:16:34.356874Z","title":"Gie-bench: Towards grounded evaluation for text-guided image editing.arXiv preprint","venue":"cs.CV","work_id":"489714b6-180f-4824-aef6-be392bf3adda","year":2025},"citing_paper":{"arxiv_id":"2606.01804","last_updated":"2026-06-03T03:45:51Z","snapshot_observed_at":"2026-08-15T02:27:40.157038Z","submitted_at":"2026-06-01T07:21:02Z","title":"SpeechEditBench: A Bilingual Multi-Attribute Benchmark for Instruction-Guided Speech Editing","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-06-28T12:54:20.815371Z"},"links":{"cited_paper":"/paper/2505.11493","citing_paper":"/paper/2606.01804"},"observation_digest":"sha256:d63e1b9569cbf8987a5f3eba3d7d9a7283b2b42d044309efd4b0284046bb03f4","observation_id":"b4d2bef4-f75e-4b7f-973a-790357831bd4","resolution":{"observed_at":"2026-07-02T01:06:23.932282Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"cited_work":{"arxiv_id":"2505.11493","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.11493","snapshot_observed_at":"2026-07-09T21:16:34.356874Z","title":"Gie-bench: Towards grounded evaluation for text-guided image editing.arXiv preprint","venue":"cs.CV","work_id":"489714b6-180f-4824-aef6-be392bf3adda","year":2025},"citing_paper":{"arxiv_id":"2606.01985","last_updated":"2026-06-01T09:46:10Z","snapshot_observed_at":"2026-08-13T08:27:20.774408Z","submitted_at":"2026-06-01T09:46:10Z","title":"MT-EditFlow: Reinforcement Learning for Multi-Turn Image Editing with Flow Matching","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-06-28T15:07:29.897089Z"},"links":{"cited_paper":"/paper/2505.11493","citing_paper":"/paper/2606.01985"},"observation_digest":"sha256:9163577eab01e1e6ac7ef6f4372fc7d690e6fe796069c0f53db380c2eace712f","observation_id":"77a83c56-d90d-4dc5-9c64-6974e62a771e","resolution":{"observed_at":"2026-07-01T22:46:18.509247Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"cited_work":{"arxiv_id":"2505.11493","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.11493","snapshot_observed_at":"2026-07-09T21:16:34.356874Z","title":"Gie-bench: Towards grounded evaluation for text-guided image editing.arXiv preprint","venue":"cs.CV","work_id":"489714b6-180f-4824-aef6-be392bf3adda","year":2025},"citing_paper":{"arxiv_id":"2606.13303","last_updated":"2026-07-31T12:46:43Z","snapshot_observed_at":"2026-08-15T09:09:02.230144Z","submitted_at":"2026-06-11T12:58:48Z","title":"DuET: Dual Expert Trajectories for Diffusion Image Editing","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-27T06:58:39.254427Z"},"links":{"cited_paper":"/paper/2505.11493","citing_paper":"/paper/2606.13303"},"observation_digest":"sha256:39540b9068fab7757f5204f9b3c74516a317b82f6df0444ff1155c9f25704f93","observation_id":"a395a2c7-1ebb-42ab-8525-4400c3b08dad","resolution":{"observed_at":"2026-07-03T14:38:29.287541Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.11493","snapshot_observed_at":"2026-08-03T02:13:11.616524Z","title":"GIE-Bench: Towards grounded evaluation for text-guided image editing.arXiv preprint arXiv:2505.11493,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.13303","last_updated":"2026-07-31T12:46:43Z","snapshot_observed_at":"2026-08-15T09:09:02.230144Z","submitted_at":"2026-06-11T12:58:48Z","title":"DuET: Dual Expert Trajectories for Diffusion Image Editing","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-03T02:13:11.616524Z"},"links":{"cited_paper":"/paper/2505.11493","citing_paper":"/paper/2606.13303"},"observation_digest":"sha256:f591309d05148c811d06945c963a609c2fa2bfd05f5f09b502b66c3ea7af6cf5","observation_id":"00b27926-85f0-4366-b1e9-7066df7cd3a4","resolution":{"observed_at":"2026-08-03T02:13:11.616524Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"cited_work":{"arxiv_id":"2505.11493","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.11493","snapshot_observed_at":"2026-07-09T21:16:34.356874Z","title":"Gie-bench: Towards grounded evaluation for text-guided image editing.arXiv preprint","venue":"cs.CV","work_id":"489714b6-180f-4824-aef6-be392bf3adda","year":2025},"citing_paper":{"arxiv_id":"2607.07051","last_updated":"2026-07-08T06:26:28Z","snapshot_observed_at":"2026-08-13T05:01:32.634587Z","submitted_at":"2026-07-08T06:26:28Z","title":"Making Implicit Preservation Intent Explicit in Conversational Image Editing","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-07-09T21:07:55.368221Z"},"links":{"cited_paper":"/paper/2505.11493","citing_paper":"/paper/2607.07051"},"observation_digest":"sha256:d96c08b4b370afae061763105aa8d481e60a5a6064053ecb456ba266ff8cf635","observation_id":"b98832a1-4a4a-4dae-a501-32b7d9362ab8","resolution":{"observed_at":"2026-07-09T21:16:34.358163Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.11493/citation-record","integrity":"/paper/2505.11493/integrity","json":"/paper/2505.11493/citation-record.json","paper":"/paper/2505.11493"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:58.085103Z","title":"Gpt image 1: State-of-the-art image generation model","venue":null,"work_id":"fe72ca1f-5f0f-4eae-9596-d82ee11c06d5","year":2025},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.399461Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:f82ca19fb051141239b47e4d91b8564b3472f5109326d95cee69116b7e49a967","observation_id":"81bec61a-efb1-4fdf-8924-39a59b4639d4","resolution":{"observed_at":"2026-08-15T20:56:58.088081Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.404174Z","title":"Learning transferable visual models from natural language supervision, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.404174Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:b0653fb498d7177905b34aba32739aad7aa6a9a79e719dfa2f95b796a4180810","observation_id":"8e51f168-1b17-4138-86eb-2d1067fbf0d1","resolution":{"observed_at":"2026-08-15T20:56:57.404174Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:58.068217Z","title":"I2ebench: A comprehensive benchmark for instruction-based image editing","venue":null,"work_id":"18df3334-5290-4ff8-bee3-c1332aafa4fb","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.409226Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:39d62feb6080365955e86e04f8d37b6cb91079a8f80f24467c828a27c5fdc911","observation_id":"ea3ce325-82f4-4fff-99a8-fffbc919de2e","resolution":{"observed_at":"2026-08-15T20:56:58.071704Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:58.058129Z","title":"Guiding Instruction-based Image Editing via Multimodal Large Language Models","venue":null,"work_id":"09899dd5-0034-484d-ac3f-8a4f58463086","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.414959Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:13240180f9625aeaad2b2167c6273b7d58bc547fc758c738c01bcd468792d973","observation_id":"0e1d0b3c-9bae-4861-8b7a-305deeb0d410","resolution":{"observed_at":"2026-08-15T20:56:58.061483Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.420481Z","title":"Omnigen: Unified image generation, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.420481Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:764af14bc68d01e24ec6b80486f41d5ba44e9d2f337cebf389ad7e17bc04b237","observation_id":"edffa18e-3844-49cc-a703-4131f0cb9ffd","resolution":{"observed_at":"2026-08-15T20:56:57.420481Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:58.039164Z","title":"Ie-bench: Advancing the measurement of text-driven image editing for human perception alignment, 2025","venue":null,"work_id":"eb2d3de5-f2af-4adf-ad75-e3aca7fbb6ca","year":2025},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.424669Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:8b128797c87da9672feb2b8690c1e2ce6cc562afd35198308d1dbfc29408a551","observation_id":"5b49e253-966c-4617-9d74-4ebd555264bc","resolution":{"observed_at":"2026-08-15T20:56:58.042853Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:58.028116Z","title":"Imagen editor and editbench: Advancing and evaluating text-guided image inpainting","venue":null,"work_id":"a5388c08-40b2-4c5b-beac-93fab3e4f061","year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.429264Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:c36af428e7072e4cb132b339a1e9745d8e3bc171034641c99e283f8518f5297f","observation_id":"292d0fe9-74b7-4d78-a526-77ad969ad749","resolution":{"observed_at":"2026-08-15T20:56:58.032062Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:58.016095Z","title":"Diffedit: Diffusion- based semantic image editing with mask guidance, 2022","venue":null,"work_id":"b45debd2-b226-4da4-ae8d-5da99241bb9c","year":2022},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.432964Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:d74995075cae38e14f9b2737763d09e7cf4159ec6b5506c1ebc39fc4241ee0dd","observation_id":"1b3476cb-77d4-4706-86e5-c05dd4b17285","resolution":{"observed_at":"2026-08-15T20:56:58.020312Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.436584Z","title":"Prompt-to-prompt image editing with cross attention control, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.436584Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:d47358a136f1b97acdb88c85753bef10f597c94263dd4e9ca42cf56c212bec7d","observation_id":"04edc1e5-bf60-4cdd-ba47-9a59a2916431","resolution":{"observed_at":"2026-08-15T20:56:57.436584Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.996407Z","title":"Imagic: Text-based real image editing with diffusion models","venue":null,"work_id":"e56f79b1-705b-457a-a260-c5eeba0f2bd6","year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.440652Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:3d10a1bfc8c30948447487e65f6c570bb1146adecdef89313372a0ac79046861","observation_id":"fe8a8261-08d1-4f0f-a472-82d14a885f78","resolution":{"observed_at":"2026-08-15T20:56:57.999987Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.986013Z","title":null,"venue":null,"work_id":"33a5c576-39a8-44f5-907c-16cc692d5360","year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.444888Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:21cccc0330a87dd8d685fab5baea60575daf1752dfd98fb28dfd9303c28cd005","observation_id":"38d8250c-d5bf-4362-a327-1044268d6ab2","resolution":{"observed_at":"2026-08-15T20:56:57.988997Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.975576Z","title":"Instructdiffusion: A generalist modeling interface for vision tasks, 2023","venue":null,"work_id":"06170f5c-9569-4e90-9b43-68e075f16037","year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.448967Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:41866a698f04a4436674f2dc812316e9a46c07acaa5c62caeb0f0fda99929ba0","observation_id":"90d2569a-f3df-4153-9c53-3e8c297fba24","resolution":{"observed_at":"2026-08-15T20:56:57.979055Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.964131Z","title":"Univg: A generalist diffusion model for unified image generation and editing, 2025","venue":null,"work_id":"70d95199-1a45-4b08-9f7f-5ff0d1f77abd","year":2025},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.453122Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:7c26f2d25970ffb4745f2dad3faa06c40b64aa08075c8ec18182fa7c654df228","observation_id":"22cb097a-c808-489a-bc9e-9afe15a04dd9","resolution":{"observed_at":"2026-08-15T20:56:57.968824Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.943047Z","title":"Le, Tuan Pham, Sangho Lee, Christopher Clark, Aniruddha Kembhavi, Stephan Mandt, Ranjay Krishna, and Jiasen Lu","venue":null,"work_id":"aaa75bca-9e26-4cd7-abd8-bceb11ea3845","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.457133Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:f6fc8f1d68a38196739fc9e97090794e71f32bcb050556c84bc58935e26f2189","observation_id":"6a63aadc-a642-45b6-b163-93333a47992c","resolution":{"observed_at":"2026-08-15T20:56:57.951692Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.460898Z","title":"Magicbrush: A manually annotated dataset for instruction-guided image editing","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.460898Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:fee9dfc418944ae49b97adae1a32b342a551b0ae33bb1185f3bd100500988fd4","observation_id":"5c2c23ee-63a1-499a-8eaa-3e40ce7f39e4","resolution":{"observed_at":"2026-08-15T20:56:57.460898Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.465188Z","title":"Smartedit: Exploring complex instruction-based image editing with multimodal large language models, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.465188Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:e36c6a64bfad5019ffc22d1add79623832f3b532ee9ba8588f2257e7fc4ae6d8","observation_id":"904db2a7-cc27-472a-aad3-8083b2d385b5","resolution":{"observed_at":"2026-08-15T20:56:57.465188Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.913601Z","title":"Instructany2pix: Flexible visual editing via multimodal instruction following, 2024","venue":null,"work_id":"54cc1814-85d9-43e0-901d-9085fa17ee1a","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.469801Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:a00d7dc972bc9b92fc3691d98ac06fb979a6aec46d9f3ef0c542b1a996470bdb","observation_id":"18801c2b-ca81-4644-9be6-fdde186b864d","resolution":{"observed_at":"2026-08-15T20:56:57.919556Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.896174Z","title":"Text-driven image editing via learnable regions, 2024","venue":null,"work_id":"f7c2d1ff-f120-495f-baca-8d616130f029","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.474014Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:7f089421231d50e6e4b3cc7e7a5a8876efe4caaff03a34803bad5c2a8e1479c6","observation_id":"a15a49bd-243a-47b1-8f1c-f3a82bccf851","resolution":{"observed_at":"2026-08-15T20:56:57.903014Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.881910Z","title":"Instructgie: Towards generalizable image editing, 2024","venue":null,"work_id":"6d350bc5-c507-41f7-b4be-dfcfec906e7c","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.478128Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:4753e098810856cb0826998ac68ccfaeeb9ff3840d29f783db56efc198589e07","observation_id":"b616a069-7e0d-442e-bbf9-3172312a60df","resolution":{"observed_at":"2026-08-15T20:56:57.887144Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.864345Z","title":"Stylebooth: Image style editing with multimodal instruction, 2024","venue":null,"work_id":"80cf4d7f-22ea-4856-aee3-6856e8742740","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.482939Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:4e91d9647a6bbfd4b51db73d898599ffc8e76dac0c612a830aebe0fe4c71aa0e","observation_id":"98fb505f-7302-4cb6-9f27-cbb0631fc769","resolution":{"observed_at":"2026-08-15T20:56:57.871091Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.848424Z","title":"Zone: Zero-shot instruction-guided local editing, 2024","venue":null,"work_id":"0e50a559-0f67-4fdf-993f-f87159fd97ca","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.487507Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:9845095aba79c468bc2224ec42b53bfa307988911266aa8f89c8ae920bc533a1","observation_id":"1b6415d6-1b57-4c04-b166-9ef783b8645b","resolution":{"observed_at":"2026-08-15T20:56:57.853875Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.830613Z","title":"Focus on your instruction: Fine-grained and multi-instruction image editing by attention modulation, 2023","venue":null,"work_id":"456ee38f-a175-445d-a342-cd2b673cdaad","year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.492060Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:f590345c799b6b6cc3d055b4b02a96691f1441f9afa4fbe91437759b69aa90e1","observation_id":"12df2a12-e338-456f-9d52-e04088a022e4","resolution":{"observed_at":"2026-08-15T20:56:57.837795Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.812787Z","title":"Leveraging llms for on-the-fly instruction guided image editing, 2024","venue":null,"work_id":"8bbfe944-c178-4168-abc6-5bdd7706d06b","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.496543Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:b4db6aaacc195d195be1a0d3d7f0f733cfd4bda456a0a7e82ea0b9b1a5022aab","observation_id":"1a2f6614-539b-4ff5-9461-8fcf708c3834","resolution":{"observed_at":"2026-08-15T20:56:57.819116Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.794360Z","title":"Instructbrush: Learning attention-based instruction optimization for image editing, 2024","venue":null,"work_id":"76c38f16-416c-4183-812b-21e8d42f07c3","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.501338Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:e145d52d985e5d6912a0881696b7b850c67ca9ffa2de0a342d63e2c3a442c351","observation_id":"06ea91ed-6fcb-42a4-bb39-20dc873fc552","resolution":{"observed_at":"2026-08-15T20:56:57.803672Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.774685Z","title":"Editval: Benchmarking diffusion based text-guided image editing methods, 2023","venue":null,"work_id":"44105d18-deb9-47b8-8300-934143800060","year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.505991Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:836ecedf5365eece95380ee0612795867cc7a4521def92948f18f65a3e24840d","observation_id":"04f41920-0d5d-44cc-9bf9-0b247d99acc6","resolution":{"observed_at":"2026-08-15T20:56:57.783459Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.510412Z","title":"Tifa: Accurate and interpretable text-to-image faithfulness evaluation with question answering, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.510412Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:9c805c37b68f69fea8376add6963ae13cd4af011dec7e58ef32069b5f9da55c8","observation_id":"52f08ae4-435e-40f7-b20a-db5225211d38","resolution":{"observed_at":"2026-08-15T20:56:57.510412Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.748671Z","title":"Davidsonian scene graph: Improving reliability in fine-grained evaluation for text-to-image generation, 2024","venue":null,"work_id":"64f64f25-c614-4ccb-845d-db0b414c6d96","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.515157Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:000fc5b0e0c067925bcc7bd0ad1da4071c6c9c3ee17304cef61956eb422ec77b","observation_id":"822cca3a-508d-4cb8-9db1-539e755e13d6","resolution":{"observed_at":"2026-08-15T20:56:57.756125Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.05499","last_updated":"2024-07-19T06:00:41Z","snapshot_observed_at":"2026-07-06T15:00:58.804337Z","submitted_at":"2023-03-09T18:52:16Z","title":"Grounding DINO: Marrying DINO with Grounded Pre-Training for Open-Set Object Detection","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.05499","snapshot_observed_at":"2026-08-15T20:56:57.519213Z","title":"Grounding dino: Marrying dino with grounded pre-training for open-set object detection.arXiv preprint arXiv:2303.05499, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.519213Z"},"links":{"cited_paper":"/paper/2303.05499","citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:761fad2213420e1f2710abfd1d25dde58839424992c7f2a2f277770fd5bb03f5","observation_id":"45eb1feb-c2ba-4ad0-94c7-016ea48a4e02","resolution":{"observed_at":"2026-08-15T20:56:57.519213Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.02643","last_updated":"2023-04-05T17:59:46Z","snapshot_observed_at":"2026-08-08T05:14:59.435033Z","submitted_at":"2023-04-05T17:59:46Z","title":"Segment Anything","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.02643","snapshot_observed_at":"2026-08-15T20:56:57.524097Z","title":"Berg, Wan-Yen Lo, Piotr Dollár, and Ross Girshick","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.524097Z"},"links":{"cited_paper":"/paper/2304.02643","citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:1b75b453bbb10726239b480a127f4016a6a82900494f64b87aa86dc0d35d55d8","observation_id":"296ff5d6-76ba-443e-af03-9c7d35151961","resolution":{"observed_at":"2026-08-15T20:56:57.524097Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.528711Z","title":"Distinctive image features from scale-invariant keypoints.International journal of computer vision, 60:91–110, 2004","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.528711Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:31d44b20ee37a56ebdcd2295b8eace6578469e33f83a5c89b347cf78483b58d2","observation_id":"0cac0d71-66ba-4202-bddb-96add496b5c7","resolution":{"observed_at":"2026-08-15T20:56:57.528711Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.723226Z","title":"Fast approximate nearest neighbors with automatic algorithm configuration.VISAPP (1), 2(331-340):2, 2009","venue":null,"work_id":"d96f49dc-0b20-4ee5-a7f8-d6dbafdec859","year":2009},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.533491Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:34be4a02569ff50ace1da466bb4a19dc95133fbc744353adfd4766e0dbda4858","observation_id":"538d6a64-e098-4005-a249-ffbce8635586","resolution":{"observed_at":"2026-08-15T20:56:57.730638Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.710363Z","title":"Introducing gemini 2.0: our new ai model for the agentic era.https://blog.google/ technology/google-deepmind/google-gemini-ai-update-december-2024/, 2024","venue":null,"work_id":"1d07281e-6cb2-4ae5-8f83-b1413ea1d8fd","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.537914Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:71e97d337b8c16f8e887df381c86a371ddb253aa1b9b4b774a682812a80c4b9e","observation_id":"3e623d27-6cf1-427b-82fc-7e8d6861db4f","resolution":{"observed_at":"2026-08-15T20:56:57.715194Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.542483Z","title":"Sara Mahdavi, Rapha Gontijo Lopes, Tim Salimans, Jonathan Ho, David J Fleet, and Mohammad Norouzi","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.542483Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:8766851c6a32a8029b0d611292aba9fcf6f3646a3cc4ebae0a3264d3cfdd1aee","observation_id":"11ebe122-e2f9-43fe-b337-5711e560e699","resolution":{"observed_at":"2026-08-15T20:56:57.542483Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.547394Z","title":"Scaling autoregressive models for content-rich text-to-image generation, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.547394Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:d6f8a837f1a437e95d9ec0cfecc3ffee689feec0f6e602845cc95b397430164b","observation_id":"6893f86d-0714-4444-ac8f-9943816a0cf9","resolution":{"observed_at":"2026-08-15T20:56:57.547394Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.683118Z","title":"Pick-a-pic: An open dataset of user preferences for text-to-image generation, 2023","venue":null,"work_id":"913455f9-a115-4256-b1cd-f0fecbdc5e9a","year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.552362Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:d14558dcfc6abf316d9588f66b7ff3b5c1d5d4820bc3e5b919000ce109f2001d","observation_id":"2ed0388a-4c80-49e4-89b7-761203ab8a4e","resolution":{"observed_at":"2026-08-15T20:56:57.687268Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.669351Z","title":"Evalmuse-40k: A reliable and fine-grained bench- mark with comprehensive human annotations for text-to-image generation model evaluation, 2024","venue":null,"work_id":"3381308e-0f0b-461d-b403-a7a4a8299fa5","year":2024},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.556859Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:ce06ca402c77789efc73ba30c5237f953c7a3976c95cf16a535235a6b221c279","observation_id":"813d41cb-f06e-4307-ac17-27c410811c57","resolution":{"observed_at":"2026-08-15T20:56:57.674666Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.656984Z","title":"Agiqa-3k: An open database for ai-generated image quality assessment, 2023","venue":null,"work_id":"3391601e-118d-4d09-a708-ab74896b9d8d","year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.562003Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:e48d90d3a1e454dd0298406ace61f7e8efa5a299cb5cebec1645a033e0ab683e","observation_id":"eafd5d98-a6d3-46db-ba60-593cbaf778ca","resolution":{"observed_at":"2026-08-15T20:56:57.661018Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.644008Z","title":"Lmm4lmm: Benchmarking and evaluating large-multimodal image generation with lmms, 2025","venue":null,"work_id":"458a3930-b9f2-456c-aa09-59de6947b3a3","year":2025},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.566533Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:2408cbb432725e90d96fbb3225e2ecff74adb69074cc4ecd73f3beede214665c","observation_id":"2e04d391-29b4-48fe-a388-b5c475e50fce","resolution":{"observed_at":"2026-08-15T20:56:57.648883Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:56:57.628051Z","title":"Color Change","venue":null,"work_id":"80901f14-c7c5-49ca-82d8-bf47c9731873","year":2023},"citing_paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-15T20:56:57.572238Z"},"links":{"citing_paper":"/paper/2505.11493"},"observation_digest":"sha256:5f8e1951c9f11a8afb77d675b488e5bbb4db387d6b257323644777e0247eed9c","observation_id":"e7ecedde-d109-4ecf-8d32-df6cc6d325ee","resolution":{"observed_at":"2026-08-15T20:56:57.635000Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.11493","last_updated":"2025-07-25T08:24:07Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-16T21:36:27.436550Z","submitted_at":"2025-05-16T17:55:54Z","title":"GIE-Bench: Towards Grounded Evaluation for Text-Guided Image Editing"},"reference_resolution":{"displayed":39,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":12,"verified_exact":0,"verified_fuzzy":26},"total_outbound_references":39},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 39 of 39 outbound references and 8 inbound Pith citation observations for arXiv:2505.11493."}