{"as_of":"2026-08-09T16:32:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:97636b49728ab179e6705f9b0cb3eee1e034d1882fb0b4fe5357f607b6e85d17","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":15,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":15,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":15,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":15,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:29:14.789184Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T13:09:50.423322Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":"2406.07546","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-07-04T13:09:50.423322Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?","venue":null,"work_id":"4e3b9283-2aa2-4823-b228-dc40ffd87114","year":2024},"citing_paper":{"arxiv_id":"2410.05363","last_updated":"2024-10-07T17:56:04Z","snapshot_observed_at":"2026-08-08T09:48:34.424815Z","submitted_at":"2024-10-07T17:56:04Z","title":"Towards World Simulator: Crafting Physical Commonsense-Based Benchmark for Video Generation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-18T14:39:59.870039Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2410.05363"},"observation_digest":"sha256:227af88ab13b0a361f6c2d33efe366ded0f797b72a1126a7ded41c663582af72","observation_id":"54752231-047b-4069-be2e-242b0b1df1de","resolution":{"observed_at":"2026-05-18T14:40:00.094694Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":"2406.07546","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-07-04T13:09:50.423322Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?","venue":null,"work_id":"4e3b9283-2aa2-4823-b228-dc40ffd87114","year":2024},"citing_paper":{"arxiv_id":"2412.04300","last_updated":"2025-06-04T03:29:18Z","snapshot_observed_at":"2026-08-08T07:13:29.035650Z","submitted_at":"2024-12-05T16:21:01Z","title":"T2I-FactualBench: Benchmarking the Factuality of Text-to-Image Models with Knowledge-Intensive Concepts","version":3},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-23T08:00:12.781392Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2412.04300"},"observation_digest":"sha256:d06e0ecfb82a9d959c1446bf656c8149717d97bd3ba44662643696d90b8cbe8f","observation_id":"afdcc556-0802-4452-b94d-39fe4267e5f9","resolution":{"observed_at":"2026-05-23T08:02:43.281911Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":"2406.07546","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-07-04T13:09:50.423322Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?","venue":null,"work_id":"4e3b9283-2aa2-4823-b228-dc40ffd87114","year":2024},"citing_paper":{"arxiv_id":"2503.07265","last_updated":"2026-06-02T17:11:50Z","snapshot_observed_at":"2026-08-07T17:17:00.060047Z","submitted_at":"2025-03-10T12:47:53Z","title":"WISE: A World Knowledge-Informed Semantic Evaluation for Text-to-Image Generation","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-15T16:24:27.407376Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2503.07265"},"observation_digest":"sha256:0c5862b9261689a84091a3a89cdd1a60d32d782e5c594f00fb21e358137feda9","observation_id":"adf6ea4b-8d45-4b8b-85df-de20ad9611e0","resolution":{"observed_at":"2026-05-15T16:24:27.579185Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T14:29:14.789184Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense? arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.18730","last_updated":"2025-05-24T14:56:09Z","snapshot_observed_at":"2026-08-07T17:06:29.055596Z","submitted_at":"2025-05-24T14:56:09Z","title":"Align Beyond Prompts: Evaluating World Knowledge Alignment in Text-to-Image Generation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T14:29:14.789184Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2505.18730"},"observation_digest":"sha256:06110c1021ed1b95bdd96361f8e044e538d77a099e707a3604905ae6fca87311","observation_id":"cc46d3a9-9e75-4f0f-83e2-6b198eacd0d1","resolution":{"observed_at":"2026-08-07T14:29:14.789184Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T14:17:32.854514Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense? arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19415","last_updated":"2025-05-27T20:10:09Z","snapshot_observed_at":"2026-08-07T14:12:18.054397Z","submitted_at":"2025-05-26T02:07:24Z","title":"MMIG-Bench: Towards Comprehensive and Explainable Evaluation of Multi-Modal Image Generation Models","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T14:17:32.854514Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2505.19415"},"observation_digest":"sha256:210a827d2ead899f4869b03caee0be96325063026a649b900ed78f4150048717","observation_id":"9b784e3f-2441-487a-a237-cf01ea15bdca","resolution":{"observed_at":"2026-08-07T14:17:32.854514Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T12:48:50.711956Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.23493","last_updated":"2025-05-29T14:43:46Z","snapshot_observed_at":"2026-08-07T20:36:09.670028Z","submitted_at":"2025-05-29T14:43:46Z","title":"R2I-Bench: Benchmarking Reasoning-Driven Text-to-Image Generation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T12:48:50.711956Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2505.23493"},"observation_digest":"sha256:fc982585a8df95d3e59327d3e518cf2dd1cc9f39e0541adaf14e479801d22b79","observation_id":"2a1e1bb1-ba1a-4820-90c6-e681d26cbd59","resolution":{"observed_at":"2026-08-07T12:48:50.711956Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T12:21:21.501770Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24870","last_updated":"2025-06-06T14:51:40Z","snapshot_observed_at":"2026-08-09T05:59:14.538896Z","submitted_at":"2025-05-30T17:59:26Z","title":"GenSpace: Benchmarking Spatially-Aware Image Generation","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T12:21:21.501770Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2505.24870"},"observation_digest":"sha256:ce613208b0012141caf0835754a3a7d9786009f6d175b4bd4a6c5cf5f782de79","observation_id":"5042c595-d8b2-4957-8d6e-e989a06c2c7b","resolution":{"observed_at":"2026-08-07T12:21:21.501770Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T05:25:44.095229Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense? arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07977","last_updated":"2025-06-26T15:47:09Z","snapshot_observed_at":"2026-08-09T14:38:22.529752Z","submitted_at":"2025-06-09T17:50:21Z","title":"OneIG-Bench: Omni-dimensional Nuanced Evaluation for Image Generation","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T05:25:44.095229Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2506.07977"},"observation_digest":"sha256:d47ef1db75d8bf327d558062726a73621eca471bd0b2951d358d51c12f10e56b","observation_id":"16ad961e-f0ce-402b-b9f1-d84bc5b2ff81","resolution":{"observed_at":"2026-08-07T05:25:44.095229Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-06T20:32:59.658017Z","title":"Commonsense-t2i challenge: Can text-to- image generation models understand commonsense? arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.02664","last_updated":"2025-07-07T08:00:38Z","snapshot_observed_at":"2026-08-09T10:11:16.716146Z","submitted_at":"2025-07-03T14:26:31Z","title":"AIGI-Holmes: Towards Explainable and Generalizable AI-Generated Image Detection via Multimodal Large Language Models","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T20:32:59.658017Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2507.02664"},"observation_digest":"sha256:372c9b6e8a9afda5b3c7e597c1e60f6df6e572e2ab5128a44138410b3ab801f2","observation_id":"ba2461e4-cd1c-4b4a-9292-c133b1004b66","resolution":{"observed_at":"2026-08-06T20:32:59.658017Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-04T18:48:03.655845Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.09680","last_updated":"2025-09-11T17:59:59Z","snapshot_observed_at":"2026-08-04T18:48:02.197558Z","submitted_at":"2025-09-11T17:59:59Z","title":"FLUX-Reason-6M & PRISM-Bench: A Million-Scale Text-to-Image Reasoning Dataset and Comprehensive Benchmark","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-04T18:48:03.655845Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2509.09680"},"observation_digest":"sha256:6a4d4da91c7345e2c80b61cf925367b71445d0ce1e05759d6b8af7eb1b92f2eb","observation_id":"66c32426-e674-488e-8e1d-aaad6cc0f225","resolution":{"observed_at":"2026-08-04T18:48:03.655845Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":"2406.07546","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-07-04T13:09:50.423322Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?","venue":null,"work_id":"4e3b9283-2aa2-4823-b228-dc40ffd87114","year":2024},"citing_paper":{"arxiv_id":"2606.26738","last_updated":"2026-07-15T11:15:44Z","snapshot_observed_at":"2026-08-02T10:10:07.945306Z","submitted_at":"2026-06-25T08:20:34Z","title":"Do Image Editing Models Understand Lighting?","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-26T05:29:42.024146Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2606.26738"},"observation_digest":"sha256:576cc178c4240606e331df85ddcef1c6da443b11f624abd26b91b9c195c52834","observation_id":"eb48bd77-e076-41ca-a7fe-926222af3742","resolution":{"observed_at":"2026-07-04T13:09:50.425449Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-02T10:10:09.175457Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.26738","last_updated":"2026-07-15T11:15:44Z","snapshot_observed_at":"2026-08-02T10:10:07.945306Z","submitted_at":"2026-06-25T08:20:34Z","title":"Do Image Editing Models Understand Lighting?","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-02T10:10:09.175457Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2606.26738"},"observation_digest":"sha256:112e05e255851f80ab969ccefb69ad3391a8efe283c78edc3a8e351cb4dae313","observation_id":"9bee8750-2446-469c-90c2-3cc1218ad00e","resolution":{"observed_at":"2026-08-02T10:10:09.175457Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":"2406.07546","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-07-04T13:09:50.423322Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?","venue":null,"work_id":"4e3b9283-2aa2-4823-b228-dc40ffd87114","year":2024},"citing_paper":{"arxiv_id":"2606.30262","last_updated":"2026-06-29T13:09:39Z","snapshot_observed_at":"2026-08-06T12:26:22.536829Z","submitted_at":"2026-06-29T13:09:39Z","title":"Intermediate Text Representation Guided Text-to-Image Generation for Enhancing One-and-Only Alignment","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-30T06:04:54.816368Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2606.30262"},"observation_digest":"sha256:dcbdded4829ed0ce7e81059c8ceac09a2bf5f33aedb0b6f6c9d6c532b1554450","observation_id":"611ce9ff-7846-466d-a374-78755633f1e1","resolution":{"observed_at":"2026-06-30T08:24:27.106668Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-01T01:54:37.826437Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.25641","last_updated":"2026-07-28T12:27:25Z","snapshot_observed_at":"2026-08-09T04:59:41.655859Z","submitted_at":"2026-07-28T12:27:25Z","title":"OmniPhys: Knowledge-Graph-Driven Benchmarking and Collective Optimization for Physical Commonsense in Text-to-Image Generation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-01T01:54:37.826437Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2607.25641"},"observation_digest":"sha256:5f0aed3b40fefa39fc5089344bb945674baf5c0579032806724b1f07bec5f626","observation_id":"64f394f1-01c4-4210-9b78-a5196d16eef7","resolution":{"observed_at":"2026-08-01T01:54:37.826437Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T00:15:38.364743Z","title":"Commonsense-t2i challenge: Can text-to- image generation models understand commonsense?, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.04436","last_updated":"2026-08-05T04:26:19Z","snapshot_observed_at":"2026-08-09T05:49:30.638721Z","submitted_at":"2026-08-05T04:26:19Z","title":"ToolArtist: Tool-Using Unified Multimodal Models for Agentic Image Generation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T00:15:38.364743Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2608.04436"},"observation_digest":"sha256:8aab5bb495f6f3f9fcac5a523a473bdb91a662b362c21f7925eb1ed22a8a9ddd","observation_id":"b8e6eff9-412a-43a9-9940-376db9b11c24","resolution":{"observed_at":"2026-08-07T00:15:38.364743Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2406.07546/citation-record","integrity":"/paper/2406.07546/integrity","json":"/paper/2406.07546/citation-record.json","paper":"/paper/2406.07546"},"outbound":[],"paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-09T08:24:08.295414Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 15 inbound Pith citation observations for arXiv:2406.07546."}