{"as_of":"2026-08-09T23:15:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:71b65cf35c4f54546c4b4ac0d94dc3375ccfd92f55e7fd703fef3d56c4574bf1","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":34,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":34,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":34,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":34,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:34:53.137002Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T16:39:58.175257Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2503.09567","last_updated":"2025-07-18T15:57:54Z","snapshot_observed_at":"2026-08-08T22:33:20.124926Z","submitted_at":"2025-03-12T17:35:03Z","title":"Towards Reasoning Era: A Survey of Long Chain-of-Thought for Reasoning Large Language Models","version":5},"reference_index":234,"source":"pdf_text","source_observed_at":"2026-05-12T08:40:40.910461Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2503.09567"},"observation_digest":"sha256:ab56df4ab4ae269d82994c25b8875b658dfa0363f740a286a0bb78ed9b476f1a","observation_id":"48f17002-9e99-4caf-acf4-ae4711e6c14c","resolution":{"observed_at":"2026-05-12T08:40:41.448367Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2503.12605","last_updated":"2025-03-23T13:47:43Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-16T18:39:13Z","title":"Multimodal Chain-of-Thought Reasoning: A Comprehensive Survey","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-15T17:18:52.996467Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2503.12605"},"observation_digest":"sha256:2b5923140a3a72cfdc40dd3b659a2ae319690a94333a1a2d1746cfd26b9f21a9","observation_id":"fba5499c-c290-4ca8-a6e6-37db20ab6514","resolution":{"observed_at":"2026-05-15T17:18:53.651274Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2505.07818","last_updated":"2025-08-28T17:19:45Z","snapshot_observed_at":"2026-07-31T14:51:03.625964Z","submitted_at":"2025-05-12T17:59:34Z","title":"DanceGRPO: Unleashing GRPO on Visual Generation","version":4},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-11T22:28:24.929046Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.07818"},"observation_digest":"sha256:1d29aad69da1bba875ae62a65a9ce9ff26af9718be771dd8807603fc3cd7d9bf","observation_id":"c6f9f267-9f74-4ea7-9ee7-83052f813b8b","resolution":{"observed_at":"2026-05-11T22:28:26.031397Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T15:34:53.137002Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step.arXiv:2501.13926, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.14682","last_updated":"2025-05-20T17:59:26Z","snapshot_observed_at":"2026-08-07T20:35:28.075030Z","submitted_at":"2025-05-20T17:59:26Z","title":"UniGen: Enhanced Training & Test-Time Strategies for Unified Multimodal Understanding and Generation","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T15:34:53.137002Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.14682"},"observation_digest":"sha256:d0ac461f969ba443ce38cca270eac056266b1f3914d56dbb28d18c210e102f44","observation_id":"1fc057b5-f3db-4e60-b51b-77bf28e69907","resolution":{"observed_at":"2026-08-07T15:34:53.137002Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T14:55:51.520817Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17017","last_updated":"2025-06-10T13:46:37Z","snapshot_observed_at":"2026-08-07T20:38:00.494377Z","submitted_at":"2025-05-22T17:59:49Z","title":"Delving into RL for Image Generation with CoT: A Study on DPO vs. GRPO","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T14:55:51.520817Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.17017"},"observation_digest":"sha256:f7bc0765124f9ba450b728ddc52a561a66bd7927cef0a41110212c677e1d9ae6","observation_id":"052310a6-09eb-4896-a7dd-72062cfa63af","resolution":{"observed_at":"2026-08-07T14:55:51.520817Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T14:49:03.001071Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17540","last_updated":"2025-05-23T06:44:26Z","snapshot_observed_at":"2026-08-07T20:36:17.159381Z","submitted_at":"2025-05-23T06:44:26Z","title":"RePrompt: Reasoning-Augmented Reprompting for Text-to-Image Generation via Reinforcement Learning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T14:49:03.001071Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.17540"},"observation_digest":"sha256:6a489b0810edd16a0d2fd7939cef76954ef64cd8beb05f456515719eabf1ac99","observation_id":"d2ea6fd7-5eff-4b53-b889-422edfe5faf5","resolution":{"observed_at":"2026-08-07T14:49:03.001071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T14:31:13.310672Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.18536","last_updated":"2025-05-24T06:01:48Z","snapshot_observed_at":"2026-08-07T22:01:21.366221Z","submitted_at":"2025-05-24T06:01:48Z","title":"Reinforcement Fine-Tuning Powers Reasoning Capability of Multimodal Large Language Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T14:31:13.310672Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.18536"},"observation_digest":"sha256:8c74b41bb0c8cc0a17989841689904a0dc6546270614fe625c08cb0cd4de5d04","observation_id":"051c8f17-2752-4dd2-86b8-fe78be8377d2","resolution":{"observed_at":"2026-08-07T14:31:13.310672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T13:14:36.335219Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.22407","last_updated":"2025-05-28T14:37:21Z","snapshot_observed_at":"2026-08-08T22:51:12.746309Z","submitted_at":"2025-05-28T14:37:21Z","title":"Self-Reflective Reinforcement Learning for Diffusion-based Image Reasoning Generation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T13:14:36.335219Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.22407"},"observation_digest":"sha256:821aed9481a76d7ba7e544a2817052612b01918a961f73438c6d1c28ac2dfca1","observation_id":"5c976fdc-86e1-4356-bcef-5e81b52d8b8e","resolution":{"observed_at":"2026-08-07T13:14:36.335219Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T12:53:22.895103Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step.CoRR, abs/2501.13926, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.23380","last_updated":"2025-05-29T12:00:15Z","snapshot_observed_at":"2026-08-09T02:35:27.089432Z","submitted_at":"2025-05-29T12:00:15Z","title":"UniRL: Self-Improving Unified Multimodal Models via Supervised and Reinforcement Learning","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T12:53:22.895103Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.23380"},"observation_digest":"sha256:6f57dff3e93d87635701e4891647a5c639421afc89c02bc35cbd64f7c4029e2d","observation_id":"3e8053ea-4c2d-445e-aebf-e7b2b1ee8a09","resolution":{"observed_at":"2026-08-07T12:53:22.895103Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T12:48:51.044539Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.23493","last_updated":"2025-05-29T14:43:46Z","snapshot_observed_at":"2026-08-07T20:36:09.670028Z","submitted_at":"2025-05-29T14:43:46Z","title":"R2I-Bench: Benchmarking Reasoning-Driven Text-to-Image Generation","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T12:48:51.044539Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.23493"},"observation_digest":"sha256:5a45a16ad4bf3731e33d648a480f55b6488324e1a255da2fe8d73154fcb8e82a","observation_id":"c30ed327-5349-46e7-b27a-bb7d9de92dad","resolution":{"observed_at":"2026-08-07T12:48:51.044539Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T12:45:45.740034Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.23693","last_updated":"2025-05-29T17:31:13Z","snapshot_observed_at":"2026-08-07T18:41:28.769570Z","submitted_at":"2025-05-29T17:31:13Z","title":"VF-Eval: Evaluating Multimodal LLMs for Generating Feedback on AIGC Videos","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T12:45:45.740034Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.23693"},"observation_digest":"sha256:9923e8fdde756f7d435c7d84e9e8d8de492203dd2f30b005cd3d22d45295ff53","observation_id":"8a0adae8-f90e-48c0-a83b-bca2917ef604","resolution":{"observed_at":"2026-08-07T12:45:45.740034Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T12:18:32.622198Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24787","last_updated":"2025-05-30T16:48:14Z","snapshot_observed_at":"2026-08-07T15:30:21.145991Z","submitted_at":"2025-05-30T16:48:14Z","title":"Draw ALL Your Imagine: A Holistic Benchmark and Agent Framework for Complex Instruction-based Image Generation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T12:18:32.622198Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.24787"},"observation_digest":"sha256:ff270e4f00c8328198fe01d70ad5c22c1dc52eab212947c25691e13b7cb69c04","observation_id":"b3f0c90b-2bef-45b1-972f-8a68e39d3793","resolution":{"observed_at":"2026-08-07T12:18:32.622198Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T12:19:21.683953Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step.arXiv preprint arXiv:2501.13926, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24875","last_updated":"2025-06-05T17:51:58Z","snapshot_observed_at":"2026-08-08T15:12:18.884282Z","submitted_at":"2025-05-30T17:59:48Z","title":"ReasonGen-R1: CoT for Autoregressive Image generation models through SFT and RL","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T12:19:21.683953Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2505.24875"},"observation_digest":"sha256:202bef38574784975fe35497fbc9865d4faa53766aa97997eb4227adbb822fba","observation_id":"081cfe8c-5871-4cda-a8be-8a5bfe6693bf","resolution":{"observed_at":"2026-08-07T12:19:21.683953Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T11:33:20.782033Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step.arXiv preprint arXiv:2501.13926, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.02161","last_updated":"2026-07-12T08:53:18Z","snapshot_observed_at":"2026-08-07T11:27:10.636069Z","submitted_at":"2025-06-02T18:44:07Z","title":"TIIF-Bench: How Does Your T2I Model Follow Your Instructions?","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T11:33:20.782033Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2506.02161"},"observation_digest":"sha256:e8bbae5460c16d4b6ddb0c6ee6258b17ae9bf96542c2dc197e74573965613efa","observation_id":"553d2b69-b567-4d94-9e80-6517d6fa6cf0","resolution":{"observed_at":"2026-08-07T11:33:20.782033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2506.03530","last_updated":"2026-05-22T06:55:17Z","snapshot_observed_at":"2026-08-01T13:49:24.957131Z","submitted_at":"2025-06-04T03:22:44Z","title":"How Far Are We from Generating Missing Modalities with Foundation Models?","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-25T08:15:12.947854Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2506.03530"},"observation_digest":"sha256:8cb2990bf429896cf80e6bcf3a22c09cdb4076496a2bd5ea6ad94854e591d2ab","observation_id":"f932e840-318a-45e1-9f6f-875b17f5d0a0","resolution":{"observed_at":"2026-05-25T08:15:33.584502Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T10:28:48.738793Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step.arXiv preprint arXiv:2501.13926, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.05331","last_updated":"2025-06-05T17:59:02Z","snapshot_observed_at":"2026-08-09T02:08:19.234840Z","submitted_at":"2025-06-05T17:59:02Z","title":"MINT-CoT: Enabling Interleaved Visual Tokens in Mathematical Chain-of-Thought Reasoning","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T10:28:48.738793Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2506.05331"},"observation_digest":"sha256:96a12ae7e38d9d44a2439604c562c6808338a3624a9edb7ef9d2c53f801626a3","observation_id":"d4c71618-6fcb-43ad-bb39-3eb3b7ef5313","resolution":{"observed_at":"2026-08-07T10:28:48.738793Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T11:22:44.051500Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.05384","last_updated":"2025-06-12T16:38:10Z","snapshot_observed_at":"2026-08-07T20:38:22.705426Z","submitted_at":"2025-06-03T10:11:51Z","title":"Q-Ponder: A Unified Training Pipeline for Reasoning-based Visual Quality Assessment","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:22:44.051500Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2506.05384"},"observation_digest":"sha256:348979c8f180ee5e51c6c19e4aa17263be7e79e42694bd1c5b968dea8cee9103","observation_id":"23dca6e3-ebf9-46de-8494-cb092f44b9b5","resolution":{"observed_at":"2026-08-07T11:22:44.051500Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T10:23:12.511095Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.05501","last_updated":"2025-06-05T18:36:33Z","snapshot_observed_at":"2026-08-07T20:35:29.372904Z","submitted_at":"2025-06-05T18:36:33Z","title":"FocusDiff: Advancing Fine-Grained Text-Image Alignment for Autoregressive Visual Generation through RL","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T10:23:12.511095Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2506.05501"},"observation_digest":"sha256:5c857220cb4a84af28daebc0fea1528e16ba207b9b7acb36314544f98e00f665","observation_id":"96b75dea-f716-4ad6-8b6c-2b525eb43572","resolution":{"observed_at":"2026-08-07T10:23:12.511095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-07T00:40:20.308003Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12830","last_updated":"2025-06-15T12:22:55Z","snapshot_observed_at":"2026-08-09T11:19:36.846193Z","submitted_at":"2025-06-15T12:22:55Z","title":"ComplexBench-Edit: Benchmarking Complex Instruction-Driven Image Editing via Compositional Dependencies","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T00:40:20.308003Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2506.12830"},"observation_digest":"sha256:c766a2215fe073151d04a82be6f56a2f8f69da3e0e3c15bd6f5ef972bba5d66f","observation_id":"294919e4-d9e3-4c01-9070-ad012af279f8","resolution":{"observed_at":"2026-08-07T00:40:20.308003Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2506.16796","last_updated":"2026-04-13T02:17:17Z","snapshot_observed_at":"2026-07-30T08:45:44.946546Z","submitted_at":"2025-06-20T07:21:21Z","title":"RealSR-R1: Reinforcement Learning for Real-World Image Super-Resolution with Vision-Language Chain-of-Thought","version":4},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-19T08:32:20.566798Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2506.16796"},"observation_digest":"sha256:7b3deae0745ef40e26d81a4f8482d5bf87da85b34ee2f216ff6e71e50789f98b","observation_id":"26c2409b-046a-4e3e-87f2-a4338ad6850b","resolution":{"observed_at":"2026-05-19T08:33:02.188477Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2506.18871","last_updated":"2026-04-21T17:32:47Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-23T17:38:54Z","title":"OmniGen2: Towards Instruction-Aligned Multimodal Generation","version":4},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-19T07:47:34.464711Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2506.18871"},"observation_digest":"sha256:4d4558cff9fecbdf772223a218c8c7c08d28ba7c3a9ed1bb7c7136b86d5bdfed","observation_id":"b8fadb8b-3881-4257-9666-e71df3074dc1","resolution":{"observed_at":"2026-05-19T07:52:10.744905Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-06T16:42:42.906185Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.12761","last_updated":"2025-07-17T03:33:46Z","snapshot_observed_at":"2026-08-07T20:51:13.041685Z","submitted_at":"2025-07-17T03:33:46Z","title":"Think-Before-Draw: Decomposing Emotion Semantics & Fine-Grained Controllable Expressive Talking Head Generation","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T16:42:42.906185Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2507.12761"},"observation_digest":"sha256:885176271ac4dd4d59ca73e512cb48bb8c8dd4a05e55b26e20252b28f8057683","observation_id":"1255e509-c867-452e-a7d0-a6c4e527b6b2","resolution":{"observed_at":"2026-08-06T16:42:42.906185Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-05T22:32:51.804152Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.06905","last_updated":"2025-08-26T12:18:14Z","snapshot_observed_at":"2026-08-07T03:28:03.679250Z","submitted_at":"2025-08-09T09:36:21Z","title":"MultiRef: Controllable Image Generation with Multiple Visual References","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T22:32:51.804152Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2508.06905"},"observation_digest":"sha256:53bdb1e75e9b0bf31f576b7aa4069798670d96c915da356e3c4d0ba1015110d1","observation_id":"e73c98f9-504a-442d-99f6-92ef727bdfd2","resolution":{"observed_at":"2026-08-05T22:32:51.804152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-05T20:44:12.432203Z","title":"Can we generate images with cot? let's verify and reinforce image generation step by step","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.09987","last_updated":"2025-08-13T17:59:28Z","snapshot_observed_at":"2026-08-09T07:45:34.680113Z","submitted_at":"2025-08-13T17:59:28Z","title":"Echo-4o: Harnessing the Power of GPT-4o Synthetic Images for Improved Image Generation","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-05T20:44:12.432203Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2508.09987"},"observation_digest":"sha256:234de3038956fdcb0d967a7f2bce699525880ddb4c7f633be0b1f0443ce26d57","observation_id":"1c41b22a-8423-4dd8-9ff9-b14b3a0b11d8","resolution":{"observed_at":"2026-08-05T20:44:12.432203Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-04T22:36:07.963130Z","title":"Can we generate images with cot? let's verify and reinforce image generation step by step","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.07295","last_updated":"2026-06-25T06:17:40Z","snapshot_observed_at":"2026-08-04T22:36:03.033298Z","submitted_at":"2025-09-08T23:59:32Z","title":"Reconstruction Alignment Improves Unified Multimodal Models","version":4},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-04T22:36:07.963130Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2509.07295"},"observation_digest":"sha256:187f5a3485a1f8c90e93269e5235c13b26d271374e6bf3ef33a864516e023ab9","observation_id":"c7d72a8c-5a3b-47d9-8018-1bc7d72d3564","resolution":{"observed_at":"2026-08-04T22:36:07.963130Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2512.07348","last_updated":"2026-04-28T10:02:14Z","snapshot_observed_at":"2026-08-06T12:32:11.849043Z","submitted_at":"2025-12-08T09:40:11Z","title":"MICo-150K: A Comprehensive Dataset Advancing Multi-Image Composition","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-17T00:20:58.483350Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2512.07348"},"observation_digest":"sha256:e79329188c8f47eff1e96f857513ef6cf5c88589cc72bb1489575f173ce688ef","observation_id":"23d37186-de7a-4292-9018-bf899153e531","resolution":{"observed_at":"2026-05-17T00:21:23.343091Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-13T23:27:11.006580Z","title":"arXiv preprint arXiv:2501.13926 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.16870","last_updated":"2026-07-31T08:55:01Z","snapshot_observed_at":"2026-08-05T23:10:24.488750Z","submitted_at":"2026-03-17T17:59:55Z","title":"Demystifying Video Reasoning","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-13T23:27:11.006580Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2603.16870"},"observation_digest":"sha256:2cab191310ce6137c647f218250f4b3345475a94d0ff5e64bdd80ca5b70e9041","observation_id":"d3c4c922-b22d-4177-b0c6-590127eb25af","resolution":{"observed_at":"2026-07-13T23:27:11.006580Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-08-03T02:33:54.197765Z","title":"arXiv preprint arXiv:2501.13926 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.16870","last_updated":"2026-07-31T08:55:01Z","snapshot_observed_at":"2026-08-05T23:10:24.488750Z","submitted_at":"2026-03-17T17:59:55Z","title":"Demystifying Video Reasoning","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-03T02:33:54.197765Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2603.16870"},"observation_digest":"sha256:c7e5044ee33e2a60cedb23af40dd0fd96d7077a08d175871c3df5df60144d1e1","observation_id":"17c05c7b-cb86-4bb6-a657-19a2bcbeb692","resolution":{"observed_at":"2026-08-03T02:33:54.197765Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2604.02355","last_updated":"2026-03-12T12:49:26Z","snapshot_observed_at":"2026-07-06T22:51:44.663367Z","submitted_at":"2026-03-12T12:49:26Z","title":"From Broad Exploration to Stable Synthesis: Entropy-Guided Optimization for Autoregressive Image Generation","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-15T12:50:13.764159Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2604.02355"},"observation_digest":"sha256:86d422c57f52f1031bce6ba5c4b973536fff32e8cc91b963605558392a9ae97f","observation_id":"53b29446-6ab5-4120-bea8-1451dcd9a279","resolution":{"observed_at":"2026-05-15T12:50:37.069714Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2606.04264","last_updated":"2026-06-02T22:30:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-06-02T22:30:46Z","title":"UniCanvas: A Diffusion-base Unified Model for Text-in-Image Joint Generation","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-28T10:23:43.501656Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2606.04264"},"observation_digest":"sha256:6aaf7768dc118a49eaad0b13f39d39e955a8a86b59bc8bd4777d8143ae45bd73","observation_id":"aafb7341-2c9b-4aa3-a2c7-b6e392123764","resolution":{"observed_at":"2026-07-02T03:06:29.520773Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2606.08231","last_updated":"2026-06-06T15:39:29Z","snapshot_observed_at":"2026-08-07T16:58:05.550639Z","submitted_at":"2026-06-06T15:39:29Z","title":"Test-Time Scaling in Multimodal Foundation Models: A Comprehensive Survey of Generation and Reasoning","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-06-27T19:36:57.231932Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2606.08231"},"observation_digest":"sha256:6b72cf16eb4a91b7dd860044d80a483384e04eb0c08b44fdbb0821faa48c9f9f","observation_id":"1ea53590-1d6f-4963-b40a-5219ead2ee88","resolution":{"observed_at":"2026-07-02T21:37:25.295363Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2606.17888","last_updated":"2026-06-16T13:09:32Z","snapshot_observed_at":"2026-08-08T01:57:17.698279Z","submitted_at":"2026-06-16T13:09:32Z","title":"MathVis-Fine: Aligning Visual Supervision with Necessity via Progressive Dependency-Guided Training for Multimodal Mathematical Reasoning","version":1},"reference_index":101,"source":"arxiv_source","source_observed_at":"2026-06-27T01:23:40.564561Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2606.17888"},"observation_digest":"sha256:53796dda7c19a3c89e208b1198ec0e9e3d50944fc110372822f55ade55aaa638","observation_id":"e823c830-b45d-49d7-92a8-b36f8676a353","resolution":{"observed_at":"2026-07-03T20:18:57.837081Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":"2501.13926","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-04T16:39:58.175257Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step","venue":null,"work_id":"42bc4633-5505-4914-a32c-2e3d11cc6e03","year":2025},"citing_paper":{"arxiv_id":"2606.24849","last_updated":"2026-06-23T17:28:00Z","snapshot_observed_at":"2026-08-03T11:30:36.402988Z","submitted_at":"2026-06-23T17:28:00Z","title":"IV-CoT: Implicit Visual Chain-of-Thought for Structure-Aware Text-to-Image Generation","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-06-26T00:19:49.071495Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2606.24849"},"observation_digest":"sha256:475941a5404b5cc4d1e61236c9ca00092bac2e08ca4dacfc5e833cc700fbc7ce","observation_id":"fdf1b1a6-3013-4842-8d7f-e210294a46c2","resolution":{"observed_at":"2026-07-04T16:39:58.177090Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13926","snapshot_observed_at":"2026-07-14T01:07:06.861298Z","title":"Can we generate images with cot? let’s verify and reinforce image generation step by step.arXiv preprint arXiv:2501.13926, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.10004","last_updated":"2026-07-10T22:03:01Z","snapshot_observed_at":"2026-08-03T02:26:02.235143Z","submitted_at":"2026-07-10T22:03:01Z","title":"Model Guides You How to Draw: Adaptive Visual Gating for Unified Multimodal Reasoning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-14T01:07:06.861298Z"},"links":{"cited_paper":"/paper/2501.13926","citing_paper":"/paper/2607.10004"},"observation_digest":"sha256:c38b347da7534def91a58e006415f1622184e223301864a55e6765070214c08c","observation_id":"1730a604-eb8c-49ea-9e3f-50043610d767","resolution":{"observed_at":"2026-07-14T01:07:06.861298Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2501.13926/citation-record","integrity":"/paper/2501.13926/integrity","json":"/paper/2501.13926/citation-record.json","paper":"/paper/2501.13926"},"outbound":[],"paper":{"arxiv_id":"2501.13926","last_updated":"2025-07-23T16:09:10Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T20:25:08.527025Z","submitted_at":"2025-01-23T18:59:43Z","title":"Can We Generate Images with CoT? Let's Verify and Reinforce Image Generation Step by Step"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 34 inbound Pith citation observations for arXiv:2501.13926."}