{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:37TFQU37MB77YIXMKLX7Y63DWC","short_pith_number":"pith:37TFQU37","schema_version":"1.0","canonical_sha256":"dfe658537f607ffc22ec52effc7b63b087d4afc399d44b04f1b4f55cd46eb956","source":{"kind":"arxiv","id":"2302.03027","version":1},"attestation_state":"computed","paper":{"title":"Zero-shot Image-to-Image Translation","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.GR","cs.LG"],"primary_cat":"cs.CV","authors_text":"Gaurav Parmar, Jingwan Lu, Jun-Yan Zhu, Krishna Kumar Singh, Richard Zhang, Yijun Li","submitted_at":"2023-02-06T18:59:51Z","abstract_excerpt":"Large-scale text-to-image generative models have shown their remarkable ability to synthesize diverse and high-quality images. However, it is still challenging to directly apply these models for editing real images for two reasons. First, it is hard for users to come up with a perfect text prompt that accurately describes every visual detail in the input image. Second, while existing models can introduce desirable changes in certain regions, they often dramatically alter the input content and introduce unexpected changes in unwanted regions. In this work, we propose pix2pix-zero, an image-to-i"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2302.03027","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CV","submitted_at":"2023-02-06T18:59:51Z","cross_cats_sorted":["cs.GR","cs.LG"],"title_canon_sha256":"ccf6be8e61222e5120afade4233058d4a2993308a07c58bf97b7d7c78b27ee69","abstract_canon_sha256":"3ab76f1cd5cf969aadde45929756b7abf8c5fa363b2475fbdc782fcca38f95d3"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T05:39:09.220035Z","signature_b64":"4Ii8SLPlv2NLbngnkpmwey6REHUt4LBYA5yghFEVeJ9ZffL4kiCzFGPOjLQGxq62IcDqMk2704/Ynm8uQ5t8BQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"dfe658537f607ffc22ec52effc7b63b087d4afc399d44b04f1b4f55cd46eb956","last_reissued_at":"2026-07-05T05:39:09.219570Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T05:39:09.219570Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Zero-shot Image-to-Image Translation","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.GR","cs.LG"],"primary_cat":"cs.CV","authors_text":"Gaurav Parmar, Jingwan Lu, Jun-Yan Zhu, Krishna Kumar Singh, Richard Zhang, Yijun Li","submitted_at":"2023-02-06T18:59:51Z","abstract_excerpt":"Large-scale text-to-image generative models have shown their remarkable ability to synthesize diverse and high-quality images. However, it is still challenging to directly apply these models for editing real images for two reasons. First, it is hard for users to come up with a perfect text prompt that accurately describes every visual detail in the input image. Second, while existing models can introduce desirable changes in certain regions, they often dramatically alter the input content and introduce unexpected changes in unwanted regions. In this work, we propose pix2pix-zero, an image-to-i"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2302.03027","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2302.03027/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2302.03027","created_at":"2026-07-05T05:39:09.219629+00:00"},{"alias_kind":"arxiv_version","alias_value":"2302.03027v1","created_at":"2026-07-05T05:39:09.219629+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2302.03027","created_at":"2026-07-05T05:39:09.219629+00:00"},{"alias_kind":"pith_short_12","alias_value":"37TFQU37MB77","created_at":"2026-07-05T05:39:09.219629+00:00"},{"alias_kind":"pith_short_16","alias_value":"37TFQU37MB77YIXM","created_at":"2026-07-05T05:39:09.219629+00:00"},{"alias_kind":"pith_short_8","alias_value":"37TFQU37","created_at":"2026-07-05T05:39:09.219629+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":5,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.01825","citing_title":"ROGLE: Robust Global-Local Alignment with Automated Region Supervision for Text-Based Person Search","ref_index":140,"is_internal_anchor":false},{"citing_arxiv_id":"2606.04306","citing_title":"Organizational Control Layer: Governance Infrastructure at the Execution Boundary of LLM Agent Systems","ref_index":119,"is_internal_anchor":false},{"citing_arxiv_id":"2606.01825","citing_title":"ROGLE: Robust Global-Local Alignment with Automated Region Supervision for Text-Based Person Search","ref_index":140,"is_internal_anchor":false},{"citing_arxiv_id":"2605.18010","citing_title":"Functionalization via Structure Completion and Motion Rectification","ref_index":111,"is_internal_anchor":false},{"citing_arxiv_id":"2302.05543","citing_title":"Adding Conditional Control to Text-to-Image Diffusion Models","ref_index":64,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/37TFQU37MB77YIXMKLX7Y63DWC","json":"https://pith.science/pith/37TFQU37MB77YIXMKLX7Y63DWC.json","graph_json":"https://pith.science/api/pith-number/37TFQU37MB77YIXMKLX7Y63DWC/graph.json","events_json":"https://pith.science/api/pith-number/37TFQU37MB77YIXMKLX7Y63DWC/events.json","paper":"https://pith.science/paper/37TFQU37"},"agent_actions":{"view_html":"https://pith.science/pith/37TFQU37MB77YIXMKLX7Y63DWC","download_json":"https://pith.science/pith/37TFQU37MB77YIXMKLX7Y63DWC.json","view_paper":"https://pith.science/paper/37TFQU37","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2302.03027&json=true","fetch_graph":"https://pith.science/api/pith-number/37TFQU37MB77YIXMKLX7Y63DWC/graph.json","fetch_events":"https://pith.science/api/pith-number/37TFQU37MB77YIXMKLX7Y63DWC/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/37TFQU37MB77YIXMKLX7Y63DWC/action/timestamp_anchor","attest_storage":"https://pith.science/pith/37TFQU37MB77YIXMKLX7Y63DWC/action/storage_attestation","attest_author":"https://pith.science/pith/37TFQU37MB77YIXMKLX7Y63DWC/action/author_attestation","sign_citation":"https://pith.science/pith/37TFQU37MB77YIXMKLX7Y63DWC/action/citation_signature","submit_replication":"https://pith.science/pith/37TFQU37MB77YIXMKLX7Y63DWC/action/replication_record"}},"created_at":"2026-07-05T05:39:09.219629+00:00","updated_at":"2026-07-05T05:39:09.219629+00:00"}