{"as_of":"2026-08-13T12:34:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:24f74d6f0ee260be1474a5640027ef57e1d2fcd693bd84ad4964575989b231a8","coverage":[{"denominator":71,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":71,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T08:19:36.399257Z","state":"measured"},{"denominator":71,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":71,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2510.21590/citation-record","integrity":"/paper/2510.21590/integrity","json":"/paper/2510.21590/citation-record.json","paper":"/paper/2510.21590"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.194082Z","title":"Dream- clear: High-capacity real-world image restoration with privacy-safe dataset curation.Advances in Neural Informa- tion Processing Systems, 37:55443–55469, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.194082Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:6019cf070cc5b7d7e6b011b6cfa7f656fbec2200ae1baa4bfc525bbecf3004b1","observation_id":"339e10e0-4ef9-4394-b1fb-457f03e8f21f","resolution":{"observed_at":"2026-08-04T08:19:31.194082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.222889Z","title":"Toward real-world single image super-resolution: A new benchmark and a new model","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.222889Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:820d5af16543934b44e3d532d5af3499a38e23c77dad03b4d371a999f4ff17c8","observation_id":"b1fc2489-0b0e-42e1-a917-22b040a99764","resolution":{"observed_at":"2026-08-04T08:19:31.222889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.276024Z","title":"Freeman, Michael Ru- binstein, Yuanzhen Li, and Dilip Krishnan","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.276024Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:ac79b46d9e265ba6fca030743ed86f8031ca179e78485ac425d59b682e4e0bb6","observation_id":"c509bd76-212e-4a25-82f8-3da4db42ed7c","resolution":{"observed_at":"2026-08-04T08:19:31.276024Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.353902Z","title":"Scene text tele- scope: Text-focused scene image super-resolution","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.353902Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:9e554965704fd1844c665a65aa5311d8627cccf65ca8e92f76eaa46a50e1b469","observation_id":"02470c48-0048-4b42-affd-13b4a01106ad","resolution":{"observed_at":"2026-08-04T08:19:31.353902Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.428776Z","title":"Activating more pixels in image super- resolution transformer","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.428776Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:ce69add20c60a05ba23ebe9ded471652434294390165ba9bb241968f88a3b29a","observation_id":"e4c23b2f-2c5d-4cca-b85f-6a697f21a4be","resolution":{"observed_at":"2026-08-04T08:19:31.428776Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.484383Z","title":"Effective diffusion transformer architecture for image super- resolution","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.484383Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:455c34daa4f720a16bf8b4142131b7d23bf387c228a3de98fa02bbf672f45980","observation_id":"f458d417-7ba3-4be2-bd50-c9b783b68d0e","resolution":{"observed_at":"2026-08-04T08:19:31.484383Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.533436Z","title":"Paddleocr 3.0 technical report, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.533436Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:2d18c6f6d767f0274f2539a25e6a6698785ac1a1485bfaec037fce3613b9cf36","observation_id":"eb947f1d-f953-4352-93d6-9fa545f3dea2","resolution":{"observed_at":"2026-08-04T08:19:31.533436Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.591317Z","title":"Textual alchemy: Coformer for scene text understanding","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.591317Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:fa865bd145090d2f01473e965f50f5717115eba124ade9da6f2519b3730beaf5","observation_id":"61886070-40c6-4a7f-a750-a2b28084e7e0","resolution":{"observed_at":"2026-08-04T08:19:31.591317Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.659151Z","title":"Diffusion models beat gans on image synthesis.Advances in neural informa- tion processing systems, 34:8780–8794, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.659151Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:81e54ccf8b3c241c099cd072ea321762ccee9eb25d1e6bf6d494cfcf0d4f7501","observation_id":"568d0e4d-6615-46ee-b00f-909464b03970","resolution":{"observed_at":"2026-08-04T08:19:31.659151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.720207Z","title":"Image quality assessment: Unifying structure and texture similarity.TPAMI, 44(5):2567–2581, 2020","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.720207Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:a485713f2fdfa4688bd8912dd6842f0efaf019512d81dcdb4b818097db12b534","observation_id":"09320c05-db9e-4db7-95be-5e1163dd77e7","resolution":{"observed_at":"2026-08-04T08:19:31.720207Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.807662Z","title":"Learning a deep convolutional network for image super-resolution","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.807662Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:4fb983f56b8e3dd62b89948252723fa7334b4028e6eefa1cd1a865ea4e3feb7d","observation_id":"d5b4aad2-a779-429b-b485-cc0a2051f043","resolution":{"observed_at":"2026-08-04T08:19:31.807662Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.864552Z","title":"Boosting optical character recognition: A super- resolution approach.arXiv preprint, 2015","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.864552Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:ff97b536eae173eccbd10575c276c54c573f2c90560a326d665786f345d97422","observation_id":"4fae3781-7f4b-45d9-8af5-c2e3019b0f20","resolution":{"observed_at":"2026-08-04T08:19:31.864552Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.946966Z","title":"TSD-SR: one-step diffusion with target score distillation for real-world image super-resolution","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.946966Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:20c0c80a863caed3ddadb5f43d14ac83d715c9a935c14aea1c73303b4f70701f","observation_id":"e55e8cf6-6edf-4958-800d-d33253458006","resolution":{"observed_at":"2026-08-04T08:19:31.946966Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.991837Z","title":"Tsd-sr: One-step diffusion with target score distillation for real-world image super-resolution","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.991837Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:e26b811bc8f6cd101bf203fcfb1eaf6ca0daefeba1603785b9e8d3c101846d07","observation_id":"b7094c7f-cc01-4149-b68a-5cba8e50a29a","resolution":{"observed_at":"2026-08-04T08:19:31.991837Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.044812Z","title":"Dit4sr: Taming diffusion transformer for real-world image super-resolution","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.044812Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:3b3974d571e25d2c388871142a5bc74fb26ea673c1d2137e1dc3a116f5c8dc03","observation_id":"1aaee2f4-566e-4e17-bd21-382c8a4f5eb4","resolution":{"observed_at":"2026-08-04T08:19:32.044812Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.117991Z","title":"Scaling recti- fied flow transformers for high-resolution image synthesis","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.117991Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:a08d7361c50e29f1ad507ff8b3f9970c020ef67a909e38286302ed7032abe80e","observation_id":"fcbe7bb3-3202-4f60-b4af-0eefb139a9b6","resolution":{"observed_at":"2026-08-04T08:19:32.117991Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.169021Z","title":"Gans trained by a two time-scale update rule converge to a local nash equilib- rium.NeurIPS, 30, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.169021Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:05745dad3e65237df4cf4df54169c73aff39a17de5ef8be281406533e83ad0fe","observation_id":"7c33c5bd-ff65-4071-b090-ac3deb740080","resolution":{"observed_at":"2026-08-04T08:19:32.169021Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.240563Z","title":"Denoising diffu- sion probabilistic models.Advances in Neural Information Processing Systems, 33:6840–6851, 2020","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.240563Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:fa6d18099c77322a480ae038566ef6d91e5a675e513e8145901fe616f2d7146f","observation_id":"42b45c8c-d465-4491-9962-27d9b97fe3c8","resolution":{"observed_at":"2026-08-04T08:19:32.240563Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04641","last_updated":"2025-06-05T05:23:10Z","snapshot_observed_at":"2026-08-07T10:34:22.875988Z","submitted_at":"2025-06-05T05:23:10Z","title":"Text-Aware Real-World Image Super-Resolution via Diffusion Model with Joint Segmentation Decoders","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.04641","snapshot_observed_at":"2026-08-04T08:19:32.285046Z","title":"Text-aware real-world image super- resolution via diffusion model with joint segmentation de- coders.arXiv preprint arXiv:2506.04641, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.285046Z"},"links":{"cited_paper":"/paper/2506.04641","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:dbbc443c8124a504ffafb1f94795e7f8d7e2bde7413bf68cbb8cf4a1cb3a0753","observation_id":"c00f79ab-8327-4421-93b3-bb8b8a4f2a82","resolution":{"observed_at":"2026-08-04T08:19:32.285046Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.364415Z","title":"Prestu: Pre-training for scene-text understanding","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.364415Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:e02eac1a05ebedfebecaa47e5c4172b839f45e6146f7cecc8c3827423557987e","observation_id":"bc606baa-e225-4427-8c33-732e9a90ecee","resolution":{"observed_at":"2026-08-04T08:19:32.364415Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1312.6114","last_updated":"2022-12-10T21:04:00Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2013-12-20T20:58:10Z","title":"Auto-Encoding Variational Bayes","version":11},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1312.6114","snapshot_observed_at":"2026-08-04T08:19:32.439906Z","title":"Auto-encoding varia- tional bayes.arXiv preprint arXiv:1312.6114, 2013","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.439906Z"},"links":{"cited_paper":"/paper/1312.6114","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:e725125a47a96e78b4ab890b0ea6e47465d9ea671cdcebce6fda0e6395f23e08","observation_id":"a2d60cf0-07e9-4a41-8e75-b559c9fde05b","resolution":{"observed_at":"2026-08-04T08:19:32.439906Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.513849Z","title":"RIPE: Reinforcement Learning on Unlabeled Image Pairs for Ro- bust Keypoint Extraction.arXiv, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.513849Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:d17c29250f8da99e1ebe2694c4bc8401cf34397d82046eb633efe88cf3e75f4c","observation_id":"c166ea59-cb96-4089-95a1-bbe7815e4cf4","resolution":{"observed_at":"2026-08-04T08:19:32.513849Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.03001","last_updated":"2022-06-14T12:00:10Z","snapshot_observed_at":"2026-07-06T13:18:06.376608Z","submitted_at":"2022-06-07T04:33:50Z","title":"PP-OCRv3: More Attempts for the Improvement of Ultra Lightweight OCR System","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.03001","snapshot_observed_at":"2026-08-04T08:19:32.565063Z","title":"Pp-ocrv3: More attempts for the improvement of ultra lightweight OCR sys- tem.CoRR, abs/2206.03001, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.565063Z"},"links":{"cited_paper":"/paper/2206.03001","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:d34a7e70e89aa456a0db735fff9162a425cd9269ced3fa98dbfe816d2fe945d8","observation_id":"be8d879f-a4fb-4f96-a016-5755649661b4","resolution":{"observed_at":"2026-08-04T08:19:32.565063Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.617024Z","title":"Navigation-guided sparse scene representation for end-to-end autonomous driving","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.617024Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:f4b2ea29da5b1f17275ac4b9a2153663553ba6760eba3c23c3e63e68e1d263c6","observation_id":"8f45f607-c906-45d1-8561-9861da0f51c6","resolution":{"observed_at":"2026-08-04T08:19:32.617024Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.686514Z","title":"Learning generative structure prior for blind text image super-resolution","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.686514Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:772d6804def7db5020afad34a0337861018ed822d0677a44e85c313c6943928e","observation_id":"bd83cb58-86b8-4e90-83bb-e4ed47600c10","resolution":{"observed_at":"2026-08-04T08:19:32.686514Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.754934Z","title":"Lsdir: A large scale dataset for image restoration","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.754934Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:4063c35f6aee6bf8a59b55d0263cde362828748690ea2a61efa51e26deee0140","observation_id":"448d4a26-2f17-4f1a-9d7f-d700e1bb4fab","resolution":{"observed_at":"2026-08-04T08:19:32.754934Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.824139Z","title":"Diff- bir: Toward blind image restoration with generative diffusion prior","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.824139Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:1bbf06cdd4910c6b4a92e581c6b697ddc4402aebf9a5c33dbd6c97c1bfbfa8dc","observation_id":"5f337f5d-ef1d-4a4b-b2e7-47e5e5f367f8","resolution":{"observed_at":"2026-08-04T08:19:32.824139Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1711.05101","last_updated":"2019-01-04T21:01:49Z","snapshot_observed_at":"2026-08-09T20:34:52.923500Z","submitted_at":"2017-11-14T14:24:06Z","title":"Decoupled Weight Decay Regularization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.05101","snapshot_observed_at":"2026-08-04T08:19:32.884783Z","title":"Decoupled weight decay regularization.arXiv preprint arXiv:1711.05101, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.884783Z"},"links":{"cited_paper":"/paper/1711.05101","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:5483a3001e172701e579a2c4377dd3d28d904d88bfc03bbdc12033ba29bf23b4","observation_id":"b3930b09-82d0-4dd2-88ee-de562766c2c5","resolution":{"observed_at":"2026-08-04T08:19:32.884783Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.963596Z","title":"Object recognition from local scale-invariant features","venue":null,"work_id":null,"year":1999},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.963596Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:0a2c2f407c6fe14b0acd362bbfe695273d09c09cb7646ccfe683eb627b5386c6","observation_id":"47a70fbd-7091-4c5d-a38c-0da74c36b3d4","resolution":{"observed_at":"2026-08-04T08:19:32.963596Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.051138Z","title":"A text atten- tion network for spatial deformation robust scene text image super-resolution","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.051138Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:0e6c2222a68cceb0213567fcca032a3c7fdf499345966bef225d934362b13217","observation_id":"10572ebb-263a-443d-9bb5-dc431a9712c7","resolution":{"observed_at":"2026-08-04T08:19:33.051138Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.197237Z","title":"A benchmark for chinese-english scene text im- age super-resolution","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.197237Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:dd57bfef0795f21a319c2c84e127305ababbd023447390023ad796a836c510c3","observation_id":"1543cf0e-9035-4da0-9b63-e4ab5e4bd30c","resolution":{"observed_at":"2026-08-04T08:19:33.197237Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.311140Z","title":"Plugnet: Degrada- tion aware scene text recognition supervised by a pluggable super-resolution unit","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.311140Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:3a8475bca613eeafdf59e10aaa02a1d8a491e11db5549bf0f17d88c1ec3e0e16","observation_id":"830c53f4-9dc2-41d7-bf5a-0227118cde29","resolution":{"observed_at":"2026-08-04T08:19:33.311140Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2204.06125","last_updated":"2022-04-13T01:10:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-04-13T01:10:33Z","title":"Hierarchical Text-Conditional Image Generation with CLIP Latents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.06125","snapshot_observed_at":"2026-08-04T08:19:33.457577Z","title":"Hierarchical text-conditional image gen- eration with clip latents.arXiv preprint arXiv:2204.06125,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.457577Z"},"links":{"cited_paper":"/paper/2204.06125","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:f8aa0e8e2f1a815a06279e365bcaebda623b68171dd0a5440cecb409b1418510","observation_id":"4bbcec5c-b217-4dd8-90dc-946d82aad756","resolution":{"observed_at":"2026-08-04T08:19:33.457577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.533426Z","title":"Addison-Wesley Longman Publishing Co., Inc.,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.533426Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:4b643f376f254959eb287f6d8836dcefda8e25c8fcb94c213d3383f7dbf1d0d3","observation_id":"c58e2bf1-12a3-4cf2-923c-8189858c953a","resolution":{"observed_at":"2026-08-04T08:19:33.533426Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.591347Z","title":"High-resolution image syn- thesis with latent diffusion models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.591347Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:868f3df7bc8eaae6da511746df076b1862695e740af37f6fc0e8d756719a29f2","observation_id":"9de8e90c-bedb-42d7-a2c9-2be0da4ce694","resolution":{"observed_at":"2026-08-04T08:19:33.591347Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.665453Z","title":"Photorealistic text-to-image diffusion models with deep language understanding.Advances in neural information processing systems, 35:36479–36494, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.665453Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:6b10ab04dc6bcb4faf4b3143c277f17430a75300aaf11633c385d5bcb891657e","observation_id":"d1f0abbb-3735-497f-9c13-4be5073a1b84","resolution":{"observed_at":"2026-08-04T08:19:33.665453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.741912Z","title":"Text-diae: a self-supervised degradation invariant autoencoder for text recognition and document enhancement","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.741912Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:db85623fceeffcd0dcccc2ab42304f1044ca7ca3c61f1b2e4700cd017fd74b91","observation_id":"2db7aa16-f13c-4f67-9433-77f717818526","resolution":{"observed_at":"2026-08-04T08:19:33.741912Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.788983Z","title":"A scene-text synthesis engine achieved through learning from decomposed real-world data.IEEE Transactions on Image Processing, 32:5837–5851, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.788983Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:0f22d0ac6d5dd20facbd905658605e582c84ec64dceac0c7252fc1d6c0d078cd","observation_id":"c321a4ed-5e89-4577-8e85-1883b76bcd97","resolution":{"observed_at":"2026-08-04T08:19:33.788983Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.889591Z","title":"Anytext: Multilingual visual text genera- tion and editing","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.889591Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:f3f42f64546ad00cb4db1eadae0da8a4f81406948f754399e367a86b92172271","observation_id":"8e603c16-7da3-4343-ba52-5aa4aa91bfe7","resolution":{"observed_at":"2026-08-04T08:19:33.889591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.937874Z","title":"Multi-task dif- fusion model for simultaneous text and image inpainting","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.937874Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:1f81773f073c8442e5b23a9beedc1a454eff9eb333ccc1e6ec28c69a7ce63b20","observation_id":"91fb1674-a558-4c7d-9f79-489a2555c2ff","resolution":{"observed_at":"2026-08-04T08:19:33.937874Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.001968Z","title":"Exploiting diffusion prior for real-world image super-resolution.IJCV, 132(12):5929– 5949, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.001968Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:4c5faf6a3ce181db9b0a1f3ead519ead478450c549003fb23886f6834bb871c5","observation_id":"e42494ad-78f3-427c-95f6-e63352de8932","resolution":{"observed_at":"2026-08-04T08:19:34.001968Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.049171Z","title":"Textsr: Content-aware text super-resolution guided by recognition.arXiv preprint,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.049171Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:39720a4660a5d2613e7afa1198461e8f24a26ec1845c062b7018070bc3949244","observation_id":"26865975-2f3e-4b95-8d26-ea10d750fc92","resolution":{"observed_at":"2026-08-04T08:19:34.049171Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.103393Z","title":"Scene text image super-resolution in the wild","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.103393Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:156e2dcd71bf9580ee1f88c95714151ae93cc5784dd0e0a22fa9340be8cde605","observation_id":"12717e45-a89d-4f44-8727-f564c75c27a2","resolution":{"observed_at":"2026-08-04T08:19:34.103393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.136432Z","title":"Real-esrgan: Training real-world blind super-resolution with pure synthetic data","venue":null,"work_id":null,"year":1905},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.136432Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:cbc052980f2685da2be2dbb8c4c9400873d68a309c412d9246619c99e97d692a","observation_id":"de2bf15e-0107-4f48-9a34-b013b4c5071c","resolution":{"observed_at":"2026-08-04T08:19:34.136432Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.211895Z","title":"Sinsr: diffusion-based image super- resolution in a single step","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.211895Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:8dff23945d5724109c8217891867f48332bcf36fdbae5a5956d14345ca4161ac","observation_id":"eca33495-30b4-472f-9152-7e25cd79f934","resolution":{"observed_at":"2026-08-04T08:19:34.211895Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.292340Z","title":"One-step effective diffusion network for real-world image super-resolution.NeurIPS, 37:92529–92553, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.292340Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:dd685e20b21246d117d63fd36dbff140ae68eaa671bb5ff22c5e6c49401631cc","observation_id":"e28cd67e-8737-460b-a70b-946ff73e9c16","resolution":{"observed_at":"2026-08-04T08:19:34.292340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.369047Z","title":"Seesr: Towards semantics-aware real-world image super-resolution","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.369047Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:cdc9e9d7fe1b4fbe5c2323daac815eceb8253e215728d81df062d36d31993bbd","observation_id":"86c57d81-16c1-47f4-9bec-a07505d3cba3","resolution":{"observed_at":"2026-08-04T08:19:34.369047Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.439149Z","title":"Pixel-aware stable diffusion for realistic im- age super-resolution and personalized stylization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.439149Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:1ded9ae99cce7c2ed65be67422974b3e058b36bb8bd391b090c31957557bc228","observation_id":"8d6bbf64-9ccc-4ebe-9457-69aec876033c","resolution":{"observed_at":"2026-08-04T08:19:34.439149Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.491001Z","title":"Hi-sam: Marrying segment anything model for hierarchical text segmentation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.491001Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:e2d976307fa2acc187cd1120e51d80accb9379170742de1aefdee3792331c247","observation_id":"18e1a187-d42b-4d9e-91bd-5d81899dbbbe","resolution":{"observed_at":"2026-08-04T08:19:34.491001Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.594826Z","title":"Scaling up to excellence: Practicing model scaling for photo- realistic image restoration in the wild","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.594826Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:8e3171d900d2e72793f1ad5fd81a429958b950250dae55ff2c0b0d6cecb69731","observation_id":"ea34aee5-847d-4eaa-ab54-b0c4a4cd8473","resolution":{"observed_at":"2026-08-04T08:19:34.594826Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2112.15093","last_updated":"2022-11-25T12:03:17Z","snapshot_observed_at":"2026-08-06T12:07:45.000627Z","submitted_at":"2021-12-30T15:30:52Z","title":"Benchmarking Chinese Text Recognition: Datasets, Baselines, and an Empirical Study","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.15093","snapshot_observed_at":"2026-08-04T08:19:34.701538Z","title":"Benchmarking chinese text recognition: Datasets, baselines, and an empirical study.arXiv preprint arXiv:2112.15093, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.701538Z"},"links":{"cited_paper":"/paper/2112.15093","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:4b8ac4ad0c423a3311a9fc87ea82e0d14e9a9be6026e03c444c5dc044e43c385","observation_id":"0c91679c-ab11-479d-b776-6d4fff8e48f1","resolution":{"observed_at":"2026-08-04T08:19:34.701538Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.760835Z","title":"Chinese text recognition with a pre-trained clip-like model through image-ids aligning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.760835Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:fda15e77a36bd7677528467c4c55a49439cdaaab71c1aebb340494323a4fb410","observation_id":"4c9e92ba-f460-4fda-8586-3887b04bda04","resolution":{"observed_at":"2026-08-04T08:19:34.760835Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.818672Z","title":"Resshift: Efficient diffusion model for image super- 10 resolution by residual shifting.NeurIPS, 36:13294–13307,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.818672Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:0cf021c6ac9740076f53759f87aa0e22b651c6699e8536271be3ad0beda2fc1d","observation_id":"4fd39a34-b503-409a-bc34-4a3974267e90","resolution":{"observed_at":"2026-08-04T08:19:34.818672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.866274Z","title":"A normalized levenshtein distance metric.TPAMI, 29(6):1091–1095, 2007","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.866274Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:5e900b80a9244eb966689ba43c6eb291901024e6e871f2f82e48015d955a292a","observation_id":"9f24f488-9b6f-4643-8ea1-417da3701a5f","resolution":{"observed_at":"2026-08-04T08:19:34.866274Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.992679Z","title":"Designing a practical degradation model for deep blind image super-resolution","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.992679Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:b772a87529a23a3f55d680fdf33aebab544e635d700b115c3b167b50a0260b78","observation_id":"a2dfb829-8879-45cf-b4c0-c7558092e146","resolution":{"observed_at":"2026-08-04T08:19:34.992679Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.064112Z","title":"Adding conditional control to text-to-image diffusion models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.064112Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:eb916fe399027690c451ea995467a8bdaadcc02532eec2058e38ae4387fb5b37","observation_id":"ca9d5dc9-8ee9-4cda-a369-8625cb875a5b","resolution":{"observed_at":"2026-08-04T08:19:35.064112Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.137372Z","title":"The unreasonable effectiveness of deep features as a perceptual metric","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.137372Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:4243d6713cedc80692b72f7662d91f9ee56195d9d43a35c454aa0d8a971a2d1f","observation_id":"aa1795bb-ebc3-4bfa-87e8-f7ed0c2ad58e","resolution":{"observed_at":"2026-08-04T08:19:35.137372Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.223800Z","title":"Diffusion-based blind text image super-resolution","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.223800Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:a6aa7c9fb2ead0298d89630a79aa5366ee91f5668a753d29c5b9e12095f8a0ed","observation_id":"a09e546a-66b0-43dd-a195-e909e50655e9","resolution":{"observed_at":"2026-08-04T08:19:35.223800Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.340971Z","title":"Artbank: Artistic style transfer with pre-trained diffusion model and implicit style prompt bank","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.340971Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:1b1dc80f82929dfc2116ca8a6057dfe1ef4374b190ed00baaf8da8481664205c","observation_id":"3222ed53-4ab9-46f4-8a63-a2372e7a2c78","resolution":{"observed_at":"2026-08-04T08:19:35.340971Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.415603Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.415603Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:bf035710bf2bd60f32b12b26767f8be43bb1c2f1fbf267bf72f9e3516f1119fe","observation_id":"78ad9f21-1b4f-4d7a-9e79-c51d45818f7d","resolution":{"observed_at":"2026-08-04T08:19:35.415603Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.489374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.489374Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:7c5fae10774513b4dcb1ef70ca00b85ead873cd298184b32a01f2d434ebf9066","observation_id":"1893056b-f313-4214-9503-45db8e2872c1","resolution":{"observed_at":"2026-08-04T08:19:35.489374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.579021Z","title":"5.1 and Sec","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.579021Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:c136195c6848f2ece11e8b3c3d15dddada30ce0b6070bc72a2f07625825ef1ad","observation_id":"748a2c85-054e-45e6-ac00-f406aae0ac77","resolution":{"observed_at":"2026-08-04T08:19:35.579021Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.688612Z","title":"4 and Sec","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.688612Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:72af1cb29d463cae9cd41de4968450584adc79400c329ec8d5080f87aa923860","observation_id":"1d882a1e-211d-46ab-be18-6b5faf7a7a29","resolution":{"observed_at":"2026-08-04T08:19:35.688612Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.764990Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.764990Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:22e8a7f800c2d2a594fb331d2a69dd3b9515987f7fdd958d78159eed238a77f3","observation_id":"defd6b68-c4fb-406a-8c9c-bd40f7b7a1c2","resolution":{"observed_at":"2026-08-04T08:19:35.764990Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.826491Z","title":"All wording and factual content were reviewed and approved by the au- thors","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.826491Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:2e8e72b8436bfd926b258efba563212bc23d0b47ad4e09ff6595d394ca6afcd6","observation_id":"1cea1a4a-9000-4d19-aa9b-0100f7d06dc2","resolution":{"observed_at":"2026-08-04T08:19:35.826491Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.924205Z","title":"We first use PP-OCRV5 [7] for the rough annotation, then we manually filter the images and annotations","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.924205Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:58ad6302a6c4d4206c4bf23a84f7550cddd51f1ed190e04e9b2fe1dc320860ba","observation_id":"2a58c2cf-736a-4cda-b454-17c75dbc077c","resolution":{"observed_at":"2026-08-04T08:19:35.924205Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:36.042595Z","title":"Our synthetic dataset builds upon LSDIR [26], con- taining 27,000 triplets(x H , xL, xm)","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:36.042595Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:6f9eaa69590430a395c9fd39f657804016eeb1b9c1ad899acbd9cb0e183c05cd","observation_id":"5fe5d8c3-568c-423e-83bb-8c1ff23d6ba6","resolution":{"observed_at":"2026-08-04T08:19:36.042595Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:36.143455Z","title":"7, we provide detailed statistics on the composition of the UZ-ST dataset","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:36.143455Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:8d0eac928e1ebfeb0b9645b98b04fab9e5d3fcdbb2cf787262ac804c01b3a6b7","observation_id":"23485b49-bc97-4686-af4a-901a53855642","resolution":{"observed_at":"2026-08-04T08:19:36.143455Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:36.214091Z","title":"During the comparison, we implement our strategy using SIFT and uti- lize the raw images from the 35mm dataset of our proposed UZ-ST dataset for evaluation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:36.214091Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:58a69d173a52c4c63a656b458f76d0688a40c429de47a205aabf6e011ea8e829","observation_id":"959de467-f888-473e-a592-f7d1367a4c3f","resolution":{"observed_at":"2026-08-04T08:19:36.214091Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:36.296738Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:36.296738Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:9aadbd4f0671d11b9e48cf3a725c579a926fa4533fef75d5ef850666c27a96f5","observation_id":"357c050c-096c-4d3c-97af-4c68efe1d752","resolution":{"observed_at":"2026-08-04T08:19:36.296738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:36.399257Z","title":"9, stage 1 is a standard diffusion process that takes multiple steps in inference, the efficiency of our model may be suboptimal compared to one-step methods","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:36.399257Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:38ee78b2022e0861d0a67c19e98d763fa5e2734895cd05dcc52f56d7724c2a8e","observation_id":"38043389-953e-4edd-b6c2-2fdba5aedb73","resolution":{"observed_at":"2026-08-04T08:19:36.399257Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-10T16:34:39.341164Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance"},"reference_resolution":{"displayed":71,"state_counts":{"malformed_identifier":2,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":69,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":71},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 71 of 71 outbound references and 0 inbound Pith citation observations for arXiv:2510.21590."}