{"as_of":"2026-08-08T10:15:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:92c20f687aad76730df39a20de545170d6244000a9e7d5dfb8688d4045519121","coverage":[{"denominator":71,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":71,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T08:19:36.399257Z","state":"measured"},{"denominator":71,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":71,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2510.21590/citation-record","integrity":"/paper/2510.21590/integrity","json":"/paper/2510.21590/citation-record.json","paper":"/paper/2510.21590"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.194082Z","title":"Dream- clear: High-capacity real-world image restoration with privacy-safe dataset curation.Advances in Neural Informa- tion Processing Systems, 37:55443–55469, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.194082Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:608deb5e28a10088ea26ac3ce246d7e2a7ccc9df9fd615bdb9f6c81cd96b6123","observation_id":"339e10e0-4ef9-4394-b1fb-457f03e8f21f","resolution":{"observed_at":"2026-08-04T08:19:31.194082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.222889Z","title":"Toward real-world single image super-resolution: A new benchmark and a new model","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.222889Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:7f714d6c4c0f91ea9b40a7a15e569a3f61866b955f14e50d84a35ca220ca8a2a","observation_id":"b1fc2489-0b0e-42e1-a917-22b040a99764","resolution":{"observed_at":"2026-08-04T08:19:31.222889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.276024Z","title":"Freeman, Michael Ru- binstein, Yuanzhen Li, and Dilip Krishnan","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.276024Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:8ce84cfa9260c8c27370906d2583241b30421639939e442b3b9cf8400c2d733b","observation_id":"c509bd76-212e-4a25-82f8-3da4db42ed7c","resolution":{"observed_at":"2026-08-04T08:19:31.276024Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.353902Z","title":"Scene text tele- scope: Text-focused scene image super-resolution","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.353902Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:e4e16f6c165a755360e8694f793dfbef0f668f140a85a6e10221b151c73eee70","observation_id":"02470c48-0048-4b42-affd-13b4a01106ad","resolution":{"observed_at":"2026-08-04T08:19:31.353902Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.428776Z","title":"Activating more pixels in image super- resolution transformer","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.428776Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:30ff8441adb5d15ea0567220b9a6e69311ba560f64168cfeb5f0e6c520649684","observation_id":"e4c23b2f-2c5d-4cca-b85f-6a697f21a4be","resolution":{"observed_at":"2026-08-04T08:19:31.428776Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.484383Z","title":"Effective diffusion transformer architecture for image super- resolution","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.484383Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:a4a0c2432fd393aa9066a7ac983ca0a9bf7a0805829cf67e8bd903140ff346b7","observation_id":"f458d417-7ba3-4be2-bd50-c9b783b68d0e","resolution":{"observed_at":"2026-08-04T08:19:31.484383Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.533436Z","title":"Paddleocr 3.0 technical report, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.533436Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:07621b27cbcf9d02218b84ed6a6e2cf5079d6e14b55ee5a1ee68f91f03b0e7c5","observation_id":"eb947f1d-f953-4352-93d6-9fa545f3dea2","resolution":{"observed_at":"2026-08-04T08:19:31.533436Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.591317Z","title":"Textual alchemy: Coformer for scene text understanding","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.591317Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:d26fc518225c5180363f543782028c7a30aae27e6af07234de6f67939836815d","observation_id":"61886070-40c6-4a7f-a750-a2b28084e7e0","resolution":{"observed_at":"2026-08-04T08:19:31.591317Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.659151Z","title":"Diffusion models beat gans on image synthesis.Advances in neural informa- tion processing systems, 34:8780–8794, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.659151Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:01fc737285532336ce1109a673824843659139bc742a21e026cf7191dedca1b9","observation_id":"568d0e4d-6615-46ee-b00f-909464b03970","resolution":{"observed_at":"2026-08-04T08:19:31.659151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.720207Z","title":"Image quality assessment: Unifying structure and texture similarity.TPAMI, 44(5):2567–2581, 2020","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.720207Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:988262ca53f6dfc26efab6e8c23ae8c37a42a9f6f0ec836448b2ed6c93e781c3","observation_id":"09320c05-db9e-4db7-95be-5e1163dd77e7","resolution":{"observed_at":"2026-08-04T08:19:31.720207Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.807662Z","title":"Learning a deep convolutional network for image super-resolution","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.807662Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:04815ec8a56ba2075eda83e545989b8c3d603c2a216f83172bcf5adbacd6ad5d","observation_id":"d5b4aad2-a779-429b-b485-cc0a2051f043","resolution":{"observed_at":"2026-08-04T08:19:31.807662Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.864552Z","title":"Boosting optical character recognition: A super- resolution approach.arXiv preprint, 2015","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.864552Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:d12ce88178931e196e5df308b6bedb2507568a797a2517ff4742c3a9cdd9b614","observation_id":"4fae3781-7f4b-45d9-8af5-c2e3019b0f20","resolution":{"observed_at":"2026-08-04T08:19:31.864552Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.946966Z","title":"TSD-SR: one-step diffusion with target score distillation for real-world image super-resolution","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.946966Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:c29c769fc078fa6f20cb6cb296c86361ea4262b99fca4e286b6a1c221d851e24","observation_id":"e55e8cf6-6edf-4958-800d-d33253458006","resolution":{"observed_at":"2026-08-04T08:19:31.946966Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:31.991837Z","title":"Tsd-sr: One-step diffusion with target score distillation for real-world image super-resolution","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:31.991837Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:fd50aa44392fd0faa9b2cbc0f0ee63f610f08b4bf3bce89dda7256a3f0fb0ac9","observation_id":"b7094c7f-cc01-4149-b68a-5cba8e50a29a","resolution":{"observed_at":"2026-08-04T08:19:31.991837Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.044812Z","title":"Dit4sr: Taming diffusion transformer for real-world image super-resolution","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.044812Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:c0bf9546ae5310a732fabffcb18f23b9adedda32573b0a54c2ae62ea45059958","observation_id":"1aaee2f4-566e-4e17-bd21-382c8a4f5eb4","resolution":{"observed_at":"2026-08-04T08:19:32.044812Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.117991Z","title":"Scaling recti- fied flow transformers for high-resolution image synthesis","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.117991Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:647793197a314d88f9779a30579b1dc1a24991fbdb0f80bdc795c9acf512f79a","observation_id":"fcbe7bb3-3202-4f60-b4af-0eefb139a9b6","resolution":{"observed_at":"2026-08-04T08:19:32.117991Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.169021Z","title":"Gans trained by a two time-scale update rule converge to a local nash equilib- rium.NeurIPS, 30, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.169021Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:e80a6501eb6103691e121cb400abbafaa0b2fd409c03901e6ed7b88e2ff3990f","observation_id":"7c33c5bd-ff65-4071-b090-ac3deb740080","resolution":{"observed_at":"2026-08-04T08:19:32.169021Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.240563Z","title":"Denoising diffu- sion probabilistic models.Advances in Neural Information Processing Systems, 33:6840–6851, 2020","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.240563Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:abefc2bb3b05a239e2cbb9be8b769077163cc01f7c8ef57492292aa8593dbf65","observation_id":"42b45c8c-d465-4491-9962-27d9b97fe3c8","resolution":{"observed_at":"2026-08-04T08:19:32.240563Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04641","last_updated":"2025-06-05T05:23:10Z","snapshot_observed_at":"2026-08-07T10:34:22.875988Z","submitted_at":"2025-06-05T05:23:10Z","title":"Text-Aware Real-World Image Super-Resolution via Diffusion Model with Joint Segmentation Decoders","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.04641","snapshot_observed_at":"2026-08-04T08:19:32.285046Z","title":"Text-aware real-world image super- resolution via diffusion model with joint segmentation de- coders.arXiv preprint arXiv:2506.04641, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.285046Z"},"links":{"cited_paper":"/paper/2506.04641","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:58ac035d777ed439a57b46ee4015256077b347bae3dcad50430cb8fe7555eca7","observation_id":"c00f79ab-8327-4421-93b3-bb8b8a4f2a82","resolution":{"observed_at":"2026-08-04T08:19:32.285046Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.364415Z","title":"Prestu: Pre-training for scene-text understanding","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.364415Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:027e85ae2d1b8f56b555405bb75fddf02bbb70ea31e7b35ba869564d45c31d0b","observation_id":"bc606baa-e225-4427-8c33-732e9a90ecee","resolution":{"observed_at":"2026-08-04T08:19:32.364415Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1312.6114","last_updated":"2022-12-10T21:04:00Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2013-12-20T20:58:10Z","title":"Auto-Encoding Variational Bayes","version":11},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1312.6114","snapshot_observed_at":"2026-08-04T08:19:32.439906Z","title":"Auto-encoding varia- tional bayes.arXiv preprint arXiv:1312.6114, 2013","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.439906Z"},"links":{"cited_paper":"/paper/1312.6114","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:92edb040a284e3655e896edd08889582947111026d87cd18e3f1bae4b1c5c844","observation_id":"a2d60cf0-07e9-4a41-8e75-b559c9fde05b","resolution":{"observed_at":"2026-08-04T08:19:32.439906Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.513849Z","title":"RIPE: Reinforcement Learning on Unlabeled Image Pairs for Ro- bust Keypoint Extraction.arXiv, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.513849Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:9eb4df36beb9aa87aea850fc9690c5762421889ec45b877f1c301ce48654897e","observation_id":"c166ea59-cb96-4089-95a1-bbe7815e4cf4","resolution":{"observed_at":"2026-08-04T08:19:32.513849Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.03001","last_updated":"2022-06-14T12:00:10Z","snapshot_observed_at":"2026-07-06T13:18:06.376608Z","submitted_at":"2022-06-07T04:33:50Z","title":"PP-OCRv3: More Attempts for the Improvement of Ultra Lightweight OCR System","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.03001","snapshot_observed_at":"2026-08-04T08:19:32.565063Z","title":"Pp-ocrv3: More attempts for the improvement of ultra lightweight OCR sys- tem.CoRR, abs/2206.03001, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.565063Z"},"links":{"cited_paper":"/paper/2206.03001","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:c37dff5affdf6a41ffbba49852749b6f07c621c976573628249945208d77183b","observation_id":"be8d879f-a4fb-4f96-a016-5755649661b4","resolution":{"observed_at":"2026-08-04T08:19:32.565063Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.617024Z","title":"Navigation-guided sparse scene representation for end-to-end autonomous driving","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.617024Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:d0fd930b535a54960bead2fda6454473dfbe6f95994b7dec38cceca347ee35eb","observation_id":"8f45f607-c906-45d1-8561-9861da0f51c6","resolution":{"observed_at":"2026-08-04T08:19:32.617024Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.686514Z","title":"Learning generative structure prior for blind text image super-resolution","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.686514Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:81abae985002a320137dd668c53e4b36843f7aad229ecdab2a06b6bfb9f11b0a","observation_id":"bd83cb58-86b8-4e90-83bb-e4ed47600c10","resolution":{"observed_at":"2026-08-04T08:19:32.686514Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.754934Z","title":"Lsdir: A large scale dataset for image restoration","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.754934Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:132e0637d626ae4a558c05f18cabd885612773c046f2f86bf748fd546c9bd21f","observation_id":"448d4a26-2f17-4f1a-9d7f-d700e1bb4fab","resolution":{"observed_at":"2026-08-04T08:19:32.754934Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.824139Z","title":"Diff- bir: Toward blind image restoration with generative diffusion prior","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.824139Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:7b6f2ed7ebc734663ef4bcdba1512b24a351a243b33e42b518526ad474c3fa93","observation_id":"5f337f5d-ef1d-4a4b-b2e7-47e5e5f367f8","resolution":{"observed_at":"2026-08-04T08:19:32.824139Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1711.05101","last_updated":"2019-01-04T21:01:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-11-14T14:24:06Z","title":"Decoupled Weight Decay Regularization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.05101","snapshot_observed_at":"2026-08-04T08:19:32.884783Z","title":"Decoupled weight decay regularization.arXiv preprint arXiv:1711.05101, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.884783Z"},"links":{"cited_paper":"/paper/1711.05101","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:94166c9610ca6b5f244593b90108c6adbe389639ec5b7ad7b79a8911b0582f59","observation_id":"b3930b09-82d0-4dd2-88ee-de562766c2c5","resolution":{"observed_at":"2026-08-04T08:19:32.884783Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:32.963596Z","title":"Object recognition from local scale-invariant features","venue":null,"work_id":null,"year":1999},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:32.963596Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:9fa8943c6a50b732b7d898668bb8679bc1db5d915c81f8d15f08699879d4e75c","observation_id":"47a70fbd-7091-4c5d-a38c-0da74c36b3d4","resolution":{"observed_at":"2026-08-04T08:19:32.963596Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.051138Z","title":"A text atten- tion network for spatial deformation robust scene text image super-resolution","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.051138Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:f7c85b1d1e1dbbe910dfd4164442728c4241659fee38d38cd1841258763e38d1","observation_id":"10572ebb-263a-443d-9bb5-dc431a9712c7","resolution":{"observed_at":"2026-08-04T08:19:33.051138Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.197237Z","title":"A benchmark for chinese-english scene text im- age super-resolution","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.197237Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:dfac0f056809a0337dca5d1fa5360218dd31dd11a9789d80e8f72a524d5b6522","observation_id":"1543cf0e-9035-4da0-9b63-e4ab5e4bd30c","resolution":{"observed_at":"2026-08-04T08:19:33.197237Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.311140Z","title":"Plugnet: Degrada- tion aware scene text recognition supervised by a pluggable super-resolution unit","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.311140Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:58e5319a9e8eb9d4f9d1874ad452fe8e5b592f5d30597c96579c20a935e1ef27","observation_id":"830c53f4-9dc2-41d7-bf5a-0227118cde29","resolution":{"observed_at":"2026-08-04T08:19:33.311140Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2204.06125","last_updated":"2022-04-13T01:10:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-04-13T01:10:33Z","title":"Hierarchical Text-Conditional Image Generation with CLIP Latents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.06125","snapshot_observed_at":"2026-08-04T08:19:33.457577Z","title":"Hierarchical text-conditional image gen- eration with clip latents.arXiv preprint arXiv:2204.06125,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.457577Z"},"links":{"cited_paper":"/paper/2204.06125","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:b539a0342463a1e771fc2f85979259cca6797d7d250e0892b51a9b48dd877089","observation_id":"4bbcec5c-b217-4dd8-90dc-946d82aad756","resolution":{"observed_at":"2026-08-04T08:19:33.457577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.533426Z","title":"Addison-Wesley Longman Publishing Co., Inc.,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.533426Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:43d8f42b592b442334497ed72bff296446919ab631468b486a8bde032c384d4f","observation_id":"c58e2bf1-12a3-4cf2-923c-8189858c953a","resolution":{"observed_at":"2026-08-04T08:19:33.533426Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.591347Z","title":"High-resolution image syn- thesis with latent diffusion models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.591347Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:b18634121ccbf63aa8744dd6108538c79879383177f89d1943fe4a5753310e20","observation_id":"9de8e90c-bedb-42d7-a2c9-2be0da4ce694","resolution":{"observed_at":"2026-08-04T08:19:33.591347Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.665453Z","title":"Photorealistic text-to-image diffusion models with deep language understanding.Advances in neural information processing systems, 35:36479–36494, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.665453Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:8878ab7fb38f49db6a978ee240311dcd5a229d12a680603c60a36e326aa2bcab","observation_id":"d1f0abbb-3735-497f-9c13-4be5073a1b84","resolution":{"observed_at":"2026-08-04T08:19:33.665453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.741912Z","title":"Text-diae: a self-supervised degradation invariant autoencoder for text recognition and document enhancement","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.741912Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:c41aa8f85ebe732d9c0fd7d636c76c3925dfdc23ec8388a9c7470b692d83c214","observation_id":"2db7aa16-f13c-4f67-9433-77f717818526","resolution":{"observed_at":"2026-08-04T08:19:33.741912Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.788983Z","title":"A scene-text synthesis engine achieved through learning from decomposed real-world data.IEEE Transactions on Image Processing, 32:5837–5851, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.788983Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:47e881ede095ac502bfb2d011d7e1c414979d54a05498e72deebc78baa91dc0d","observation_id":"c321a4ed-5e89-4577-8e85-1883b76bcd97","resolution":{"observed_at":"2026-08-04T08:19:33.788983Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.889591Z","title":"Anytext: Multilingual visual text genera- tion and editing","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.889591Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:f874ae09d5e1d3aa07b1074b4d39b8c8c59f3cbf610cf5f39ccc79eef344be0c","observation_id":"8e603c16-7da3-4343-ba52-5aa4aa91bfe7","resolution":{"observed_at":"2026-08-04T08:19:33.889591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:33.937874Z","title":"Multi-task dif- fusion model for simultaneous text and image inpainting","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:33.937874Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:af1ab7debb61eb929274d097ab50a931e593e9fac55aa927e16e82d1241571ef","observation_id":"91fb1674-a558-4c7d-9f79-489a2555c2ff","resolution":{"observed_at":"2026-08-04T08:19:33.937874Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.001968Z","title":"Exploiting diffusion prior for real-world image super-resolution.IJCV, 132(12):5929– 5949, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.001968Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:45e610575271a6ebdac43796b3f66a7ced16412d25d3b75fa585fa785137bc61","observation_id":"e42494ad-78f3-427c-95f6-e63352de8932","resolution":{"observed_at":"2026-08-04T08:19:34.001968Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.049171Z","title":"Textsr: Content-aware text super-resolution guided by recognition.arXiv preprint,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.049171Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:18646905eb5f622542c04cde55a76e601a733a383dbe1e92f1136258cb44a846","observation_id":"26865975-2f3e-4b95-8d26-ea10d750fc92","resolution":{"observed_at":"2026-08-04T08:19:34.049171Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.103393Z","title":"Scene text image super-resolution in the wild","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.103393Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:b45c89b627847595604136b93091e60d706c6f638d408ecc03a8a40cc12b2d1d","observation_id":"12717e45-a89d-4f44-8727-f564c75c27a2","resolution":{"observed_at":"2026-08-04T08:19:34.103393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.136432Z","title":"Real-esrgan: Training real-world blind super-resolution with pure synthetic data","venue":null,"work_id":null,"year":1905},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.136432Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:903ca3c68b9d8f88b0396056d5346c78367290052d30136969d702b1a2ac5f49","observation_id":"de2bf15e-0107-4f48-9a34-b013b4c5071c","resolution":{"observed_at":"2026-08-04T08:19:34.136432Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.211895Z","title":"Sinsr: diffusion-based image super- resolution in a single step","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.211895Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:e81e957e0ba3d6d5f5c3984367e48f287e1bcf2a8a27b359034ffccb239417eb","observation_id":"eca33495-30b4-472f-9152-7e25cd79f934","resolution":{"observed_at":"2026-08-04T08:19:34.211895Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.292340Z","title":"One-step effective diffusion network for real-world image super-resolution.NeurIPS, 37:92529–92553, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.292340Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:c8e963eb47f327773ba250c438668b13e341159457e6685cf2b501a24f13e429","observation_id":"e28cd67e-8737-460b-a70b-946ff73e9c16","resolution":{"observed_at":"2026-08-04T08:19:34.292340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.369047Z","title":"Seesr: Towards semantics-aware real-world image super-resolution","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.369047Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:a44033524ba321f299c868f87996244710dc0de7c736dd1cc6feb06e38c99be6","observation_id":"86c57d81-16c1-47f4-9bec-a07505d3cba3","resolution":{"observed_at":"2026-08-04T08:19:34.369047Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.439149Z","title":"Pixel-aware stable diffusion for realistic im- age super-resolution and personalized stylization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.439149Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:86c4c449e9ca24246708ce18f9d7952e459bd51ebae051219021a3e19512031d","observation_id":"8d6bbf64-9ccc-4ebe-9457-69aec876033c","resolution":{"observed_at":"2026-08-04T08:19:34.439149Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.491001Z","title":"Hi-sam: Marrying segment anything model for hierarchical text segmentation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.491001Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:dc10f85aabec68a2617f12b0a5d2c4956238f0be7f1436ff93aba39aba26d30c","observation_id":"18e1a187-d42b-4d9e-91bd-5d81899dbbbe","resolution":{"observed_at":"2026-08-04T08:19:34.491001Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.594826Z","title":"Scaling up to excellence: Practicing model scaling for photo- realistic image restoration in the wild","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.594826Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:c7cd459f29a0b6b7a9ebb37d612915a77d83864e9a4fd628cd19b913cc5ab81f","observation_id":"ea34aee5-847d-4eaa-ab54-b0c4a4cd8473","resolution":{"observed_at":"2026-08-04T08:19:34.594826Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2112.15093","last_updated":"2022-11-25T12:03:17Z","snapshot_observed_at":"2026-08-06T12:07:45.000627Z","submitted_at":"2021-12-30T15:30:52Z","title":"Benchmarking Chinese Text Recognition: Datasets, Baselines, and an Empirical Study","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.15093","snapshot_observed_at":"2026-08-04T08:19:34.701538Z","title":"Benchmarking chinese text recognition: Datasets, baselines, and an empirical study.arXiv preprint arXiv:2112.15093, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.701538Z"},"links":{"cited_paper":"/paper/2112.15093","citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:317534b097800454e3a394903160663aa1b16fad602bc1b2b88ebd88517cc72b","observation_id":"0c91679c-ab11-479d-b776-6d4fff8e48f1","resolution":{"observed_at":"2026-08-04T08:19:34.701538Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.760835Z","title":"Chinese text recognition with a pre-trained clip-like model through image-ids aligning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.760835Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:6f17b3395c0e06beadcb19556896ad9cabbcd2260bce491ff36bd28216b55986","observation_id":"4c9e92ba-f460-4fda-8586-3887b04bda04","resolution":{"observed_at":"2026-08-04T08:19:34.760835Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.818672Z","title":"Resshift: Efficient diffusion model for image super- 10 resolution by residual shifting.NeurIPS, 36:13294–13307,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.818672Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:f02af3e7a70b8188c757b8f4a40c59efd9122a0875cd77734ce29e63310d0c81","observation_id":"4fd39a34-b503-409a-bc34-4a3974267e90","resolution":{"observed_at":"2026-08-04T08:19:34.818672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.866274Z","title":"A normalized levenshtein distance metric.TPAMI, 29(6):1091–1095, 2007","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.866274Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:4359e894afbca45b29865e18d6503b7139cd4a5a5a987b1538da4edf0174ad71","observation_id":"9f24f488-9b6f-4643-8ea1-417da3701a5f","resolution":{"observed_at":"2026-08-04T08:19:34.866274Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:34.992679Z","title":"Designing a practical degradation model for deep blind image super-resolution","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:34.992679Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:db8993aafcffe7673a37d688ae97112af04612d6b8d1159a4330cdf4c10d0bf0","observation_id":"a2dfb829-8879-45cf-b4c0-c7558092e146","resolution":{"observed_at":"2026-08-04T08:19:34.992679Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.064112Z","title":"Adding conditional control to text-to-image diffusion models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.064112Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:2c1a69b8551b5bcec1fe660b547a40f2f1a20150786adaf36ce6f0cf443af4e0","observation_id":"ca9d5dc9-8ee9-4cda-a369-8625cb875a5b","resolution":{"observed_at":"2026-08-04T08:19:35.064112Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.137372Z","title":"The unreasonable effectiveness of deep features as a perceptual metric","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.137372Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:503e9c5b34341d34686d4160b66dc114fc41821cda839b665df750e0c2790c93","observation_id":"aa1795bb-ebc3-4bfa-87e8-f7ed0c2ad58e","resolution":{"observed_at":"2026-08-04T08:19:35.137372Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.223800Z","title":"Diffusion-based blind text image super-resolution","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.223800Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:9603347fceeeac72eb5536e15f1e2dc9a5cafefd76c3b64cac82796c60af6764","observation_id":"a09e546a-66b0-43dd-a195-e909e50655e9","resolution":{"observed_at":"2026-08-04T08:19:35.223800Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.340971Z","title":"Artbank: Artistic style transfer with pre-trained diffusion model and implicit style prompt bank","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.340971Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:3d20f73b3f50b0e58c47ec22dd75adf40f776b0afe41b32e6cf42b291ba459b6","observation_id":"3222ed53-4ab9-46f4-8a63-a2372e7a2c78","resolution":{"observed_at":"2026-08-04T08:19:35.340971Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.415603Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.415603Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:422b092d833bef0e71538e7582477e7974c71c3d0ad15c75399307c931f5042b","observation_id":"78ad9f21-1b4f-4d7a-9e79-c51d45818f7d","resolution":{"observed_at":"2026-08-04T08:19:35.415603Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.489374Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.489374Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:047789fa521880cc88aaf7cd483716f1ef503f7a8ffeebcaae517bb54660e05d","observation_id":"1893056b-f313-4214-9503-45db8e2872c1","resolution":{"observed_at":"2026-08-04T08:19:35.489374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.579021Z","title":"5.1 and Sec","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.579021Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:f48ae626280eb5f11f8e5029daffdd71205a3a04dc9813f272f792b441ed82e4","observation_id":"748a2c85-054e-45e6-ac00-f406aae0ac77","resolution":{"observed_at":"2026-08-04T08:19:35.579021Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.688612Z","title":"4 and Sec","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.688612Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:996c8735d6a9ada909f98d2f50279f9136aa14ad21e158b8af475b648730ab8a","observation_id":"1d882a1e-211d-46ab-be18-6b5faf7a7a29","resolution":{"observed_at":"2026-08-04T08:19:35.688612Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.764990Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.764990Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:ce15f44ca48b02b7a6aaa962ee4f54b2907bf3c2e4f0a0559e9d09c4fe2009ab","observation_id":"defd6b68-c4fb-406a-8c9c-bd40f7b7a1c2","resolution":{"observed_at":"2026-08-04T08:19:35.764990Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.826491Z","title":"All wording and factual content were reviewed and approved by the au- thors","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.826491Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:cdf67c0d95460ac733b5776d960545a1d8e6a703f611577b65118e47934ee644","observation_id":"1cea1a4a-9000-4d19-aa9b-0100f7d06dc2","resolution":{"observed_at":"2026-08-04T08:19:35.826491Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:35.924205Z","title":"We first use PP-OCRV5 [7] for the rough annotation, then we manually filter the images and annotations","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:35.924205Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:238dbe53f455023ca99663db90de8bd55e21ecf2fe4d8ba46997800097a7db59","observation_id":"2a58c2cf-736a-4cda-b454-17c75dbc077c","resolution":{"observed_at":"2026-08-04T08:19:35.924205Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:36.042595Z","title":"Our synthetic dataset builds upon LSDIR [26], con- taining 27,000 triplets(x H , xL, xm)","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:36.042595Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:4c35f14b51ea052c1163c923eb720db2030ae3c4a300aebacf56a2d4c369832d","observation_id":"5fe5d8c3-568c-423e-83bb-8c1ff23d6ba6","resolution":{"observed_at":"2026-08-04T08:19:36.042595Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:36.143455Z","title":"7, we provide detailed statistics on the composition of the UZ-ST dataset","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:36.143455Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:9911a158060611bb92e3c4b7af742d905b2e39bcbcbea78fae0a981ea2446bfa","observation_id":"23485b49-bc97-4686-af4a-901a53855642","resolution":{"observed_at":"2026-08-04T08:19:36.143455Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:36.214091Z","title":"During the comparison, we implement our strategy using SIFT and uti- lize the raw images from the 35mm dataset of our proposed UZ-ST dataset for evaluation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:36.214091Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:636a09ba2e98621f618e97457a1362a680c3e9bf5dbf92c691507d20a809134f","observation_id":"959de467-f888-473e-a592-f7d1367a4c3f","resolution":{"observed_at":"2026-08-04T08:19:36.214091Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:36.296738Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:36.296738Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:b7092732fe3ca8813ac04aa99b7b1ab11111ba25923810644ebdbc64daa4b424","observation_id":"357c050c-096c-4d3c-97af-4c68efe1d752","resolution":{"observed_at":"2026-08-04T08:19:36.296738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T08:19:36.399257Z","title":"9, stage 1 is a standard diffusion process that takes multiple steps in inference, the efficiency of our model may be suboptimal compared to one-step methods","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance","version":3},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-04T08:19:36.399257Z"},"links":{"citing_paper":"/paper/2510.21590"},"observation_digest":"sha256:f8580a3aaec5be1c97f3aca40a54482f4b6140289b422317c37f7edd073e8994","observation_id":"38043389-953e-4edd-b6c2-2fdba5aedb73","resolution":{"observed_at":"2026-08-04T08:19:36.399257Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2510.21590","last_updated":"2026-08-03T06:53:36Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-06T23:28:48.341710Z","submitted_at":"2025-10-24T15:59:04Z","title":"Restore Text First, Enhance Image Later: Two-Stage Scene Text Image Super-Resolution with Glyph Structure Guidance"},"reference_resolution":{"displayed":71,"state_counts":{"malformed_identifier":2,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":69,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":71},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 71 of 71 outbound references and 0 inbound Pith citation observations for arXiv:2510.21590."}