{"as_of":"2026-08-21T08:51:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:19dd63205efa82bfda730175218d05d8f9912eae5d7adac5d86538c2dfddd029","coverage":[{"denominator":52,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":52,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T18:10:40.699063Z","state":"measured"},{"denominator":54,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":54,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T22:58:08.761058Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T22:17:26.248429Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.19002","snapshot_observed_at":"2026-08-04T22:58:08.761058Z","title":"Enhancing reward models for high-quality image generation: Beyond text-image align- ment.arXiv preprint arXiv:2507.19002, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.06942","last_updated":"2025-09-11T17:14:11Z","snapshot_observed_at":"2026-08-14T05:10:53.900739Z","submitted_at":"2025-09-08T17:54:08Z","title":"Directly Aligning the Full Diffusion Trajectory with Fine-Grained Human Preference","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-04T22:58:08.761058Z"},"links":{"cited_paper":"/paper/2507.19002","citing_paper":"/paper/2509.06942"},"observation_digest":"sha256:51ab160e8ca4e8c478118f266978d765659a8a86a6e7feb6bbfdf17022318dda","observation_id":"dde1f6ee-f58c-4ec4-829d-5167820bbb4f","resolution":{"observed_at":"2026-08-04T22:58:08.761058Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"cited_work":{"arxiv_id":"2507.19002","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.19002","snapshot_observed_at":"2026-07-02T22:17:26.248429Z","title":null,"venue":null,"work_id":"35ac14cf-e1e4-433f-8277-f3680dc0edf1","year":2025},"citing_paper":{"arxiv_id":"2606.08147","last_updated":"2026-06-06T12:56:08Z","snapshot_observed_at":"2026-08-12T12:07:31.674697Z","submitted_at":"2026-06-06T12:56:08Z","title":"Biological Reasoning-Informed Regression for Interpretable Regulatory DNA Activity Prediction","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-27T19:01:43.892354Z"},"links":{"cited_paper":"/paper/2507.19002","citing_paper":"/paper/2606.08147"},"observation_digest":"sha256:1b479a3ca7a7a124b2023bc117e9fb430f05e3decc495019bd00add8286fb87d","observation_id":"e036bf74-7a24-4ac0-9d53-3addf87407b4","resolution":{"observed_at":"2026-07-02T22:17:26.249802Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2507.19002/citation-record","integrity":"/paper/2507.19002/integrity","json":"/paper/2507.19002/citation-record.json","paper":"/paper/2507.19002"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.749981Z","title":"Stable diffusion 3.5 large","venue":null,"work_id":"4bf7fd65-58db-4fda-aba3-aa30a536a49b","year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.446332Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:c3929fbd9b28b2abd8816e5c91d8fa768369460ca4df0fe65f221588969b3370","observation_id":"bda621ab-8dc1-488e-b203-7ac91700a7aa","resolution":{"observed_at":"2026-08-15T18:10:41.754640Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.732824Z","title":"Stable diffusion 3.5 large turbo","venue":null,"work_id":"17033a0d-597f-4494-9716-34c10ccb2da6","year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.451776Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:a3d5fa93e8666571bb3e44a1fe410bb1232303bc8cf5c967ecba0997fda5bcc2","observation_id":"bdcdf7f6-40b0-4ff7-9b88-1194774267e5","resolution":{"observed_at":"2026-08-15T18:10:41.738730Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.13942","last_updated":"2024-05-10T06:03:33Z","snapshot_observed_at":"2026-08-18T03:46:04.128908Z","submitted_at":"2024-01-25T04:53:03Z","title":"StyleInject: Parameter Efficient Tuning of Text-to-Image Diffusion Models","version":2},"cited_work":{"arxiv_id":"2401.13942","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.13942","snapshot_observed_at":"2026-08-15T18:10:40.997849Z","title":"StyleInject: Parameter Efficient Tuning of Text-to-Image Diffusion Models","venue":"cs.CV","work_id":"dc026d0e-68b1-4729-a29e-bbe552092a83","year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.456972Z"},"links":{"cited_paper":"/paper/2401.13942","citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:02186e42af9045453634510f14f0bcd56844121fb701ebcfac53362449e0f3c7","observation_id":"573944af-e5ef-42c8-bbd1-54de38fbdbc2","resolution":{"observed_at":"2026-08-15T18:10:41.004662Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.714008Z","title":"Reproducible scaling laws for contrastive language-image learning","venue":null,"work_id":"cce978e8-fbdd-431c-b9b5-3ac2c28fe9e5","year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.462867Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:d28384861312e7b3bc9e082a1ff673b7b36d283a3027f86b21a85cf5bfbc69e7","observation_id":"e5b4c844-71c8-4833-b917-2b45b1f955a8","resolution":{"observed_at":"2026-08-15T18:10:41.720578Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.694434Z","title":"Directly fine-tuning diffusion models on differ- entiable rewards, 2024","venue":null,"work_id":"372a7541-2ab3-4214-9b59-4b21dda7a740","year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.468064Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:cde1a32f6ed8b3c101c858202165e26c1adaaf1e34afb9af8646e1eaf9ccbfd8","observation_id":"e6d0d5cb-a730-44b2-8556-b97a1a7254b5","resolution":{"observed_at":"2026-08-15T18:10:41.700929Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.669305Z","title":"Cogview2: Faster and better text-to-image generation via hierarchical transformers, 2022","venue":null,"work_id":"e2482b24-2fdb-4d9d-a894-609bc7828f46","year":2022},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.473505Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:220e5b7ee1bb4f19c8af2a3148f42ece5f40f4dd59d78e519f47d8a34b802624","observation_id":"500d36d0-40c8-4c7b-8961-0073801a2dd8","resolution":{"observed_at":"2026-08-15T18:10:41.678216Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.03206","last_updated":"2024-03-05T18:45:39Z","snapshot_observed_at":"2026-08-13T01:47:21.043519Z","submitted_at":"2024-03-05T18:45:39Z","title":"Scaling Rectified Flow Transformers for High-Resolution Image Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.03206","snapshot_observed_at":"2026-08-15T18:10:40.479119Z","title":"Scaling rectified flow transformers for high-resolution image synthesis","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.479119Z"},"links":{"cited_paper":"/paper/2403.03206","citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:0c605984f67f6d7a861d47f87d70a75648025b3dcd520331e35fb2be865a178b","observation_id":"390dba54-4fc4-41bd-8b5d-d2d02b28660e","resolution":{"observed_at":"2026-08-15T18:10:40.479119Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.11513","last_updated":"2023-10-17T18:20:03Z","snapshot_observed_at":"2026-08-16T14:51:16.293174Z","submitted_at":"2023-10-17T18:20:03Z","title":"GenEval: An Object-Focused Framework for Evaluating Text-to-Image Alignment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.11513","snapshot_observed_at":"2026-08-15T18:10:40.485271Z","title":"Geneval: An object-focused framework for evaluating text-to-image alignment","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.485271Z"},"links":{"cited_paper":"/paper/2310.11513","citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:03d28f83e9048b9f86d23f6f5fe6179e5a83f602cefade85e5e57f124f870997","observation_id":"e8b89537-d542-4e7f-89d2-74c6d3e1af27","resolution":{"observed_at":"2026-08-15T18:10:40.485271Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.647796Z","title":"Op- timizing prompts for text-to-image generation","venue":null,"work_id":"a69c2ad9-fdd8-4573-b94a-798dd0470c6d","year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.490467Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:fd596ae7d4562d18c70ec602714f252b2844c44ed9cadbca991e3eec13c7c75d","observation_id":"d936a42a-9b81-40ca-9540-0f26a149814e","resolution":{"observed_at":"2026-08-15T18:10:41.653449Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.630705Z","title":"Gans trained by a two time-scale update rule converge to a local nash equilibrium, 2018","venue":null,"work_id":"1fe2e589-c62f-43ff-8a9e-66039aa1fb1b","year":2018},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.495128Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:ce2888817f8e5d9ddaf642526237f4b434192d7b314cc58f13255a6727b03dc2","observation_id":"b83c7890-169a-4d37-a81d-ef36c7559271","resolution":{"observed_at":"2026-08-15T18:10:41.636242Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.613275Z","title":"Denois- ing diffusion probabilistic models","venue":null,"work_id":"d7e2d4ef-9f5e-4076-8d9e-d13b06193001","year":2020},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.499777Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:5fdc4b945c65d6b546ddd026da77507e7fde2397114be5a380fde57530d5a430","observation_id":"6386b555-4c89-4fd2-8d19-d0b47a1ab62f","resolution":{"observed_at":"2026-08-15T18:10:41.619395Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.596811Z","title":"Information technology – digital compres- sion and coding of continuous-tone still images – re- quirements and guidelines","venue":null,"work_id":"bc21541a-d48b-4353-ab01-a67dc3aee2c7","year":1992},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.505101Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:d8a194099b902ce3cb6042fd1b5cda0e7fb3a913c521755382ec3279613f57a7","observation_id":"6ba79372-4dc9-40b6-b274-2962f55a6f43","resolution":{"observed_at":"2026-08-15T18:10:41.602272Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.578243Z","title":"Pick-a-pic: An open dataset of user preferences for text-to-image generation, 2023","venue":null,"work_id":"14d0a666-d097-4d27-be8a-cdbaf6c1db19","year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.510212Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:7ab33d52c237c637bb90b8b216c675d88eaf401fd1db8c047bc49df64c03cf20","observation_id":"b03907d3-128d-40ba-a1ec-83d5f879e178","resolution":{"observed_at":"2026-08-15T18:10:41.584502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:40.515062Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.515062Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:678a522b424b6c11a3a75fe4abc2300611a5c2013d0a67c705ec06442ca8b14f","observation_id":"769f8c15-147b-48ac-a842-c86b9ff2fbc4","resolution":{"observed_at":"2026-08-15T18:10:40.515062Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.547813Z","title":"Blip: Bootstrapping language-image pre-training for unified vision-language understanding and genera- tion, 2022","venue":null,"work_id":"3bc721b8-68f8-4f3c-a1e2-bf9708e3e432","year":2022},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.519800Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:fdb8397be67c55ee08763710548938ba47098c0c0e12d561b7ac8a8ccb1a09b9","observation_id":"962c4e59-e349-49d5-bbb8-73578e36677c","resolution":{"observed_at":"2026-08-15T18:10:41.553444Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.528096Z","title":"Decoupled weight decay regularization, 2019","venue":null,"work_id":"cda2c0fb-6aee-4167-aeea-c44a52054346","year":2019},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.524752Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:fe1a8156f73d296d622a71daafffd331c3a8c179da6eb55e7a4467355243f10b","observation_id":"27b89135-56f9-46db-a087-232aad1431f1","resolution":{"observed_at":"2026-08-15T18:10:41.534131Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10797","last_updated":"2025-02-19T06:00:55Z","snapshot_observed_at":"2026-08-18T11:01:01.238927Z","submitted_at":"2024-06-16T03:45:45Z","title":"STAR: Scale-wise Text-conditioned AutoRegressive image generation","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10797","snapshot_observed_at":"2026-08-15T18:10:40.529384Z","title":"Star: Scale-wise text-conditioned autoregressive image gen- eration","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.529384Z"},"links":{"cited_paper":"/paper/2406.10797","citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:a5fe70f0c57dc9d83ff53c3c6ab6559344fe59b7c8afbdaea933f98a33277f42","observation_id":"b29cf751-bb8a-4257-99e8-1e612bbf9054","resolution":{"observed_at":"2026-08-15T18:10:40.529384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.512735Z","title":"Ai-powered image generation","venue":null,"work_id":"528f883a-07d1-43d8-997d-33ca58269579","year":2022},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.534273Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:81aa67fcc79c3f8e1b4fc7165dfd36628f3db35159986499e396439a0cda840b","observation_id":"83200ca7-2e13-4347-b0c5-e6898a0405df","resolution":{"observed_at":"2026-08-15T18:10:41.517387Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.494453Z","title":"Dynamic prompt opti- mizing for text-to-image generation","venue":null,"work_id":"bb16df5c-287e-4eed-8fc7-044e83db1a32","year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.538619Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:11917416b9ed46e5736701f8bd42a0196f48ca350e31520edbf02cfae3d83401","observation_id":"b11280d3-7343-4fee-975d-7a9eb3ea07bd","resolution":{"observed_at":"2026-08-15T18:10:41.500240Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.473594Z","title":"Uniform attention maps: Boost- ing image fidelity in reconstruction and editing","venue":null,"work_id":"b97abec7-781e-46fe-a4ab-bb6e35368527","year":2025},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.543362Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:555697eb77aa539e70f787ea365f58b89110fd654c3c4a04a45a8b39040123f7","observation_id":"e8803778-3333-4553-8a54-d807a2670e44","resolution":{"observed_at":"2026-08-15T18:10:41.483148Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.454614Z","title":"Dall-e 3","venue":null,"work_id":"efb010cf-f352-4342-b496-8a075f68dac1","year":null},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.547911Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:0ff7ea9c46df596b583ee77ed9aab901604ebdd4bb3151ad48d064d837560046","observation_id":"d534e4a4-2524-4258-b8bc-5d9eb2a3898b","resolution":{"observed_at":"2026-08-15T18:10:41.460934Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.436381Z","title":"Sdxl: Improving latent diffusion models for high-resolution image synthesis, 2023","venue":null,"work_id":"a361d54b-2e7d-4d82-ac5a-06fbb3ef566c","year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.552219Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:354eb3466ad68a683894122f2c53a2063d2cb64ad3a2c596eea53768498d9fde","observation_id":"0b8eccd2-f8c1-43e2-aff7-d8e44539e1c2","resolution":{"observed_at":"2026-08-15T18:10:41.442024Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.419253Z","title":"SDXL: Improving latent diffu- sion models for high-resolution image synthesis","venue":null,"work_id":"c871f0d9-3fee-49e0-9609-ebe87a16ff49","year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.556650Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:38a7d453fc4ec64dde9b9941b82b02728dd4ec2fd7f82a4c89f5e945150ab46e","observation_id":"8193513d-2967-403d-a725-2790736b0eea","resolution":{"observed_at":"2026-08-15T18:10:41.424727Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.16675","last_updated":"2025-05-22T13:40:00Z","snapshot_observed_at":"2026-08-17T12:52:24.699964Z","submitted_at":"2025-05-22T13:40:00Z","title":"On the Out-of-Distribution Generalization of Self-Supervised Learning","version":1},"cited_work":{"arxiv_id":"2505.16675","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.16675","snapshot_observed_at":"2026-08-15T18:10:40.909668Z","title":"On the Out-of-Distribution Generalization of Self-Supervised Learning","venue":"cs.LG","work_id":"35b050c3-d18f-4465-9ade-b335463a8854","year":2025},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.560832Z"},"links":{"cited_paper":"/paper/2505.16675","citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:67100284afe59753de9489151d554f4484f6f5537118fb7a9d02de6aa7cf5ab8","observation_id":"95a7d0a0-182c-4a33-b28d-115c0cc34527","resolution":{"observed_at":"2026-08-15T18:10:40.918364Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.401199Z","title":"Controlling text-to-image dif- fusion by orthogonal finetuning","venue":null,"work_id":"33635be1-1e46-4a7a-9d7a-4ff9fb4c5fca","year":null},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.565654Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:0332ca803bc56ebe8bdecca75fa060440961a09fa7e0dc70fb645c8bd7516b45","observation_id":"241ed4cb-ab43-4dea-a55d-a0486437b56a","resolution":{"observed_at":"2026-08-15T18:10:41.407480Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.382696Z","title":"Learning transferable visual models from natural lan- guage supervision","venue":null,"work_id":"951e5cd2-e504-477f-a7e4-a0a8684aab53","year":2021},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.570172Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:1758477bf10d5ac5aea1c315b4ff3ca63a861472ca6c7f35d498749f5914ff4b","observation_id":"0bcdf783-ac8e-4981-a232-0bb8671de2ce","resolution":{"observed_at":"2026-08-15T18:10:41.389597Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2204.06125","last_updated":"2022-04-13T01:10:33Z","snapshot_observed_at":"2026-08-15T12:50:58.405488Z","submitted_at":"2022-04-13T01:10:33Z","title":"Hierarchical Text-Conditional Image Generation with CLIP Latents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.06125","snapshot_observed_at":"2026-08-15T18:10:40.574563Z","title":"Hierarchical text- conditional image generation with CLIP latents","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.574563Z"},"links":{"cited_paper":"/paper/2204.06125","citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:edba93787c18e62208018997a85208a0c61193b15c666a6478b573bb4207edbf","observation_id":"2b1f95dc-063a-4ffd-a800-f2f3c10421d3","resolution":{"observed_at":"2026-08-15T18:10:40.574563Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.364361Z","title":"Hierarchical text-conditional image generation with clip latents, 2022","venue":null,"work_id":"26e9252d-06df-4960-a38b-2593d1811b78","year":2022},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.580724Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:2ea7561678699d11e359f4c44fe2cf91ab96212ff2c90ef45923aeb5032e9ba0","observation_id":"7d6ffa08-b8e5-4602-aee5-674fa98d7e3c","resolution":{"observed_at":"2026-08-15T18:10:41.369765Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.345653Z","title":"High- resolution image synthesis with latent diffusion mod- els","venue":null,"work_id":"1ef191a1-c85f-430b-bf63-d9010421b05e","year":2022},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.585426Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:0d0029341a2059f0f665367837cb47f0489d5d461bfbe27ff79e8a2348c0ca48","observation_id":"dab4b69e-e5f9-4cec-9630-fa4e0b6ea8e4","resolution":{"observed_at":"2026-08-15T18:10:41.352125Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.325560Z","title":"Dreambooth: Fine tuning text-to-image diffusion models for subject-driven generation","venue":null,"work_id":"b8b0c8df-d368-44dc-8512-0882f3b0a17e","year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.590126Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:8c08ce58bacdfc1aacb80c6d020a2c85c212c58b11c39d1e6531518522b41df6","observation_id":"cf220621-bdf2-43d0-96dc-b076c525382c","resolution":{"observed_at":"2026-08-15T18:10:41.332071Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.304569Z","title":"Improved techniques for training gans, 2016","venue":null,"work_id":"618b1cfc-911e-4d16-a298-08a561611e93","year":2016},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.595805Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:ecf0075d845d74fe083cd6cc5eeae9b575b398a18c1639c71143e05939667c99","observation_id":"20a1b4e6-b69d-4397-bef2-a9d881fe116a","resolution":{"observed_at":"2026-08-15T18:10:41.311385Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.288467Z","title":"Denoising diffusion implicit models","venue":null,"work_id":"49d61d8b-3ea9-4cde-b472-eadb4be3d283","year":2021},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.600643Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:ed215067dc727e1213640e674f5bc0c45a75ba1ded459229f50d22268b365130","observation_id":"9dfd8afa-b185-4e2c-ab3a-98d42ff09c6b","resolution":{"observed_at":"2026-08-15T18:10:41.293592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.15289","last_updated":"2025-06-17T11:43:46Z","snapshot_observed_at":"2026-08-20T02:55:12.742282Z","submitted_at":"2024-05-24T07:22:35Z","title":"Learning Invariant Causal Mechanism from Vision-Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.15289","snapshot_observed_at":"2026-08-15T18:10:40.605066Z","title":"Learning in- variant causal mechanism from vision-language mod- els","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.605066Z"},"links":{"cited_paper":"/paper/2405.15289","citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:abb4823fcb0167603ae782688105abc7e8b1d30d20bc15617fe20eff20760994","observation_id":"5a8ea8a6-9659-4053-90af-43ae343ced84","resolution":{"observed_at":"2026-08-15T18:10:40.605066Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.271097Z","title":"Joty, and Nikhil Naik","venue":null,"work_id":"b3fec42b-a5f8-49e0-ac71-a5bb16007196","year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.610095Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:0fde72dad4e841b1c30201347a85d7f232d4deeb86312c3f51e042c0e445fa1a","observation_id":"3124e14a-7a62-48c0-97d0-775dfbcce62f","resolution":{"observed_at":"2026-08-15T18:10:41.275936Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.02733","last_updated":"2024-04-04T19:42:32Z","snapshot_observed_at":"2026-08-16T14:03:57.454330Z","submitted_at":"2024-04-03T13:34:09Z","title":"InstantStyle: Free Lunch towards Style-Preserving in Text-to-Image Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.02733","snapshot_observed_at":"2026-08-15T18:10:40.614720Z","title":"Instantstyle: Free lunch towards style- preserving in text-to-image generation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.614720Z"},"links":{"cited_paper":"/paper/2404.02733","citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:38341dcc1932b94614bf9c0280e119897d5baacda102210d8a1618b507e5ad53","observation_id":"f93b09e9-603f-46ee-ae38-e2c1a3a7394e","resolution":{"observed_at":"2026-08-15T18:10:40.614720Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.254782Z","title":"Wang, Evan Montoya, David Munechika, Haoyang Yang, Benjamin Hoover, and Duen Horng Chau","venue":null,"work_id":"1a65beb6-2682-4f4a-88b5-d415cafffcbd","year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.620907Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:eb4df71ff619a4df4b5467a5fa0c07e2db053aec49b0f1e01769a48c457d16c3","observation_id":"df88327b-33c2-45e9-bbf3-ca87f60aef6e","resolution":{"observed_at":"2026-08-15T18:10:41.260491Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.09341","last_updated":"2023-09-25T08:19:23Z","snapshot_observed_at":"2026-08-14T02:46:25.655503Z","submitted_at":"2023-06-15T17:59:31Z","title":"Human Preference Score v2: A Solid Benchmark for Evaluating Human Preferences of Text-to-Image Synthesis","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.09341","snapshot_observed_at":"2026-08-15T18:10:40.625415Z","title":"Human prefer- ence score v2: A solid benchmark for evaluating human preferences of text-to-image synthesis","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.625415Z"},"links":{"cited_paper":"/paper/2306.09341","citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:5feb76e1f1a04be72d0eb22cb1302728a19e2ae2978378d5779284cc7e33c9f0","observation_id":"0493d612-f5c9-43b0-9d1b-be298253be9e","resolution":{"observed_at":"2026-08-15T18:10:40.625415Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.238704Z","title":"Human pref- erence score v2: A solid benchmark for evaluating hu- man preferences of text-to-image synthesis, 2023","venue":null,"work_id":"bababc0f-2d59-4313-92d1-3c889520aad4","year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.630277Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:1a87949c92b479a2acf0d927b025a831ccb341d15db192abb4308ac0464a2683","observation_id":"f548df4e-78e5-4707-941f-3686d8c975aa","resolution":{"observed_at":"2026-08-15T18:10:41.244031Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.220086Z","title":"Human preference score: Better align- ing text-to-image models with human preference","venue":null,"work_id":"44409269-4400-46f7-b667-daa04581440d","year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.635791Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:47f292fdd07f0ba6ba75152acc979568422990f7bc5101e2b2cd487d05137579","observation_id":"483df671-d50f-4372-b2e3-df941c890d18","resolution":{"observed_at":"2026-08-15T18:10:41.226892Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.200907Z","title":"Imagereward: Learning and evaluating human prefer- ences for text-to-image generation, 2023","venue":null,"work_id":"6dd020ad-35ad-42e0-80de-cb78a732652a","year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.641202Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:9d72991acd6b292b8095a832f502d869aee249d8e78fb7ebd132fe95cbf8298c","observation_id":"0872a63d-1ed5-4dfd-82a6-1b56bf6df2c7","resolution":{"observed_at":"2026-08-15T18:10:41.207570Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.181113Z","title":"Attngan: Fine-grained text to image generation with attentional generative adversarial networks","venue":null,"work_id":"d4fd3c99-7530-4752-908d-6909e63d7eaf","year":2018},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.646067Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:f242f9dbf949431497d993dfd72165aa6ba8ef0b91e3583dc4b90130dc57275c","observation_id":"f826b0fa-3e4c-49dd-97b9-dda28b8f5086","resolution":{"observed_at":"2026-08-15T18:10:41.188153Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.164879Z","title":"Synthesizing long-term human motions with diffusion models via coherent sampling","venue":null,"work_id":"387789ec-68d6-4ba5-a0d9-981d1d15484e","year":2023},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.650387Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:9673c438e5b04342ea1035484c179b0c960227e68b9bf1ad58f59e78db8e9a92","observation_id":"3cb89925-643e-4fa7-bf5d-78fdb20e8551","resolution":{"observed_at":"2026-08-15T18:10:41.170787Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.148953Z","title":"Scaling au- toregressive models for content-rich text-to-image gen- eration","venue":null,"work_id":"d2f5b3de-2639-48c3-9c2d-4623d806da32","year":null},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.654714Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:6402ae1d4056cc5e0109e79f6243c2c07e115c950ec3d581ea02b171430c1859","observation_id":"30f553e4-8861-4a70-8fc4-60c462c53ea9","resolution":{"observed_at":"2026-08-15T18:10:41.154490Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.07493","last_updated":"2025-03-10T16:12:50Z","snapshot_observed_at":"2026-08-16T12:51:27.176479Z","submitted_at":"2025-03-10T16:12:50Z","title":"V2Flow: Unifying Visual Tokenization and Large Language Model Vocabularies for Autoregressive Image Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.07493","snapshot_observed_at":"2026-08-15T18:10:40.664888Z","title":"V2flow: Unifying visual tokenization and large language model vocabularies for autoregressive image generation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.664888Z"},"links":{"cited_paper":"/paper/2503.07493","citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:116d520b2613943a88d0baef9207ce44046d2c13aa50bc56165e59daf8d81e48","observation_id":"b7b3b827-0a92-4d26-9640-1371687f615d","resolution":{"observed_at":"2026-08-15T18:10:40.664888Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.111323Z","title":"Rethinking misalignment in vision-language model adaptation from a causal perspective","venue":null,"work_id":"4c1f2f9a-8b2e-41e4-b718-40f6e770db6f","year":null},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.669545Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:24712eb9022bd62325b003d988ebdccd9816d9a7007df8363d19d2890d173fed","observation_id":"f214499b-233e-4606-bc11-1e07807ea24d","resolution":{"observed_at":"2026-08-15T18:10:41.118053Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:40.674628Z","title":"Aligning few-step diffusion models with dense reward difference learning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.674628Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:2a7fbf2574cf93709392733ca839c1dae0a1d73e923e6d2e7a7637d54ee6483a","observation_id":"5cad291d-6381-4c09-b2d0-fe88c68470be","resolution":{"observed_at":"2026-08-15T18:10:40.674628Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.092703Z","title":"Llada 1.5: Variance-reduced preference optimization for large lan- guage diffusion models, 2025","venue":null,"work_id":"0d887703-78d5-4304-b9d8-6e7acfc746b7","year":2025},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.679135Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:5ded4773f7fe3026251bd4e8833e58e2f5452982958df1a95016954593436be7","observation_id":"70538c97-c069-487a-9c71-2b78efb41ac8","resolution":{"observed_at":"2026-08-15T18:10:41.098371Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.074543Z","title":"e” to denote images generated with base prompts and “r","venue":null,"work_id":"d2a2844a-5763-4f31-94da-f47509ac0070","year":2019},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.683699Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:c2f0e44e66ecd76fe1a0e2f2043c1759adf0a647b618615e641e20a733b00f2f","observation_id":"3946fe8e-d1f3-48bc-a628-e0c27dff35a4","resolution":{"observed_at":"2026-08-15T18:10:41.080695Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.058041Z","title":null,"venue":null,"work_id":"214e5e8b-e902-41ff-b399-c99cadbaca70","year":null},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.689539Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:61dfc9c41d917d9e38f376a54908f6e3eb281bdda49fbcbb701d6d9d1d95df72","observation_id":"1357ebee-c7cb-4df7-80e2-e95756f7e3fd","resolution":{"observed_at":"2026-08-15T18:10:41.062701Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.040304Z","title":null,"venue":null,"work_id":"7264c074-b56c-41d9-b1e8-06f02089acba","year":null},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.694092Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:cdd8e0c83381537d7d41b5140cd1d23e56694f6f0c9f99b03ce36053e5c295e9","observation_id":"1977e587-5552-4101-8cc1-72b810e41b5a","resolution":{"observed_at":"2026-08-15T18:10:41.045459Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.017134Z","title":"simple black hexagon, single geometric shape","venue":null,"work_id":"a0902e51-a48b-479b-bc12-b70bb60f4840","year":null},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.699063Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:3cdc9f51d1276058025598c42b6ab8fa58fdc9f76b3d3174f88504ee044dfe49","observation_id":"cff9ba47-debb-48b0-ae58-a27d5c56c1be","resolution":{"observed_at":"2026-08-15T18:10:41.028078Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:10:41.131973Z","title":null,"venue":null,"work_id":"fb70876f-80b9-4965-8209-50a5b427cfa4","year":null},"citing_paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-15T18:10:40.660408Z"},"links":{"citing_paper":"/paper/2507.19002"},"observation_digest":"sha256:2105a5f77005434f3d1e0f46ecce8897decf575e0b75cd92db2fd39e36bfad53","observation_id":"2783e9ae-4894-4ff8-9a5e-188445dae9a6","resolution":{"observed_at":"2026-08-15T18:10:41.137474Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2507.19002","last_updated":"2025-07-25T07:01:50Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-18T03:44:47.640284Z","submitted_at":"2025-07-25T07:01:50Z","title":"Enhancing Reward Models for High-quality Image Generation: Beyond Text-Image Alignment"},"reference_resolution":{"displayed":52,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":13,"verified_exact":2,"verified_fuzzy":37},"total_outbound_references":52},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 21 August 2026, this Paper Citation Record lists 52 of 52 outbound references and 2 inbound Pith citation observations for arXiv:2507.19002."}