{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:SGHWTYXQH5KNG37HZQTXKPNY6C","short_pith_number":"pith:SGHWTYXQ","schema_version":"1.0","canonical_sha256":"918f69e2f03f54d36fe7cc27753db8f0b7c22adeae5a88237036bbcdae8704ac","source":{"kind":"arxiv","id":"2411.07199","version":2},"attestation_state":"computed","paper":{"title":"OmniEdit: Building Image Editing Generalist Models Through Specialist Supervision","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CV","authors_text":"Cong Wei, Ge Zhang, Weiming Ren, Wenhu Chen, Xinrun Du, Zheyang Xiong","submitted_at":"2024-11-11T18:21:43Z","abstract_excerpt":"Instruction-guided image editing methods have demonstrated significant potential by training diffusion models on automatically synthesized or manually annotated image editing pairs. However, these methods remain far from practical, real-life applications. We identify three primary challenges contributing to this gap. Firstly, existing models have limited editing skills due to the biased synthesis process. Secondly, these methods are trained with datasets with a high volume of noise and artifacts. This is due to the application of simple filtering methods like CLIP-score. Thirdly, all these dat"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2411.07199","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2024-11-11T18:21:43Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"4bff5afac04c6f15279f02c5965657a1b54181e84a5e8f840757e4c2b3decd71","abstract_canon_sha256":"8d4f96ca33a1d898387797f91e3b5b8e076be6bbe460addb39edcd6d54570933"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T10:54:39.251444Z","signature_b64":"dLTJak+Fzlv7eJM1vgV0DkaibgBXDDYUtGQhhOZBTCpcg7yLYuoPKNUXcILUDKZuO2JmZ8pLnlyAygCGcBrMDg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"918f69e2f03f54d36fe7cc27753db8f0b7c22adeae5a88237036bbcdae8704ac","last_reissued_at":"2026-07-05T10:54:39.250939Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T10:54:39.250939Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"OmniEdit: Building Image Editing Generalist Models Through Specialist Supervision","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CV","authors_text":"Cong Wei, Ge Zhang, Weiming Ren, Wenhu Chen, Xinrun Du, Zheyang Xiong","submitted_at":"2024-11-11T18:21:43Z","abstract_excerpt":"Instruction-guided image editing methods have demonstrated significant potential by training diffusion models on automatically synthesized or manually annotated image editing pairs. However, these methods remain far from practical, real-life applications. We identify three primary challenges contributing to this gap. Firstly, existing models have limited editing skills due to the biased synthesis process. Secondly, these methods are trained with datasets with a high volume of noise and artifacts. This is due to the application of simple filtering methods like CLIP-score. Thirdly, all these dat"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2411.07199","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2411.07199/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2411.07199","created_at":"2026-07-05T10:54:39.251002+00:00"},{"alias_kind":"arxiv_version","alias_value":"2411.07199v2","created_at":"2026-07-05T10:54:39.251002+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2411.07199","created_at":"2026-07-05T10:54:39.251002+00:00"},{"alias_kind":"pith_short_12","alias_value":"SGHWTYXQH5KN","created_at":"2026-07-05T10:54:39.251002+00:00"},{"alias_kind":"pith_short_16","alias_value":"SGHWTYXQH5KNG37H","created_at":"2026-07-05T10:54:39.251002+00:00"},{"alias_kind":"pith_short_8","alias_value":"SGHWTYXQ","created_at":"2026-07-05T10:54:39.251002+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":8,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2607.06445","citing_title":"Analysis-by-Proxy: Localization Signals in VLMs Operating as Condition Encoders","ref_index":25,"is_internal_anchor":true},{"citing_arxiv_id":"2606.04264","citing_title":"UniCanvas: A Diffusion-base Unified Model for Text-in-Image Joint Generation","ref_index":53,"is_internal_anchor":false},{"citing_arxiv_id":"2605.09233","citing_title":"Towards Robust Sequential Decomposition for Complex Image Editing","ref_index":36,"is_internal_anchor":false},{"citing_arxiv_id":"2504.13109","citing_title":"UniEdit-Flow: Unleashing Inversion and Editing in the Era of Flow Models","ref_index":74,"is_internal_anchor":false},{"citing_arxiv_id":"2504.20690","citing_title":"In-Context Edit: Enabling Instructional Image Editing with In-Context Generation in Large Scale Diffusion Transformer","ref_index":34,"is_internal_anchor":false},{"citing_arxiv_id":"2505.20275","citing_title":"ImgEdit: A Unified Image Editing Dataset and Benchmark","ref_index":73,"is_internal_anchor":false},{"citing_arxiv_id":"2605.09233","citing_title":"Towards Robust Sequential Decomposition for Complex Image Editing","ref_index":36,"is_internal_anchor":false},{"citing_arxiv_id":"2504.17761","citing_title":"Step1X-Edit: A Practical Framework for General Image Editing","ref_index":59,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/SGHWTYXQH5KNG37HZQTXKPNY6C","json":"https://pith.science/pith/SGHWTYXQH5KNG37HZQTXKPNY6C.json","graph_json":"https://pith.science/api/pith-number/SGHWTYXQH5KNG37HZQTXKPNY6C/graph.json","events_json":"https://pith.science/api/pith-number/SGHWTYXQH5KNG37HZQTXKPNY6C/events.json","paper":"https://pith.science/paper/SGHWTYXQ"},"agent_actions":{"view_html":"https://pith.science/pith/SGHWTYXQH5KNG37HZQTXKPNY6C","download_json":"https://pith.science/pith/SGHWTYXQH5KNG37HZQTXKPNY6C.json","view_paper":"https://pith.science/paper/SGHWTYXQ","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2411.07199&json=true","fetch_graph":"https://pith.science/api/pith-number/SGHWTYXQH5KNG37HZQTXKPNY6C/graph.json","fetch_events":"https://pith.science/api/pith-number/SGHWTYXQH5KNG37HZQTXKPNY6C/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/SGHWTYXQH5KNG37HZQTXKPNY6C/action/timestamp_anchor","attest_storage":"https://pith.science/pith/SGHWTYXQH5KNG37HZQTXKPNY6C/action/storage_attestation","attest_author":"https://pith.science/pith/SGHWTYXQH5KNG37HZQTXKPNY6C/action/author_attestation","sign_citation":"https://pith.science/pith/SGHWTYXQH5KNG37HZQTXKPNY6C/action/citation_signature","submit_replication":"https://pith.science/pith/SGHWTYXQH5KNG37HZQTXKPNY6C/action/replication_record"}},"created_at":"2026-07-05T10:54:39.251002+00:00","updated_at":"2026-07-05T10:54:39.251002+00:00"}