{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:S6N3SDGZ2LBGNM3WRN2DJ4WXWO","short_pith_number":"pith:S6N3SDGZ","schema_version":"1.0","canonical_sha256":"979bb90cd9d2c266b3768b7434f2d7b3890c83693290eb935515d3d4e3459092","source":{"kind":"arxiv","id":"2310.00158","version":2},"attestation_state":"computed","paper":{"title":"Feedback-guided Data Synthesis for Imbalanced Classification","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI","cs.LG"],"primary_cat":"cs.CV","authors_text":"Adriana Romero-Soriano, Florian Bordes, Michal Drozdzal, Mohammad Pezeshki, Reyhane Askari Hemmat","submitted_at":"2023-09-29T21:47:57Z","abstract_excerpt":"Current status quo in machine learning is to use static datasets of real images for training, which often come from long-tailed distributions. With the recent advances in generative models, researchers have started augmenting these static datasets with synthetic data, reporting moderate performance improvements on classification tasks. We hypothesize that these performance gains are limited by the lack of feedback from the classifier to the generative model, which would promote the usefulness of the generated samples to improve the classifier's performance. In this work, we introduce a framewo"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2310.00158","kind":"arxiv","version":2},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CV","submitted_at":"2023-09-29T21:47:57Z","cross_cats_sorted":["cs.AI","cs.LG"],"title_canon_sha256":"9fffcec8a11d128c90b1fca51314d9ffe26a6a2dfa08e5bd0ae1a1d7e84d55e1","abstract_canon_sha256":"a25d52d983f220ab0fab5013736caba4b4f941e3b301959c1d95273cf648d99f"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T09:04:56.170311Z","signature_b64":"8DI4hnE/sff3j+dH9DNeam5lyQcDHcWabPQqieRslf3emIgN2gzqg/ULA+bbCNtZOLUhfZceRKmZiO2srX/5Cg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"979bb90cd9d2c266b3768b7434f2d7b3890c83693290eb935515d3d4e3459092","last_reissued_at":"2026-07-05T09:04:56.169820Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T09:04:56.169820Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Feedback-guided Data Synthesis for Imbalanced Classification","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI","cs.LG"],"primary_cat":"cs.CV","authors_text":"Adriana Romero-Soriano, Florian Bordes, Michal Drozdzal, Mohammad Pezeshki, Reyhane Askari Hemmat","submitted_at":"2023-09-29T21:47:57Z","abstract_excerpt":"Current status quo in machine learning is to use static datasets of real images for training, which often come from long-tailed distributions. With the recent advances in generative models, researchers have started augmenting these static datasets with synthetic data, reporting moderate performance improvements on classification tasks. We hypothesize that these performance gains are limited by the lack of feedback from the classifier to the generative model, which would promote the usefulness of the generated samples to improve the classifier's performance. In this work, we introduce a framewo"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2310.00158","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2310.00158/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2310.00158","created_at":"2026-07-05T09:04:56.169881+00:00"},{"alias_kind":"arxiv_version","alias_value":"2310.00158v2","created_at":"2026-07-05T09:04:56.169881+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2310.00158","created_at":"2026-07-05T09:04:56.169881+00:00"},{"alias_kind":"pith_short_12","alias_value":"S6N3SDGZ2LBG","created_at":"2026-07-05T09:04:56.169881+00:00"},{"alias_kind":"pith_short_16","alias_value":"S6N3SDGZ2LBGNM3W","created_at":"2026-07-05T09:04:56.169881+00:00"},{"alias_kind":"pith_short_8","alias_value":"S6N3SDGZ","created_at":"2026-07-05T09:04:56.169881+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":3,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2605.26353","citing_title":"Personalized Generative Models for Contextual Debiasing","ref_index":23,"is_internal_anchor":false},{"citing_arxiv_id":"2605.11231","citing_title":"LiBaGS: Lightweight Boundary Gap Synthesis for Targeted Synthetic Data Selection","ref_index":18,"is_internal_anchor":false},{"citing_arxiv_id":"2605.11231","citing_title":"LiBaGS: Lightweight Boundary Gap Synthesis for Targeted Synthetic Data Selection","ref_index":18,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/S6N3SDGZ2LBGNM3WRN2DJ4WXWO","json":"https://pith.science/pith/S6N3SDGZ2LBGNM3WRN2DJ4WXWO.json","graph_json":"https://pith.science/api/pith-number/S6N3SDGZ2LBGNM3WRN2DJ4WXWO/graph.json","events_json":"https://pith.science/api/pith-number/S6N3SDGZ2LBGNM3WRN2DJ4WXWO/events.json","paper":"https://pith.science/paper/S6N3SDGZ"},"agent_actions":{"view_html":"https://pith.science/pith/S6N3SDGZ2LBGNM3WRN2DJ4WXWO","download_json":"https://pith.science/pith/S6N3SDGZ2LBGNM3WRN2DJ4WXWO.json","view_paper":"https://pith.science/paper/S6N3SDGZ","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2310.00158&json=true","fetch_graph":"https://pith.science/api/pith-number/S6N3SDGZ2LBGNM3WRN2DJ4WXWO/graph.json","fetch_events":"https://pith.science/api/pith-number/S6N3SDGZ2LBGNM3WRN2DJ4WXWO/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/S6N3SDGZ2LBGNM3WRN2DJ4WXWO/action/timestamp_anchor","attest_storage":"https://pith.science/pith/S6N3SDGZ2LBGNM3WRN2DJ4WXWO/action/storage_attestation","attest_author":"https://pith.science/pith/S6N3SDGZ2LBGNM3WRN2DJ4WXWO/action/author_attestation","sign_citation":"https://pith.science/pith/S6N3SDGZ2LBGNM3WRN2DJ4WXWO/action/citation_signature","submit_replication":"https://pith.science/pith/S6N3SDGZ2LBGNM3WRN2DJ4WXWO/action/replication_record"}},"created_at":"2026-07-05T09:04:56.169881+00:00","updated_at":"2026-07-05T09:04:56.169881+00:00"}