{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:RZWALUTGVRIBVKIBFRC7SKCPIJ","short_pith_number":"pith:RZWALUTG","schema_version":"1.0","canonical_sha256":"8e6c05d266ac501aa9012c45f9284f4258a2d5ffd84eb4a1f0491568c472f593","source":{"kind":"arxiv","id":"2406.00888","version":2},"attestation_state":"computed","paper":{"title":"Aligning Language Models with Demonstrated Feedback","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.HC"],"primary_cat":"cs.CL","authors_text":"Diyi Yang, Hyundong Cho, Joey Hejna, Michael S. Bernstein, Michelle S. Lam, Omar Shaikh, Yijia Shao","submitted_at":"2024-06-02T23:13:56Z","abstract_excerpt":"Language models are aligned to emulate the collective voice of many, resulting in outputs that align with no one in particular. Steering LLMs away from generic output is possible through supervised finetuning or RLHF, but requires prohibitively large datasets for new ad-hoc tasks. We argue that it is instead possible to align an LLM to a specific setting by leveraging a very small number (< 10) of demonstrations as feedback. Our method, Demonstration ITerated Task Optimization (DITTO), directly aligns language model outputs to a user's demonstrated behaviors. Derived using ideas from online im"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2406.00888","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CL","submitted_at":"2024-06-02T23:13:56Z","cross_cats_sorted":["cs.HC"],"title_canon_sha256":"3c6b61cd868acd48f84f9938ed35d2191ce9af05a14d759d44f6428999103ee1","abstract_canon_sha256":"abfb4b5d5e59f296bec1e13ea31c22b67661c82c227a42bff753e54c4544dcf1"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T10:51:03.804674Z","signature_b64":"0rg0UfjKlov7l5hXciBHUR7N1dvtDxN0h1QNskR86ddFVY9OIqDp49CIvfrKZldhJi1U841onQn/tjQetAwIDA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"8e6c05d266ac501aa9012c45f9284f4258a2d5ffd84eb4a1f0491568c472f593","last_reissued_at":"2026-07-05T10:51:03.804182Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T10:51:03.804182Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Aligning Language Models with Demonstrated Feedback","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.HC"],"primary_cat":"cs.CL","authors_text":"Diyi Yang, Hyundong Cho, Joey Hejna, Michael S. Bernstein, Michelle S. Lam, Omar Shaikh, Yijia Shao","submitted_at":"2024-06-02T23:13:56Z","abstract_excerpt":"Language models are aligned to emulate the collective voice of many, resulting in outputs that align with no one in particular. Steering LLMs away from generic output is possible through supervised finetuning or RLHF, but requires prohibitively large datasets for new ad-hoc tasks. We argue that it is instead possible to align an LLM to a specific setting by leveraging a very small number (< 10) of demonstrations as feedback. Our method, Demonstration ITerated Task Optimization (DITTO), directly aligns language model outputs to a user's demonstrated behaviors. Derived using ideas from online im"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2406.00888","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2406.00888/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2406.00888","created_at":"2026-07-05T10:51:03.804238+00:00"},{"alias_kind":"arxiv_version","alias_value":"2406.00888v2","created_at":"2026-07-05T10:51:03.804238+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2406.00888","created_at":"2026-07-05T10:51:03.804238+00:00"},{"alias_kind":"pith_short_12","alias_value":"RZWALUTGVRIB","created_at":"2026-07-05T10:51:03.804238+00:00"},{"alias_kind":"pith_short_16","alias_value":"RZWALUTGVRIBVKIB","created_at":"2026-07-05T10:51:03.804238+00:00"},{"alias_kind":"pith_short_8","alias_value":"RZWALUTG","created_at":"2026-07-05T10:51:03.804238+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2605.30323","citing_title":"In-Context Reward Adaptation for Robust Preference Modeling","ref_index":11,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/RZWALUTGVRIBVKIBFRC7SKCPIJ","json":"https://pith.science/pith/RZWALUTGVRIBVKIBFRC7SKCPIJ.json","graph_json":"https://pith.science/api/pith-number/RZWALUTGVRIBVKIBFRC7SKCPIJ/graph.json","events_json":"https://pith.science/api/pith-number/RZWALUTGVRIBVKIBFRC7SKCPIJ/events.json","paper":"https://pith.science/paper/RZWALUTG"},"agent_actions":{"view_html":"https://pith.science/pith/RZWALUTGVRIBVKIBFRC7SKCPIJ","download_json":"https://pith.science/pith/RZWALUTGVRIBVKIBFRC7SKCPIJ.json","view_paper":"https://pith.science/paper/RZWALUTG","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2406.00888&json=true","fetch_graph":"https://pith.science/api/pith-number/RZWALUTGVRIBVKIBFRC7SKCPIJ/graph.json","fetch_events":"https://pith.science/api/pith-number/RZWALUTGVRIBVKIBFRC7SKCPIJ/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/RZWALUTGVRIBVKIBFRC7SKCPIJ/action/timestamp_anchor","attest_storage":"https://pith.science/pith/RZWALUTGVRIBVKIBFRC7SKCPIJ/action/storage_attestation","attest_author":"https://pith.science/pith/RZWALUTGVRIBVKIBFRC7SKCPIJ/action/author_attestation","sign_citation":"https://pith.science/pith/RZWALUTGVRIBVKIBFRC7SKCPIJ/action/citation_signature","submit_replication":"https://pith.science/pith/RZWALUTGVRIBVKIBFRC7SKCPIJ/action/replication_record"}},"created_at":"2026-07-05T10:51:03.804238+00:00","updated_at":"2026-07-05T10:51:03.804238+00:00"}