{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:VAZTYQXIJ4ABWPT5J2ZWEZGOWB","short_pith_number":"pith:VAZTYQXI","schema_version":"1.0","canonical_sha256":"a8333c42e84f001b3e7d4eb36264ceb057a6c931d8094f19a31a4e834a80f1ae","source":{"kind":"arxiv","id":"2311.02147","version":1},"attestation_state":"computed","paper":{"title":"The Alignment Problem in Context","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Rapha\\\"el Milli\\`ere","submitted_at":"2023-11-03T17:57:55Z","abstract_excerpt":"A core challenge in the development of increasingly capable AI systems is to make them safe and reliable by ensuring their behaviour is consistent with human values. This challenge, known as the alignment problem, does not merely apply to hypothetical future AI systems that may pose catastrophic risks; it already applies to current systems, such as large language models, whose potential for harm is rapidly increasing. In this paper, I assess whether we are on track to solve the alignment problem for large language models, and what that means for the safety of future AI systems. I argue that ex"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2311.02147","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2023-11-03T17:57:55Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"85ad2dbe49f6d3d68c503c6f1e8a3ee6ca68d1210aad53110daecc49dc82db86","abstract_canon_sha256":"f7b03cb914db52d0dc38faf3110f7302a361987a2272cfe552fcf6ff666f44ab"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T07:09:06.085733Z","signature_b64":"Q3V+Bu8UgIiWbx4EolNmWhynkcPhX93t14nCN9OnMPdryVG9pwJ8pFV2BogtkH/K+fSRl7MifDdPn5k7UWvpAw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"a8333c42e84f001b3e7d4eb36264ceb057a6c931d8094f19a31a4e834a80f1ae","last_reissued_at":"2026-07-05T07:09:06.085260Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T07:09:06.085260Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"The Alignment Problem in Context","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Rapha\\\"el Milli\\`ere","submitted_at":"2023-11-03T17:57:55Z","abstract_excerpt":"A core challenge in the development of increasingly capable AI systems is to make them safe and reliable by ensuring their behaviour is consistent with human values. This challenge, known as the alignment problem, does not merely apply to hypothetical future AI systems that may pose catastrophic risks; it already applies to current systems, such as large language models, whose potential for harm is rapidly increasing. In this paper, I assess whether we are on track to solve the alignment problem for large language models, and what that means for the safety of future AI systems. I argue that ex"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2311.02147","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2311.02147/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2311.02147","created_at":"2026-07-05T07:09:06.085321+00:00"},{"alias_kind":"arxiv_version","alias_value":"2311.02147v1","created_at":"2026-07-05T07:09:06.085321+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2311.02147","created_at":"2026-07-05T07:09:06.085321+00:00"},{"alias_kind":"pith_short_12","alias_value":"VAZTYQXIJ4AB","created_at":"2026-07-05T07:09:06.085321+00:00"},{"alias_kind":"pith_short_16","alias_value":"VAZTYQXIJ4ABWPT5","created_at":"2026-07-05T07:09:06.085321+00:00"},{"alias_kind":"pith_short_8","alias_value":"VAZTYQXI","created_at":"2026-07-05T07:09:06.085321+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2501.11739","citing_title":"Episodic memory in AI agents poses risks that should be studied and mitigated","ref_index":101,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/VAZTYQXIJ4ABWPT5J2ZWEZGOWB","json":"https://pith.science/pith/VAZTYQXIJ4ABWPT5J2ZWEZGOWB.json","graph_json":"https://pith.science/api/pith-number/VAZTYQXIJ4ABWPT5J2ZWEZGOWB/graph.json","events_json":"https://pith.science/api/pith-number/VAZTYQXIJ4ABWPT5J2ZWEZGOWB/events.json","paper":"https://pith.science/paper/VAZTYQXI"},"agent_actions":{"view_html":"https://pith.science/pith/VAZTYQXIJ4ABWPT5J2ZWEZGOWB","download_json":"https://pith.science/pith/VAZTYQXIJ4ABWPT5J2ZWEZGOWB.json","view_paper":"https://pith.science/paper/VAZTYQXI","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2311.02147&json=true","fetch_graph":"https://pith.science/api/pith-number/VAZTYQXIJ4ABWPT5J2ZWEZGOWB/graph.json","fetch_events":"https://pith.science/api/pith-number/VAZTYQXIJ4ABWPT5J2ZWEZGOWB/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/VAZTYQXIJ4ABWPT5J2ZWEZGOWB/action/timestamp_anchor","attest_storage":"https://pith.science/pith/VAZTYQXIJ4ABWPT5J2ZWEZGOWB/action/storage_attestation","attest_author":"https://pith.science/pith/VAZTYQXIJ4ABWPT5J2ZWEZGOWB/action/author_attestation","sign_citation":"https://pith.science/pith/VAZTYQXIJ4ABWPT5J2ZWEZGOWB/action/citation_signature","submit_replication":"https://pith.science/pith/VAZTYQXIJ4ABWPT5J2ZWEZGOWB/action/replication_record"}},"created_at":"2026-07-05T07:09:06.085321+00:00","updated_at":"2026-07-05T07:09:06.085321+00:00"}