{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:4SSTXQAMEJ6Z5YKB32SBUNOPRR","short_pith_number":"pith:4SSTXQAM","schema_version":"1.0","canonical_sha256":"e4a53bc00c227d9ee141dea41a35cf8c60e8275c75bc8e64341a46a9b40ecd6b","source":{"kind":"arxiv","id":"2412.13678","version":1},"attestation_state":"computed","paper":{"title":"Clio: Privacy-Preserving Insights into Real-World AI Use","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI","cs.CL","cs.CR","cs.LG"],"primary_cat":"cs.CY","authors_text":"Alex Tamkin, Alfred Mountfield, Ankur Rathi, Brian Clarke, Deep Ganguli, Esin Durmus, Jack Clark, Jared Kaplan, Jared Mueller, Jerry Hong, Kunal Handa, Landon Goldberg, Liane Lovitt, Michael Stern, Miles McCain, Saffron Huang, Shan Carter, Stuart Ritchie, Theodore R. Sumers, Wes Mitchell, William McEachen","submitted_at":"2024-12-18T10:05:43Z","abstract_excerpt":"How are AI assistants being used in the real world? While model providers in theory have a window into this impact via their users' data, both privacy concerns and practical challenges have made analyzing this data difficult. To address these issues, we present Clio (Claude insights and observations), a privacy-preserving platform that uses AI assistants themselves to analyze and surface aggregated usage patterns across millions of conversations, without the need for human reviewers to read raw conversations. We validate this can be done with a high degree of accuracy and privacy by conducting"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2412.13678","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CY","submitted_at":"2024-12-18T10:05:43Z","cross_cats_sorted":["cs.AI","cs.CL","cs.CR","cs.LG"],"title_canon_sha256":"f5ac34ec97f376533b5f18a9a46bf44373761db8c8ddece889f72a4dd3ff3087","abstract_canon_sha256":"3209121e769c902bd174311a501b1c4041a91ac7d3ac9ade178cd7cd2f11fcbb"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T09:51:06.384792Z","signature_b64":"20mXf7CYWDS/QnFcm08FvBJ3jFXUfaS7RtZPlaKGa31nA9aPD55yRIo5j7wvFHzu1TTjhMXGFhTpamtyAmdzCQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"e4a53bc00c227d9ee141dea41a35cf8c60e8275c75bc8e64341a46a9b40ecd6b","last_reissued_at":"2026-07-05T09:51:06.384268Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T09:51:06.384268Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Clio: Privacy-Preserving Insights into Real-World AI Use","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI","cs.CL","cs.CR","cs.LG"],"primary_cat":"cs.CY","authors_text":"Alex Tamkin, Alfred Mountfield, Ankur Rathi, Brian Clarke, Deep Ganguli, Esin Durmus, Jack Clark, Jared Kaplan, Jared Mueller, Jerry Hong, Kunal Handa, Landon Goldberg, Liane Lovitt, Michael Stern, Miles McCain, Saffron Huang, Shan Carter, Stuart Ritchie, Theodore R. Sumers, Wes Mitchell, William McEachen","submitted_at":"2024-12-18T10:05:43Z","abstract_excerpt":"How are AI assistants being used in the real world? While model providers in theory have a window into this impact via their users' data, both privacy concerns and practical challenges have made analyzing this data difficult. To address these issues, we present Clio (Claude insights and observations), a privacy-preserving platform that uses AI assistants themselves to analyze and surface aggregated usage patterns across millions of conversations, without the need for human reviewers to read raw conversations. We validate this can be done with a high degree of accuracy and privacy by conducting"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2412.13678","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2412.13678/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2412.13678","created_at":"2026-07-05T09:51:06.384338+00:00"},{"alias_kind":"arxiv_version","alias_value":"2412.13678v1","created_at":"2026-07-05T09:51:06.384338+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2412.13678","created_at":"2026-07-05T09:51:06.384338+00:00"},{"alias_kind":"pith_short_12","alias_value":"4SSTXQAMEJ6Z","created_at":"2026-07-05T09:51:06.384338+00:00"},{"alias_kind":"pith_short_16","alias_value":"4SSTXQAMEJ6Z5YKB","created_at":"2026-07-05T09:51:06.384338+00:00"},{"alias_kind":"pith_short_8","alias_value":"4SSTXQAM","created_at":"2026-07-05T09:51:06.384338+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":21,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.05750","citing_title":"Three Years of r/ChatGPT: Societal Impact Evaluations from Social Media Data","ref_index":16,"is_internal_anchor":false},{"citing_arxiv_id":"2605.31593","citing_title":"Stateful Online Monitoring Catches Distributed Agent Attacks","ref_index":26,"is_internal_anchor":false},{"citing_arxiv_id":"2605.30848","citing_title":"LLM Anonymization Against Agentic Re-Identification","ref_index":39,"is_internal_anchor":false},{"citing_arxiv_id":"2605.29018","citing_title":"Adopt $\\neq$ Adapt: Longitudinal Analyses of LLM Conversations in the Wild","ref_index":9,"is_internal_anchor":false},{"citing_arxiv_id":"2412.14855","citing_title":"Position: Mind the Gap-AI Security and the Limits of Current Reporting Standards","ref_index":107,"is_internal_anchor":false},{"citing_arxiv_id":"2504.15801","citing_title":"A closer look at how large language models trust humans: patterns and biases","ref_index":13,"is_internal_anchor":false},{"citing_arxiv_id":"2505.14549","citing_title":"Can Large Language Models Really Recognize Your Name?","ref_index":70,"is_internal_anchor":false},{"citing_arxiv_id":"2605.22564","citing_title":"SynAE: A Framework for Measuring the Quality of Synthetic Data for Tool-Calling Agent Evaluations","ref_index":31,"is_internal_anchor":false},{"citing_arxiv_id":"2605.19190","citing_title":"Going PLACES: Participatory Localized Red Teaming for Text-to-Image Safety in the Global South","ref_index":91,"is_internal_anchor":false},{"citing_arxiv_id":"2506.06414","citing_title":"Benchmarking Misuse Mitigation Against Covert Adversaries","ref_index":39,"is_internal_anchor":false},{"citing_arxiv_id":"2509.11206","citing_title":"Evalet: Evaluating Large Language Models through Functional Fragmentation","ref_index":82,"is_internal_anchor":false},{"citing_arxiv_id":"2509.21267","citing_title":"Task-Dependent Evaluation of LLM Output Homogenization: A Taxonomy-Guided Framework","ref_index":27,"is_internal_anchor":false},{"citing_arxiv_id":"2603.03295","citing_title":"Language Model Goal Selection Differs from Humans' in a Self-Directed Learning Task","ref_index":19,"is_internal_anchor":false},{"citing_arxiv_id":"2604.15340","citing_title":"Restoration, Exploration and Transformation: How Youth Engage Character.AI Chatbots for Feels, Fun and Finding themselves","ref_index":75,"is_internal_anchor":false},{"citing_arxiv_id":"2605.11388","citing_title":"Deep Reasoning in General Purpose Agents via Structured Meta-Cognition","ref_index":161,"is_internal_anchor":false},{"citing_arxiv_id":"2604.23842","citing_title":"Reheat Nachos for Dinner? Evaluating AI Support for Cross-Cultural Communication of Neologisms","ref_index":7,"is_internal_anchor":false},{"citing_arxiv_id":"2605.05767","citing_title":"Priming, Path-dependence, and Plasticity: Understanding the molding of user-LLM interaction and its implications from (many) chat logs in the wild","ref_index":93,"is_internal_anchor":false},{"citing_arxiv_id":"2605.01104","citing_title":"RECAP: An End-to-End Platform for Capturing, Replaying, and Analyzing AI-Assisted Programming Interactions","ref_index":26,"is_internal_anchor":false},{"citing_arxiv_id":"2604.21769","citing_title":"Who Defines \"Best\"? Towards Interactive, User-Defined Evaluation of LLM Leaderboards","ref_index":57,"is_internal_anchor":false},{"citing_arxiv_id":"2604.20720","citing_title":"COMPASS: COntinual Multilingual PEFT with Adaptive Semantic Sampling","ref_index":68,"is_internal_anchor":false},{"citing_arxiv_id":"2604.07629","citing_title":"Behavior Latticing: Inferring User Motivations from Unstructured Interactions","ref_index":96,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/4SSTXQAMEJ6Z5YKB32SBUNOPRR","json":"https://pith.science/pith/4SSTXQAMEJ6Z5YKB32SBUNOPRR.json","graph_json":"https://pith.science/api/pith-number/4SSTXQAMEJ6Z5YKB32SBUNOPRR/graph.json","events_json":"https://pith.science/api/pith-number/4SSTXQAMEJ6Z5YKB32SBUNOPRR/events.json","paper":"https://pith.science/paper/4SSTXQAM"},"agent_actions":{"view_html":"https://pith.science/pith/4SSTXQAMEJ6Z5YKB32SBUNOPRR","download_json":"https://pith.science/pith/4SSTXQAMEJ6Z5YKB32SBUNOPRR.json","view_paper":"https://pith.science/paper/4SSTXQAM","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2412.13678&json=true","fetch_graph":"https://pith.science/api/pith-number/4SSTXQAMEJ6Z5YKB32SBUNOPRR/graph.json","fetch_events":"https://pith.science/api/pith-number/4SSTXQAMEJ6Z5YKB32SBUNOPRR/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/4SSTXQAMEJ6Z5YKB32SBUNOPRR/action/timestamp_anchor","attest_storage":"https://pith.science/pith/4SSTXQAMEJ6Z5YKB32SBUNOPRR/action/storage_attestation","attest_author":"https://pith.science/pith/4SSTXQAMEJ6Z5YKB32SBUNOPRR/action/author_attestation","sign_citation":"https://pith.science/pith/4SSTXQAMEJ6Z5YKB32SBUNOPRR/action/citation_signature","submit_replication":"https://pith.science/pith/4SSTXQAMEJ6Z5YKB32SBUNOPRR/action/replication_record"}},"created_at":"2026-07-05T09:51:06.384338+00:00","updated_at":"2026-07-05T09:51:06.384338+00:00"}