{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:LCMZUO4PARAZQYTX6FC3MC4TLZ","short_pith_number":"pith:LCMZUO4P","schema_version":"1.0","canonical_sha256":"58999a3b8f0441986277f145b60b935e60b5d1e48665a1bed5eec9a6aaa286b9","source":{"kind":"arxiv","id":"2309.14517","version":2},"attestation_state":"computed","paper":{"title":"Watch Your Language: Investigating Content Moderation with Large Language Models","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI","cs.CL","cs.CR","cs.SI"],"primary_cat":"cs.HC","authors_text":"Deepak Kumar, Yousef AbuHashem, Zakir Durumeric","submitted_at":"2023-09-25T20:23:51Z","abstract_excerpt":"Large language models (LLMs) have exploded in popularity due to their ability to perform a wide array of natural language tasks. Text-based content moderation is one LLM use case that has received recent enthusiasm, however, there is little research investigating how LLMs perform in content moderation settings. In this work, we evaluate a suite of commodity LLMs on two common content moderation tasks: rule-based community moderation and toxic content detection. For rule-based community moderation, we instantiate 95 subcommunity specific LLMs by prompting GPT-3.5 with rules from 95 Reddit subco"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2309.14517","kind":"arxiv","version":2},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.HC","submitted_at":"2023-09-25T20:23:51Z","cross_cats_sorted":["cs.AI","cs.CL","cs.CR","cs.SI"],"title_canon_sha256":"64901b3353e2682e1445ace96b5752f873a343e9c1245f3d3f1d5236a1eb0840","abstract_canon_sha256":"2aeace680063e68ff1354777e2214d91cee99f4b0e749a86d7c77bf176fe5603"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T07:34:35.551500Z","signature_b64":"XzOR+2ypA8p1j+QGP2pYf9CEAXBQz7h4bf0GiWpdfyHceYtgPh/aoVAQb8Be21gWfarP/XJkaBQ3R4CgWGzbCg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"58999a3b8f0441986277f145b60b935e60b5d1e48665a1bed5eec9a6aaa286b9","last_reissued_at":"2026-07-05T07:34:35.550993Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T07:34:35.550993Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Watch Your Language: Investigating Content Moderation with Large Language Models","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI","cs.CL","cs.CR","cs.SI"],"primary_cat":"cs.HC","authors_text":"Deepak Kumar, Yousef AbuHashem, Zakir Durumeric","submitted_at":"2023-09-25T20:23:51Z","abstract_excerpt":"Large language models (LLMs) have exploded in popularity due to their ability to perform a wide array of natural language tasks. Text-based content moderation is one LLM use case that has received recent enthusiasm, however, there is little research investigating how LLMs perform in content moderation settings. In this work, we evaluate a suite of commodity LLMs on two common content moderation tasks: rule-based community moderation and toxic content detection. For rule-based community moderation, we instantiate 95 subcommunity specific LLMs by prompting GPT-3.5 with rules from 95 Reddit subco"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2309.14517","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2309.14517/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2309.14517","created_at":"2026-07-05T07:34:35.551047+00:00"},{"alias_kind":"arxiv_version","alias_value":"2309.14517v2","created_at":"2026-07-05T07:34:35.551047+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2309.14517","created_at":"2026-07-05T07:34:35.551047+00:00"},{"alias_kind":"pith_short_12","alias_value":"LCMZUO4PARAZ","created_at":"2026-07-05T07:34:35.551047+00:00"},{"alias_kind":"pith_short_16","alias_value":"LCMZUO4PARAZQYTX","created_at":"2026-07-05T07:34:35.551047+00:00"},{"alias_kind":"pith_short_8","alias_value":"LCMZUO4P","created_at":"2026-07-05T07:34:35.551047+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":3,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.00603","citing_title":"Toward Agentic Governance: What Shapes LLM-Agent Intervention in Public Forums?","ref_index":46,"is_internal_anchor":false},{"citing_arxiv_id":"2605.23676","citing_title":"AI at the Front Lines of Platform Governance: Using LLMs to Support Illegal Content Reporting under the Digital Services Act","ref_index":75,"is_internal_anchor":false},{"citing_arxiv_id":"2409.18169","citing_title":"Harmful Fine-tuning Attacks and Defenses for Large Language Models: A Survey","ref_index":82,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/LCMZUO4PARAZQYTX6FC3MC4TLZ","json":"https://pith.science/pith/LCMZUO4PARAZQYTX6FC3MC4TLZ.json","graph_json":"https://pith.science/api/pith-number/LCMZUO4PARAZQYTX6FC3MC4TLZ/graph.json","events_json":"https://pith.science/api/pith-number/LCMZUO4PARAZQYTX6FC3MC4TLZ/events.json","paper":"https://pith.science/paper/LCMZUO4P"},"agent_actions":{"view_html":"https://pith.science/pith/LCMZUO4PARAZQYTX6FC3MC4TLZ","download_json":"https://pith.science/pith/LCMZUO4PARAZQYTX6FC3MC4TLZ.json","view_paper":"https://pith.science/paper/LCMZUO4P","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2309.14517&json=true","fetch_graph":"https://pith.science/api/pith-number/LCMZUO4PARAZQYTX6FC3MC4TLZ/graph.json","fetch_events":"https://pith.science/api/pith-number/LCMZUO4PARAZQYTX6FC3MC4TLZ/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/LCMZUO4PARAZQYTX6FC3MC4TLZ/action/timestamp_anchor","attest_storage":"https://pith.science/pith/LCMZUO4PARAZQYTX6FC3MC4TLZ/action/storage_attestation","attest_author":"https://pith.science/pith/LCMZUO4PARAZQYTX6FC3MC4TLZ/action/author_attestation","sign_citation":"https://pith.science/pith/LCMZUO4PARAZQYTX6FC3MC4TLZ/action/citation_signature","submit_replication":"https://pith.science/pith/LCMZUO4PARAZQYTX6FC3MC4TLZ/action/replication_record"}},"created_at":"2026-07-05T07:34:35.551047+00:00","updated_at":"2026-07-05T07:34:35.551047+00:00"}