{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2016:FYKBLULM3BQN3RTFETEVACCAML","short_pith_number":"pith:FYKBLULM","canonical_record":{"source":{"id":"1605.05404","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.DS","submitted_at":"2016-05-18T00:20:00Z","cross_cats_sorted":[],"title_canon_sha256":"148312762d4d5a61eaff2713946ffe64f56bf7107abd7dc2c30b9788a3c3a7f2","abstract_canon_sha256":"782ff59c4170363f8644722ef7ff41f9f65b384df573a5335ca2c4f628ec477d"},"schema_version":"1.0"},"canonical_sha256":"2e1415d16cd860ddc66524c950084062c83289ac8e8f5207e377dde94eec561a","source":{"kind":"arxiv","id":"1605.05404","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1605.05404","created_at":"2026-05-18T01:14:35Z"},{"alias_kind":"arxiv_version","alias_value":"1605.05404v1","created_at":"2026-05-18T01:14:35Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1605.05404","created_at":"2026-05-18T01:14:35Z"},{"alias_kind":"pith_short_12","alias_value":"FYKBLULM3BQN","created_at":"2026-05-18T12:30:15Z"},{"alias_kind":"pith_short_16","alias_value":"FYKBLULM3BQN3RTF","created_at":"2026-05-18T12:30:15Z"},{"alias_kind":"pith_short_8","alias_value":"FYKBLULM","created_at":"2026-05-18T12:30:15Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2016:FYKBLULM3BQN3RTFETEVACCAML","target":"record","payload":{"canonical_record":{"source":{"id":"1605.05404","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.DS","submitted_at":"2016-05-18T00:20:00Z","cross_cats_sorted":[],"title_canon_sha256":"148312762d4d5a61eaff2713946ffe64f56bf7107abd7dc2c30b9788a3c3a7f2","abstract_canon_sha256":"782ff59c4170363f8644722ef7ff41f9f65b384df573a5335ca2c4f628ec477d"},"schema_version":"1.0"},"canonical_sha256":"2e1415d16cd860ddc66524c950084062c83289ac8e8f5207e377dde94eec561a","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-18T01:14:35.090914Z","signature_b64":"bBWdlAhhSnFAJOlGaeXJtBCjpjWjf84/u+c41CyEP5eR3nXAmydp1Ca7luofeDxsGeWmN3fCjqrInQUEICAVCg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"2e1415d16cd860ddc66524c950084062c83289ac8e8f5207e377dde94eec561a","last_reissued_at":"2026-05-18T01:14:35.090186Z","signature_status":"signed_v1","first_computed_at":"2026-05-18T01:14:35.090186Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"1605.05404","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-18T01:14:35Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"0DMYw6ivgbZdv2mjAgg0ygNjWYYnI0lQvsmc/ofH/VQ8120J6WdVMYYVqJ5htLkdYOb41poCQYhPYbBYexb+AQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-05-31T23:52:42.296369Z"},"content_sha256":"6305eb3ee67d9681fd8f4d622773e3047cb40ceca257d029162654e4c8441b8c","schema_version":"1.0","event_id":"sha256:6305eb3ee67d9681fd8f4d622773e3047cb40ceca257d029162654e4c8441b8c"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2016:FYKBLULM3BQN3RTFETEVACCAML","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"CSA++: Fast Pattern Search for Large Alphabets","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.DS","authors_text":"Alistair Moffat, Matthias Petri, Simon Gog","submitted_at":"2016-05-18T00:20:00Z","abstract_excerpt":"Indexed pattern search in text has been studied for many decades. For small alphabets, the FM-Index provides unmatched performance, in terms of both space required and search speed. For large alphabets -- for example, when the tokens are words -- the situation is more complex, and FM-Index representations are compact, but potentially slow. In this paper we apply recent innovations from the field of inverted indexing and document retrieval to compressed pattern search, including for alphabets into the millions. Commencing with the practical compressed suffix array structure developed by Sadakan"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1605.05404","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-18T01:14:35Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"ChcURM1u1BkvEXGGYyiC/14ZNgGzcUfhHo2NyDH42iVY+So9SV0YD1h2TUyH9m7vTCRrrk1DeXHfbnQyZp0KCg==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-05-31T23:52:42.297045Z"},"content_sha256":"a3ccf087f562156be20eafb9bd11dd4e644a8c838ed2f6356ea17ff59d9e3bb3","schema_version":"1.0","event_id":"sha256:a3ccf087f562156be20eafb9bd11dd4e644a8c838ed2f6356ea17ff59d9e3bb3"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/FYKBLULM3BQN3RTFETEVACCAML/bundle.json","state_url":"https://pith.science/pith/FYKBLULM3BQN3RTFETEVACCAML/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/FYKBLULM3BQN3RTFETEVACCAML/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-05-31T23:52:42Z","links":{"resolver":"https://pith.science/pith/FYKBLULM3BQN3RTFETEVACCAML","bundle":"https://pith.science/pith/FYKBLULM3BQN3RTFETEVACCAML/bundle.json","state":"https://pith.science/pith/FYKBLULM3BQN3RTFETEVACCAML/state.json","well_known_bundle":"https://pith.science/.well-known/pith/FYKBLULM3BQN3RTFETEVACCAML/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2016:FYKBLULM3BQN3RTFETEVACCAML","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"782ff59c4170363f8644722ef7ff41f9f65b384df573a5335ca2c4f628ec477d","cross_cats_sorted":[],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.DS","submitted_at":"2016-05-18T00:20:00Z","title_canon_sha256":"148312762d4d5a61eaff2713946ffe64f56bf7107abd7dc2c30b9788a3c3a7f2"},"schema_version":"1.0","source":{"id":"1605.05404","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1605.05404","created_at":"2026-05-18T01:14:35Z"},{"alias_kind":"arxiv_version","alias_value":"1605.05404v1","created_at":"2026-05-18T01:14:35Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1605.05404","created_at":"2026-05-18T01:14:35Z"},{"alias_kind":"pith_short_12","alias_value":"FYKBLULM3BQN","created_at":"2026-05-18T12:30:15Z"},{"alias_kind":"pith_short_16","alias_value":"FYKBLULM3BQN3RTF","created_at":"2026-05-18T12:30:15Z"},{"alias_kind":"pith_short_8","alias_value":"FYKBLULM","created_at":"2026-05-18T12:30:15Z"}],"graph_snapshots":[{"event_id":"sha256:a3ccf087f562156be20eafb9bd11dd4e644a8c838ed2f6356ea17ff59d9e3bb3","target":"graph","created_at":"2026-05-18T01:14:35Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"paper":{"abstract_excerpt":"Indexed pattern search in text has been studied for many decades. For small alphabets, the FM-Index provides unmatched performance, in terms of both space required and search speed. For large alphabets -- for example, when the tokens are words -- the situation is more complex, and FM-Index representations are compact, but potentially slow. In this paper we apply recent innovations from the field of inverted indexing and document retrieval to compressed pattern search, including for alphabets into the millions. Commencing with the practical compressed suffix array structure developed by Sadakan","authors_text":"Alistair Moffat, Matthias Petri, Simon Gog","cross_cats":[],"headline":"","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.DS","submitted_at":"2016-05-18T00:20:00Z","title":"CSA++: Fast Pattern Search for Large Alphabets"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1605.05404","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:6305eb3ee67d9681fd8f4d622773e3047cb40ceca257d029162654e4c8441b8c","target":"record","created_at":"2026-05-18T01:14:35Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"782ff59c4170363f8644722ef7ff41f9f65b384df573a5335ca2c4f628ec477d","cross_cats_sorted":[],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.DS","submitted_at":"2016-05-18T00:20:00Z","title_canon_sha256":"148312762d4d5a61eaff2713946ffe64f56bf7107abd7dc2c30b9788a3c3a7f2"},"schema_version":"1.0","source":{"id":"1605.05404","kind":"arxiv","version":1}},"canonical_sha256":"2e1415d16cd860ddc66524c950084062c83289ac8e8f5207e377dde94eec561a","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"2e1415d16cd860ddc66524c950084062c83289ac8e8f5207e377dde94eec561a","first_computed_at":"2026-05-18T01:14:35.090186Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-18T01:14:35.090186Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"bBWdlAhhSnFAJOlGaeXJtBCjpjWjf84/u+c41CyEP5eR3nXAmydp1Ca7luofeDxsGeWmN3fCjqrInQUEICAVCg==","signature_status":"signed_v1","signed_at":"2026-05-18T01:14:35.090914Z","signed_message":"canonical_sha256_bytes"},"source_id":"1605.05404","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:6305eb3ee67d9681fd8f4d622773e3047cb40ceca257d029162654e4c8441b8c","sha256:a3ccf087f562156be20eafb9bd11dd4e644a8c838ed2f6356ea17ff59d9e3bb3"],"state_sha256":"7c49e266ae5b4b773cde76c75becaba6a583d3fc17265dece73793a53c2cebf1"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"O1l0AagVCcEwuYl5WE4sbRM3wvv7MHec9kJgZ/DFwsxToVlZhXGXmgnCs8sRlZ4312eDTfowq0hf/PStcvNyDA==","signed_message":"bundle_sha256_bytes","signed_at":"2026-05-31T23:52:42.300725Z","bundle_sha256":"327177d5e102a9000db1a0ec6caabf365b70ddb331cb5deb7bb612436c6c2de5"}}