{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2017:RLLVSZPSM46CYSLH6AB2O3KOGK","short_pith_number":"pith:RLLVSZPS","canonical_record":{"source":{"id":"1704.07329","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2017-04-24T17:07:26Z","cross_cats_sorted":[],"title_canon_sha256":"a91bf41ea7c4de5203132b6db41dddbd91a5869e692063245331f167c1e9f10d","abstract_canon_sha256":"f3ccdf9fa39eb0e94f96830c63ca383bb04b4c282b9cb6e499ba3de31644ddbb"},"schema_version":"1.0"},"canonical_sha256":"8ad75965f2673c2c4967f003a76d4e32922784b06675ad7c02ccfc1390c74838","source":{"kind":"arxiv","id":"1704.07329","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1704.07329","created_at":"2026-05-18T00:45:52Z"},{"alias_kind":"arxiv_version","alias_value":"1704.07329v1","created_at":"2026-05-18T00:45:52Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1704.07329","created_at":"2026-05-18T00:45:52Z"},{"alias_kind":"pith_short_12","alias_value":"RLLVSZPSM46C","created_at":"2026-05-18T12:31:39Z"},{"alias_kind":"pith_short_16","alias_value":"RLLVSZPSM46CYSLH","created_at":"2026-05-18T12:31:39Z"},{"alias_kind":"pith_short_8","alias_value":"RLLVSZPS","created_at":"2026-05-18T12:31:39Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2017:RLLVSZPSM46CYSLH6AB2O3KOGK","target":"record","payload":{"canonical_record":{"source":{"id":"1704.07329","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2017-04-24T17:07:26Z","cross_cats_sorted":[],"title_canon_sha256":"a91bf41ea7c4de5203132b6db41dddbd91a5869e692063245331f167c1e9f10d","abstract_canon_sha256":"f3ccdf9fa39eb0e94f96830c63ca383bb04b4c282b9cb6e499ba3de31644ddbb"},"schema_version":"1.0"},"canonical_sha256":"8ad75965f2673c2c4967f003a76d4e32922784b06675ad7c02ccfc1390c74838","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-18T00:45:52.902063Z","signature_b64":"sYo64IAqAh4KuG9pNmXqTkFkjOej3Sp1dw5TfHEuCZEjbw8QSauQ3A3czfzZowJ3Hz8WTVpoE0Ko9OP8dILoCQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"8ad75965f2673c2c4967f003a76d4e32922784b06675ad7c02ccfc1390c74838","last_reissued_at":"2026-05-18T00:45:52.900831Z","signature_status":"signed_v1","first_computed_at":"2026-05-18T00:45:52.900831Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"1704.07329","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-18T00:45:52Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"4hqTTECP6bS4iZCtdBVPoHJMDoqCHisx+J99IXJjaCS3Vc6XA3NLOomEPqqs8BBmLNpYgz2lrMKJnynJW0sLDg==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-06-25T21:46:02.246699Z"},"content_sha256":"f5437f41d4d3b961fda03f07b530e5c5e8fd248e4005a67b1e34a29f50a0fbea","schema_version":"1.0","event_id":"sha256:f5437f41d4d3b961fda03f07b530e5c5e8fd248e4005a67b1e34a29f50a0fbea"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2017:RLLVSZPSM46CYSLH6AB2O3KOGK","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"A Trie-Structured Bayesian Model for Unsupervised Morphological Segmentation","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Ahmet \\\"Ust\\\"un, Burcu Can, Murathan Kurfal{\\i}","submitted_at":"2017-04-24T17:07:26Z","abstract_excerpt":"In this paper, we introduce a trie-structured Bayesian model for unsupervised morphological segmentation. We adopt prior information from different sources in the model. We use neural word embeddings to discover words that are morphologically derived from each other and thereby that are semantically similar. We use letter successor variety counts obtained from tries that are built by neural word embeddings. Our results show that using different information sources such as neural word embeddings and letter successor variety as prior information improves morphological segmentation in a Bayesian "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1704.07329","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-18T00:45:52Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"UEyw/DATD1QtBKEiuGfecLZ8oBYCRWtmLJkmlmut+i1c6lNxPA+ueEUbYrhhrYy9cK824183OFNMnE5SoBQ+CQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-06-25T21:46:02.247048Z"},"content_sha256":"4fc5e48bb9710bb10c66141f5a9cac2f2179162e92c26a178162c8182a1ee9bc","schema_version":"1.0","event_id":"sha256:4fc5e48bb9710bb10c66141f5a9cac2f2179162e92c26a178162c8182a1ee9bc"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/RLLVSZPSM46CYSLH6AB2O3KOGK/bundle.json","state_url":"https://pith.science/pith/RLLVSZPSM46CYSLH6AB2O3KOGK/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/RLLVSZPSM46CYSLH6AB2O3KOGK/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-06-25T21:46:02Z","links":{"resolver":"https://pith.science/pith/RLLVSZPSM46CYSLH6AB2O3KOGK","bundle":"https://pith.science/pith/RLLVSZPSM46CYSLH6AB2O3KOGK/bundle.json","state":"https://pith.science/pith/RLLVSZPSM46CYSLH6AB2O3KOGK/state.json","well_known_bundle":"https://pith.science/.well-known/pith/RLLVSZPSM46CYSLH6AB2O3KOGK/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2017:RLLVSZPSM46CYSLH6AB2O3KOGK","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"f3ccdf9fa39eb0e94f96830c63ca383bb04b4c282b9cb6e499ba3de31644ddbb","cross_cats_sorted":[],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2017-04-24T17:07:26Z","title_canon_sha256":"a91bf41ea7c4de5203132b6db41dddbd91a5869e692063245331f167c1e9f10d"},"schema_version":"1.0","source":{"id":"1704.07329","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1704.07329","created_at":"2026-05-18T00:45:52Z"},{"alias_kind":"arxiv_version","alias_value":"1704.07329v1","created_at":"2026-05-18T00:45:52Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1704.07329","created_at":"2026-05-18T00:45:52Z"},{"alias_kind":"pith_short_12","alias_value":"RLLVSZPSM46C","created_at":"2026-05-18T12:31:39Z"},{"alias_kind":"pith_short_16","alias_value":"RLLVSZPSM46CYSLH","created_at":"2026-05-18T12:31:39Z"},{"alias_kind":"pith_short_8","alias_value":"RLLVSZPS","created_at":"2026-05-18T12:31:39Z"}],"graph_snapshots":[{"event_id":"sha256:4fc5e48bb9710bb10c66141f5a9cac2f2179162e92c26a178162c8182a1ee9bc","target":"graph","created_at":"2026-05-18T00:45:52Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"paper":{"abstract_excerpt":"In this paper, we introduce a trie-structured Bayesian model for unsupervised morphological segmentation. We adopt prior information from different sources in the model. We use neural word embeddings to discover words that are morphologically derived from each other and thereby that are semantically similar. We use letter successor variety counts obtained from tries that are built by neural word embeddings. Our results show that using different information sources such as neural word embeddings and letter successor variety as prior information improves morphological segmentation in a Bayesian ","authors_text":"Ahmet \\\"Ust\\\"un, Burcu Can, Murathan Kurfal{\\i}","cross_cats":[],"headline":"","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2017-04-24T17:07:26Z","title":"A Trie-Structured Bayesian Model for Unsupervised Morphological Segmentation"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1704.07329","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:f5437f41d4d3b961fda03f07b530e5c5e8fd248e4005a67b1e34a29f50a0fbea","target":"record","created_at":"2026-05-18T00:45:52Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"f3ccdf9fa39eb0e94f96830c63ca383bb04b4c282b9cb6e499ba3de31644ddbb","cross_cats_sorted":[],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2017-04-24T17:07:26Z","title_canon_sha256":"a91bf41ea7c4de5203132b6db41dddbd91a5869e692063245331f167c1e9f10d"},"schema_version":"1.0","source":{"id":"1704.07329","kind":"arxiv","version":1}},"canonical_sha256":"8ad75965f2673c2c4967f003a76d4e32922784b06675ad7c02ccfc1390c74838","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"8ad75965f2673c2c4967f003a76d4e32922784b06675ad7c02ccfc1390c74838","first_computed_at":"2026-05-18T00:45:52.900831Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-18T00:45:52.900831Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"sYo64IAqAh4KuG9pNmXqTkFkjOej3Sp1dw5TfHEuCZEjbw8QSauQ3A3czfzZowJ3Hz8WTVpoE0Ko9OP8dILoCQ==","signature_status":"signed_v1","signed_at":"2026-05-18T00:45:52.902063Z","signed_message":"canonical_sha256_bytes"},"source_id":"1704.07329","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:f5437f41d4d3b961fda03f07b530e5c5e8fd248e4005a67b1e34a29f50a0fbea","sha256:4fc5e48bb9710bb10c66141f5a9cac2f2179162e92c26a178162c8182a1ee9bc"],"state_sha256":"b8995617053ac75fce995e154947cd5c4031beb1eeb0a134f40973afcf56b5a4"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"AFOyUa8kxUELRaplvaRKtwoWKVSJhQbK1ISfqrWHA46ipdZTQpsqgGQmaqn0qIh4TZxk4U/YhNdnXWItsWuzAg==","signed_message":"bundle_sha256_bytes","signed_at":"2026-06-25T21:46:02.249094Z","bundle_sha256":"7a3262b554dc5f4e2ba291767a999f8f85be312d028c9becd6e52642dbc2b999"}}