{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2021:AAVB643T3LCOW5JGYZOA4KQUAO","short_pith_number":"pith:AAVB643T","schema_version":"1.0","canonical_sha256":"002a1f7373dac4eb7526c65c0e2a14039318db897e79e2a1a5c09f54869c428d","source":{"kind":"arxiv","id":"2103.11943","version":1},"attestation_state":"computed","paper":{"title":"BERT: A Review of Applications in Natural Language Processing and Understanding","license":"http://creativecommons.org/licenses/by-sa/4.0/","headline":"","cross_cats":["cs.AI","cs.LG"],"primary_cat":"cs.CL","authors_text":"M. V. Koroteev","submitted_at":"2021-03-22T15:34:39Z","abstract_excerpt":"In this review, we describe the application of one of the most popular deep learning-based language models - BERT. The paper describes the mechanism of operation of this model, the main areas of its application to the tasks of text analytics, comparisons with similar models in each task, as well as a description of some proprietary models. In preparing this review, the data of several dozen original scientific articles published over the past few years, which attracted the most attention in the scientific community, were systematized. This survey will be useful to all students and researchers "},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2103.11943","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.CL","submitted_at":"2021-03-22T15:34:39Z","cross_cats_sorted":["cs.AI","cs.LG"],"title_canon_sha256":"62a59b3cbfb9ba0f9b877dfb635c7ddf119edc29c6ac6cc427ea92f89d50384d","abstract_canon_sha256":"43ee2f857a4b342d8283731a4dcaeb951a7a7458542caa3b5f19a965b9c3f6b1"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T02:25:13.737024Z","signature_b64":"v55fng8eLscbaswHUDj7yN5f+LlHSiAZeE/MIB1UrRDzZDMaz9+DfG5tJRgoSLpE5dlQrsvOTkwo5HBBdU2IBg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"002a1f7373dac4eb7526c65c0e2a14039318db897e79e2a1a5c09f54869c428d","last_reissued_at":"2026-07-05T02:25:13.736577Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T02:25:13.736577Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"BERT: A Review of Applications in Natural Language Processing and Understanding","license":"http://creativecommons.org/licenses/by-sa/4.0/","headline":"","cross_cats":["cs.AI","cs.LG"],"primary_cat":"cs.CL","authors_text":"M. V. Koroteev","submitted_at":"2021-03-22T15:34:39Z","abstract_excerpt":"In this review, we describe the application of one of the most popular deep learning-based language models - BERT. The paper describes the mechanism of operation of this model, the main areas of its application to the tasks of text analytics, comparisons with similar models in each task, as well as a description of some proprietary models. In preparing this review, the data of several dozen original scientific articles published over the past few years, which attracted the most attention in the scientific community, were systematized. This survey will be useful to all students and researchers "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2103.11943","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2103.11943/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2103.11943","created_at":"2026-07-05T02:25:13.736634+00:00"},{"alias_kind":"arxiv_version","alias_value":"2103.11943v1","created_at":"2026-07-05T02:25:13.736634+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2103.11943","created_at":"2026-07-05T02:25:13.736634+00:00"},{"alias_kind":"pith_short_12","alias_value":"AAVB643T3LCO","created_at":"2026-07-05T02:25:13.736634+00:00"},{"alias_kind":"pith_short_16","alias_value":"AAVB643T3LCOW5JG","created_at":"2026-07-05T02:25:13.736634+00:00"},{"alias_kind":"pith_short_8","alias_value":"AAVB643T","created_at":"2026-07-05T02:25:13.736634+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":12,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.24696","citing_title":"A Physics-Informed Fourier-Wavelet Transformer for Multiscale Computational Fluid Dynamics Surrogate Modeling","ref_index":44,"is_internal_anchor":false},{"citing_arxiv_id":"2605.24602","citing_title":"Correcting Visual Blur Induced by Attention Distraction to Reduce Hallucinations: Algorithm and Theory","ref_index":6,"is_internal_anchor":false},{"citing_arxiv_id":"2605.25655","citing_title":"Bandwidth-Aware LLM Inference on Heterogeneous Many-Core Supercomputers","ref_index":12,"is_internal_anchor":false},{"citing_arxiv_id":"2605.25521","citing_title":"CS-PQ: Cache-Friendly SIMD Product Quantization for Large-Scale ANNS Index Construction","ref_index":36,"is_internal_anchor":false},{"citing_arxiv_id":"2605.29639","citing_title":"RTP-LLM: High-Performance Alibaba LLM Inference Engine","ref_index":29,"is_internal_anchor":false},{"citing_arxiv_id":"2508.16419","citing_title":"Can LLMs Find Bugs in Code? An Evaluation from Beginner Errors to Security Vulnerabilities in Python and C++","ref_index":6,"is_internal_anchor":false},{"citing_arxiv_id":"2603.09108","citing_title":"Composed Vision-Language Retrieval for Skin Cancer Case Search via Joint Alignment of Global and Local Representations","ref_index":15,"is_internal_anchor":false},{"citing_arxiv_id":"2306.14565","citing_title":"Mitigating Hallucination in Large Multi-Modal Models via Robust Instruction Tuning","ref_index":10,"is_internal_anchor":false},{"citing_arxiv_id":"2604.23647","citing_title":"Hardware-Efficient Softmax and Layer Normalization with Guaranteed Normalization for Edge Devices","ref_index":2,"is_internal_anchor":false},{"citing_arxiv_id":"2605.00536","citing_title":"Tempus: A Temporally Scalable Resource-Invariant GEMM Streaming Framework for Versal AI Edge","ref_index":26,"is_internal_anchor":false},{"citing_arxiv_id":"2604.19567","citing_title":"Multi-modal Reasoning with LLMs for Visual Semantic Arithmetic","ref_index":8,"is_internal_anchor":false},{"citing_arxiv_id":"2604.18003","citing_title":"SELF-EMO: Emotional Self-Evolution from Recognition to Consistent Expression","ref_index":11,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/AAVB643T3LCOW5JGYZOA4KQUAO","json":"https://pith.science/pith/AAVB643T3LCOW5JGYZOA4KQUAO.json","graph_json":"https://pith.science/api/pith-number/AAVB643T3LCOW5JGYZOA4KQUAO/graph.json","events_json":"https://pith.science/api/pith-number/AAVB643T3LCOW5JGYZOA4KQUAO/events.json","paper":"https://pith.science/paper/AAVB643T"},"agent_actions":{"view_html":"https://pith.science/pith/AAVB643T3LCOW5JGYZOA4KQUAO","download_json":"https://pith.science/pith/AAVB643T3LCOW5JGYZOA4KQUAO.json","view_paper":"https://pith.science/paper/AAVB643T","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2103.11943&json=true","fetch_graph":"https://pith.science/api/pith-number/AAVB643T3LCOW5JGYZOA4KQUAO/graph.json","fetch_events":"https://pith.science/api/pith-number/AAVB643T3LCOW5JGYZOA4KQUAO/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/AAVB643T3LCOW5JGYZOA4KQUAO/action/timestamp_anchor","attest_storage":"https://pith.science/pith/AAVB643T3LCOW5JGYZOA4KQUAO/action/storage_attestation","attest_author":"https://pith.science/pith/AAVB643T3LCOW5JGYZOA4KQUAO/action/author_attestation","sign_citation":"https://pith.science/pith/AAVB643T3LCOW5JGYZOA4KQUAO/action/citation_signature","submit_replication":"https://pith.science/pith/AAVB643T3LCOW5JGYZOA4KQUAO/action/replication_record"}},"created_at":"2026-07-05T02:25:13.736634+00:00","updated_at":"2026-07-05T02:25:13.736634+00:00"}