{"work":{"id":"fddd3111-c69a-4453-8bf2-d3517e863145","openalex_id":null,"doi":null,"arxiv_id":"2211.09085","raw_key":null,"title":"Galactica: A Large Language Model for Science","authors":null,"authors_text":"Ross Taylor, Marcin Kardas, Guillem Cucurull, Thomas Scialom, Anthony Hartshorn, Elvis Saravia","year":2022,"venue":"cs.CL","abstract":"Information overload is a major obstacle to scientific progress. The explosive growth in scientific literature and data has made it ever harder to discover useful insights in a large mass of information. Today scientific knowledge is accessed through search engines, but they are unable to organize scientific knowledge alone. In this paper we introduce Galactica: a large language model that can store, combine and reason about scientific knowledge. We train on a large scientific corpus of papers, reference material, knowledge bases and many other sources. We outperform existing models on a range of scientific tasks. On technical knowledge probes such as LaTeX equations, Galactica outperforms the latest GPT-3 by 68.2% versus 49.0%. Galactica also performs well on reasoning, outperforming Chinchilla on mathematical MMLU by 41.3% to 35.7%, and PaLM 540B on MATH with a score of 20.4% versus 8.8%. It also sets a new state-of-the-art on downstream tasks such as PubMedQA and MedMCQA dev of 77.6% and 52.9%. And despite not being trained on a general corpus, Galactica outperforms BLOOM and OPT-175B on BIG-bench. We believe these results demonstrate the potential for language models as a new interface for science. We open source the model for the benefit of the scientific community.","external_url":"https://arxiv.org/abs/2211.09085","cited_by_count":null,"metadata_source":"pith","metadata_fetched_at":"2026-06-29T08:43:15.514436+00:00","pith_arxiv_id":"2211.09085","created_at":"2026-05-09T06:36:25.487191+00:00","updated_at":"2026-06-29T08:43:15.514436+00:00","title_quality_ok":true,"display_title":"Galactica: A Large Language Model for Science","render_title":"Galactica: A Large Language Model for Science"},"hub":{"state":{"work_id":"fddd3111-c69a-4453-8bf2-d3517e863145","tier":"hub","tier_reason":"10+ Pith inbound or 1,000+ external citations","pith_inbound_count":58,"external_cited_by_count":null,"distinct_field_count":10,"first_pith_cited_at":"2023-02-23T14:02:47+00:00","last_pith_cited_at":"2026-05-28T17:12:48+00:00","author_build_status":"not_needed","summary_status":"needed","contexts_status":"needed","graph_status":"needed","ask_index_status":"not_needed","reader_status":"not_needed","recognition_status":"not_needed","updated_at":"2026-06-30T03:39:23.342963+00:00","tier_text":"hub"},"tier":"hub","role_counts":[{"context_role":"background","n":11},{"context_role":"baseline","n":1},{"context_role":"method","n":1}],"polarity_counts":[{"context_polarity":"background","n":11},{"context_polarity":"baseline","n":1},{"context_polarity":"use_method","n":1}],"runs":{},"summary":{},"graph":{},"authors":[]}}