{"as_of":"2026-08-22T01:33:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:30bb6f2f812beb1f0902c1a6830be02b0aa508deddd06ce8c9dc14644f1e4281","coverage":[{"denominator":31,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":31,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-14T13:18:47.837583Z","state":"measured"},{"denominator":31,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":31,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/1908.05672/citation-record","integrity":"/paper/1908.05672/integrity","json":"/paper/1908.05672/citation-record.json","paper":"/paper/1908.05672"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.642709Z","title":"URL: \" 'urlintro :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.642709Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:03d31532b991d0505b170b5cac092fcdfb125c691686c14db120bb2cfffaac88","observation_id":"569163d3-0598-4294-a365-63d6d32bb837","resolution":{"observed_at":"2026-08-14T13:18:47.642709Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.649704Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.649704Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:7d6089811e47c70d6b646fc83758aa04d1493418993ac6ef73cb639370dedfae","observation_id":"29f3cb42-a4ee-4306-b772-1028ddc1da4f","resolution":{"observed_at":"2026-08-14T13:18:47.649704Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.655571Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.655571Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:93d5af19b3bcc82009993fb8367c13d3bc0c640da6dd2b7ce7b2ca6597fcfd3b","observation_id":"8e17391a-1d46-40be-b3a3-8fe587514705","resolution":{"observed_at":"2026-08-14T13:18:47.655571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.406901Z","title":null,"venue":null,"work_id":"f9787c37-1646-4785-9fc9-d240107e30f7","year":2018},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.662290Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:0318bb0d18b49c06c057636759f6628641c3512fffbd970d4b8c2d7a80e67ff3","observation_id":"71590a0f-c48e-4093-891c-242ffd233316","resolution":{"observed_at":"2026-08-14T13:18:48.412406Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1412.3555","last_updated":"2014-12-11T06:46:53Z","snapshot_observed_at":"2026-08-20T00:41:09.240026Z","submitted_at":"2014-12-11T06:46:53Z","title":"Empirical Evaluation of Gated Recurrent Neural Networks on Sequence Modeling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.3555","snapshot_observed_at":"2026-08-14T13:18:47.669873Z","title":null,"venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.669873Z"},"links":{"cited_paper":"/paper/1412.3555","citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:a8456cf915893f448934c3ca6757beffd83da85882fce0fcb91ff15de5982bf5","observation_id":"2b58e37b-5252-4a14-a267-b2e7271d679b","resolution":{"observed_at":"2026-08-14T13:18:47.669873Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.378952Z","title":null,"venue":null,"work_id":"f70cb197-d559-4a82-8648-26c957f0e4f8","year":2011},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.676048Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:e09c95480e37601906c495f9da6963279b770e794bca54e0011dad794f31b499","observation_id":"3cf6eb12-a3b4-42b2-8af6-22883b1f5c08","resolution":{"observed_at":"2026-08-14T13:18:48.388483Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.352138Z","title":null,"venue":null,"work_id":"c569d321-d6dc-4b62-9506-d7518c6c19f2","year":2012},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.681466Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:90d87eb595726164d749fb459f5439fae0888ac0582f8e640e508b3f3850aa21","observation_id":"d7f4fd8c-8afc-41ae-8ec0-41c1a3934932","resolution":{"observed_at":"2026-08-14T13:18:48.361484Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1810.04805","last_updated":"2019-05-24T20:37:26Z","snapshot_observed_at":"2026-08-14T18:16:28.847993Z","submitted_at":"2018-10-11T00:50:01Z","title":"BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1810.04805","snapshot_observed_at":"2026-08-14T13:18:47.686960Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.686960Z"},"links":{"cited_paper":"/paper/1810.04805","citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:266bc3cb54838296eec518ce27d018233ae12804fabfdde0452db8e34720c16d","observation_id":"75af1952-1481-4ef9-a583-4ae2d7f22373","resolution":{"observed_at":"2026-08-14T13:18:47.686960Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1903.09722","last_updated":"2019-04-01T17:49:39Z","snapshot_observed_at":"2026-08-18T16:01:04.416009Z","submitted_at":"2019-03-22T22:14:51Z","title":"Pre-trained Language Model Representations for Language Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1903.09722","snapshot_observed_at":"2026-08-14T13:18:47.694598Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.694598Z"},"links":{"cited_paper":"/paper/1903.09722","citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:620e59885a8015421b7d4e6724825651e214d0d16e4902c266ef96d0f4a85352","observation_id":"7137753a-52c1-4a84-8888-3ff73621fcc4","resolution":{"observed_at":"2026-08-14T13:18:47.694598Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1312.6211","last_updated":"2015-03-04T01:43:31Z","snapshot_observed_at":"2026-08-14T23:50:42.566579Z","submitted_at":"2013-12-21T06:31:41Z","title":"An Empirical Investigation of Catastrophic Forgetting in Gradient-Based Neural Networks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1312.6211","snapshot_observed_at":"2026-08-14T13:18:47.700055Z","title":"Goodfellow , Mehdi Mirza , Da Xiao , Aaron Courville , and Yoshua Bengio","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.700055Z"},"links":{"cited_paper":"/paper/1312.6211","citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:274e9b1465bd421140a0d172744e73afd38fe0848ed627e369fe48fab983048e","observation_id":"bbf03def-f1b2-4c66-af26-d99044849e73","resolution":{"observed_at":"2026-08-14T13:18:47.700055Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.707228Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.707228Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:bd725268201b403d056ca1110e6f811517ec165108bfb790aa99f4815d51da81","observation_id":"fc5e9e12-3d61-4514-bdbf-28c8a88c7651","resolution":{"observed_at":"2026-08-14T13:18:47.707228Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.330000Z","title":null,"venue":null,"work_id":"bc615ee0-637d-428a-9bfb-432ead8812de","year":2018},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.714321Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:b9336b4c25d63cce4a568b7f0c3235d2b98057a3bb9b2fb43bbf20430b9481d9","observation_id":"711a31ca-f0a9-49d6-a67e-597750798f5b","resolution":{"observed_at":"2026-08-14T13:18:48.336339Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.311336Z","title":null,"venue":null,"work_id":"d8098a39-c7d1-4fbf-a2d4-5a066e623ccd","year":2015},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.721665Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:c25c80abdcdb41e170d2f476d1576330c362291627d916b46215dbda2b8a151e","observation_id":"26cdc785-ad98-4547-af47-1cc73ab1dcf5","resolution":{"observed_at":"2026-08-14T13:18:48.316760Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1901.07291","last_updated":"2019-01-22T13:22:34Z","snapshot_observed_at":"2026-08-14T17:27:41.965844Z","submitted_at":"2019-01-22T13:22:34Z","title":"Cross-lingual Language Model Pretraining","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1901.07291","snapshot_observed_at":"2026-08-14T13:18:47.728914Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.728914Z"},"links":{"cited_paper":"/paper/1901.07291","citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:a0e2be42e636af5088f304be3acccc3d6df9b943951ff2e676a1c971f41c1f10","observation_id":"41a97fcc-f2e4-45a6-a074-7a0ba8942567","resolution":{"observed_at":"2026-08-14T13:18:47.728914Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.737176Z","title":null,"venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.737176Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:131bf7e372798fb0a29de58a607cbc31ac7dbe53972f00cd96686bf83c1c0dbc","observation_id":"4c4936de-59d2-4f4b-8a47-ebf89999705f","resolution":{"observed_at":"2026-08-14T13:18:47.737176Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.280243Z","title":null,"venue":null,"work_id":"0d06581e-7291-416c-a50b-17eb6b9e22d4","year":2013},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.742886Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:05ee6015b6a4c9249882ac4cf120da1116a1b34b3d50ac0e35808421efdf5502","observation_id":"7da14f6a-d5f0-45be-ae44-930a25ce4a9d","resolution":{"observed_at":"2026-08-14T13:18:48.286250Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.749664Z","title":null,"venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.749664Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:81d4e9e0b3aaf296c9e1985269d35e623df11a6cf7633e63d5b41266eb0010d8","observation_id":"d0812dc7-f960-4be5-a99f-e8fc6178f33b","resolution":{"observed_at":"2026-08-14T13:18:47.749664Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.754905Z","title":null,"venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.754905Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:8d8943131810176a00a2acb1f32ac1d16e017fa825d9bdf4e7273aa5724150ab","observation_id":"b52c9d14-5ac0-45ee-94c1-0df8e14fa39c","resolution":{"observed_at":"2026-08-14T13:18:47.754905Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1701.06548","last_updated":"2017-01-23T18:35:28Z","snapshot_observed_at":"2026-08-14T21:19:55.038511Z","submitted_at":"2017-01-23T18:35:28Z","title":"Regularizing Neural Networks by Penalizing Confident Output Distributions","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1701.06548","snapshot_observed_at":"2026-08-14T13:18:47.760358Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.760358Z"},"links":{"cited_paper":"/paper/1701.06548","citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:269ff07ce3ebd17cec75e5726c14cbc9078ecbd7b54e2c41af0d397d877dd0d9","observation_id":"c5277ebf-9b3c-4b77-b819-3eaedc773766","resolution":{"observed_at":"2026-08-14T13:18:47.760358Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.765936Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.765936Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:4f4188bc0a3ff4dab69603c5270ddd622914bb75c13a1d6257925e1bcbf2ed1b","observation_id":"c1294628-9244-4f55-8d13-40058be09cab","resolution":{"observed_at":"2026-08-14T13:18:47.765936Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.772185Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.772185Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:2e6f2749439d369a9e3ad87851616dc3390fe7bc4361240f3d9f7e66d7df5279","observation_id":"69dfeb86-aab4-42e2-a2a5-7d955c43071b","resolution":{"observed_at":"2026-08-14T13:18:47.772185Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.778861Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.778861Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:b99eb4cba013d7bc14df244bf771b02c1a410fe97435fd5eed0efd8a9ac15750","observation_id":"c9c67669-e309-45ef-8960-0115e77deb0f","resolution":{"observed_at":"2026-08-14T13:18:47.778861Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.784821Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.784821Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:a81fe2274e92daf6c0a295b854c1efbb81f07e925471db08a45a792c4f1d8947","observation_id":"5d9ff0c2-2dce-415f-982c-157d8d898705","resolution":{"observed_at":"2026-08-14T13:18:47.784821Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/d17-1039","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.907191Z","title":null,"venue":null,"work_id":"e2fc94e5-a5d7-4d3d-b0dc-8f23bf4210e5","year":2017},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.790048Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:7817bb0966c074e6a2ec159848add77ad312bdfa85e331aa5aa6f39de624ca52","observation_id":"d4917625-3f9b-4075-8e10-a041768b2d9e","resolution":{"observed_at":"2026-08-14T13:18:47.914164Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.225690Z","title":null,"venue":null,"work_id":"c921afd2-b7a4-4232-8d80-f269712afc0d","year":2017},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.797318Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:470280ea18a22a0828148b9d15a06d876e713576d57d1463914c9e1f2de840dc","observation_id":"e691870c-078a-4c2f-b39c-1476f8abe682","resolution":{"observed_at":"2026-08-14T13:18:48.231334Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.210064Z","title":null,"venue":null,"work_id":"e7d5abf0-517a-4cd7-b84f-ea54d4c8d7dc","year":2013},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.803991Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:8ef7cfcc0a2af8fa36d600f7a6dc93c22ad1ecd14d4f50ec71f657ed50ccf4c7","observation_id":"b9ffb226-96e0-4c17-89d5-2abd3b210a18","resolution":{"observed_at":"2026-08-14T13:18:48.214758Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.191997Z","title":null,"venue":null,"work_id":"1a4617b6-0db8-4af8-9cd5-c40742f9547a","year":2018},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.809588Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:54a79080162a5cd647779adf2ff7b9ea3ab909fc0eae9251de364fe62d69275a","observation_id":"0ec9e62d-556a-48a1-8caa-805b9b344c48","resolution":{"observed_at":"2026-08-14T13:18:48.197576Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.817217Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.817217Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:b242aa4f5df063c58467548ecec5722efec7de33c8e278d7dff45da8ddfb85ca","observation_id":"88b398bf-ebc0-47ac-a13c-9bfad40c6c22","resolution":{"observed_at":"2026-08-14T13:18:47.817217Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.157469Z","title":"Gomez , Lukasz Kaiser , and Illia Polosukhin","venue":null,"work_id":"0fb5e6bd-c751-4cb1-83e8-581d3fb3cae4","year":2017},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.823873Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:9276cd11d59c94cf3e6481b33c6b3cfcbb5bf939165a3141c4562f94555c44a4","observation_id":"951cd7fc-0064-444d-b696-71a319e16154","resolution":{"observed_at":"2026-08-14T13:18:48.164362Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:48.136437Z","title":null,"venue":null,"work_id":"8f71face-84df-49ce-8e43-3f91ae83b2dc","year":2015},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.829283Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:d505585af4c6120b4a06f2762fe1d2348781de97d9f0fb17e1d84423a5b83245","observation_id":"9bdf5d9d-0ea1-4483-91a0-ac08f6776e70","resolution":{"observed_at":"2026-08-14T13:18:48.143528Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T13:18:47.837583Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation","version":5},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-14T13:18:47.837583Z"},"links":{"citing_paper":"/paper/1908.05672"},"observation_digest":"sha256:93d0ab4d3603386c8f96b20d7c528eef99efb2854e16b0467a9ab4ed9335b06d","observation_id":"7970a4c2-5a7f-46b1-9ddf-4269fe7190dc","resolution":{"observed_at":"2026-08-14T13:18:47.837583Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"1908.05672","last_updated":"2022-06-20T02:58:06Z","latest_version":5,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-19T16:33:32.935475Z","submitted_at":"2019-08-15T03:33:50Z","title":"Towards Making the Most of BERT in Neural Machine Translation"},"reference_resolution":{"displayed":31,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":29,"verified_exact":1,"verified_fuzzy":1},"total_outbound_references":31},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 22 August 2026, this Paper Citation Record lists 31 of 31 outbound references and 0 inbound Pith citation observations for arXiv:1908.05672."}