{"as_of":"2026-08-20T11:37:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7a03b0bdef913c62cd2f69206cb8b93d595d108f7ab318141ecb42a89e601784","coverage":[{"denominator":41,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":41,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-11T22:20:52.422988Z","state":"measured"},{"denominator":41,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":41,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-20T06:33:59.587034+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.04011/citation-record","integrity":"/paper/2607.04011/integrity","json":"/paper/2607.04011/citation-record.json","paper":"/paper/2607.04011"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2025.emnlp-main.1804","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Training compute-optimal transformer encoder models","venue":null,"work_id":"2a4b0a6f-e011-4d2a-9821-1c21db44f638","year":2025},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:3fdf3f4bd984224fce33da48d67897e9a74c54da34a2c4bfa3ffd705588d1dd1","observation_id":"ace548c8-16a9-42c0-9245-8188fc7159ff","resolution":{"observed_at":"2026-07-11T22:28:20.496998Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:22.769462+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:22.769462+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:c260bab50d32a85377caa0b51c09b1c53c540a184dfbe6d21288725cc95163e1","observation_id":"03748d60-f13c-4598-b48c-78191ff8a90e","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.19587","last_updated":"2025-06-06T20:57:10Z","snapshot_observed_at":"2026-08-16T12:54:51.139217Z","submitted_at":"2025-02-26T22:00:22Z","title":"NeoBERT: A Next-Generation BERT","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.19587","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2502.19587 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2502.19587","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:f094811a6c37fd86648a4ec9f142cbbf253a1fae6488643d4aef6515bcb168dd","observation_id":"e81b4826-dc72-4069-a966-c9a4452c0966","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.01613","last_updated":"2025-02-03T22:26:56Z","snapshot_observed_at":"2026-08-14T15:03:32.894363Z","submitted_at":"2024-02-02T18:23:18Z","title":"Nomic Embed: Training a Reproducible Long Context Text Embedder","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.01613","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2402.01613 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2402.01613","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:6f959baf18ff5189543acb43ec6c678eef007ee32f755592c5b90eef41a8ebe3","observation_id":"c267ec25-95f7-444e-9d2a-9b7912ed6ec3","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:2fe888673eedc7aa9d210fe0239d90ad4b08f68f3ae98fa798100f2cd07ab3cc","observation_id":"c93a426c-8097-4e02-bbc0-649ce50e6f05","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Proceedings of the 17th Conference of the European Chapter of the Association for Computational Linguistics , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:e3dc04710aa6bbc5a78ff0dd30b261629c4649efb5bd5694d6c22ae937c61d84","observation_id":"7396c3bf-c2a6-4ce3-92bf-9225faedadf9","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.14391","last_updated":"2025-04-10T07:50:15Z","snapshot_observed_at":"2026-08-20T11:12:19.487221Z","submitted_at":"2024-01-25T18:49:57Z","title":"Rethinking Patch Dependence for Masked Autoencoders","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.14391","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2401.14391 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2401.14391","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:f63a221960f87d7dcd7ab04653065277cada1513068fc00ac7808df7a5be0693","observation_id":"3cf1f46a-aac4-47b4-a04c-e97ef02f6780","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.08769","last_updated":"2025-06-30T10:48:00Z","snapshot_observed_at":"2026-08-17T20:53:24.468007Z","submitted_at":"2025-02-12T20:17:10Z","title":"Cluster and Predict Latent Patches for Improved Masked Image Modeling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.08769","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2502.08769 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2502.08769","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:6553fd7b0371702c764e6ffc89c8cfa6a51f63a4d0637118ac918a7c7b94280e","observation_id":"544ba040-2d98-410e-945d-49933a15a2e2","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"International conference on machine learning , pages=","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:3ab72ea7494c0a4994ed04e5ebb82d2aac2e2ef21d4565f3fa63d30b476c6780","observation_id":"79eeae20-b7f4-4d47-97e6-c518ae867bbf","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Proceedings of the IEEE/CVF international conference on computer vision , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:462c999ae5ffc302d1195638c0017685010f1b098ed43f6e8a4698a4898d81e7","observation_id":"6daafff7-c458-4ac8-a447-05af2d25aa18","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.07193","last_updated":"2024-02-02T10:24:09Z","snapshot_observed_at":"2026-08-17T13:03:40.359628Z","submitted_at":"2023-04-14T15:12:19Z","title":"DINOv2: Learning Robust Visual Features without Supervision","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.07193","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2304.07193 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2304.07193","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:1f2422675946688c6c1f5384c161aacef047d2bb3ab7144ef9b16e71ce6447e8","observation_id":"4c28df6a-9739-4792-af68-185d9b0e5f5e","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.10104","last_updated":"2025-08-13T18:00:55Z","snapshot_observed_at":"2026-08-18T01:07:23.737664Z","submitted_at":"2025-08-13T18:00:55Z","title":"DINOv3","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.10104","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2508.10104 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2508.10104","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:d96c57dfce8fc16421b74dd774158b92e549e438d5996cdfe6538cb9828799ac","observation_id":"6379d943-2cca-4af0-98b6-1f3f65add56e","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2112.09118","last_updated":"2022-08-29T12:17:32Z","snapshot_observed_at":"2026-08-14T13:19:21.742321Z","submitted_at":"2021-12-16T18:57:37Z","title":"Unsupervised Dense Information Retrieval with Contrastive Learning","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.09118","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2112.09118 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2112.09118","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:143b27048718b7fed992ffb27b8f89e714e9786e90227d7d46dd13c6d4cd5ccd","observation_id":"f3499700-1d26-4260-896f-19330642bd36","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12035","last_updated":"2022-10-17T14:08:37Z","snapshot_observed_at":"2026-08-18T02:27:47.037722Z","submitted_at":"2022-05-24T12:43:04Z","title":"RetroMAE: Pre-Training Retrieval-oriented Language Models Via Masked Auto-Encoder","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12035","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2205.12035 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2205.12035","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:0e74e6074aec13f29369bf18461f9f97175ae3e256424f230834e152aaecf71d","observation_id":"53857998-f47c-40b9-bfa8-8b1c804b38c0","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Proceedings of the 2018 EMNLP workshop BlackboxNLP: Analyzing and interpreting neural networks for NLP , pages=","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:748cd6041315e90a3550ab63e0a8baca92cb1d7ad67977217ec715d57c7929a7","observation_id":"69b00ae8-84c5-4040-831f-bce93611320b","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1611.09268","last_updated":"2018-10-31T14:46:47Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2016-11-28T18:14:11Z","title":"MS MARCO: A Human Generated MAchine Reading COmprehension Dataset","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1611.09268","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:1611.09268 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/1611.09268","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:d652eb386f55d6e0c4b163910ac963caeea46e5aae09aa4d5aad13167936adf2","observation_id":"147c304b-e940-4abe-ba3c-dd86b45ca0da","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:a059431cd8e0b5ef6e05e1f0c1b93d6ab36c3a2268fe67858a44090fe38ca73d","observation_id":"6081b905-7ad9-451e-bd7c-9665cf5fb463","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2212.03533","last_updated":"2024-02-22T06:21:51Z","snapshot_observed_at":"2026-07-06T14:27:46.217000Z","submitted_at":"2022-12-07T09:25:54Z","title":"Text Embeddings by Weakly-Supervised Contrastive Pre-training","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.03533","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2212.03533 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2212.03533","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:094855d23e80047cc5169e2b81dea83ee2e376048ad2a5ef45128c34e85dd892","observation_id":"08f24245-a9a4-4b5a-a830-758de003b8ff","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:259e2b0c0fd6f8488634c0b01accf1e3247ed8db6dbbd3bb6ab91ff2c096df3b","observation_id":"f5754cf5-b3f6-45bd-b967-f5a772023d9d","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2105.04906","last_updated":"2022-01-28T12:23:37Z","snapshot_observed_at":"2026-08-12T14:02:29.797842Z","submitted_at":"2021-05-11T09:53:21Z","title":"VICReg: Variance-Invariance-Covariance Regularization for Self-Supervised Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2105.04906","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2105.04906 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2105.04906","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:e2509620ad34b7e77f246a8c01cf86078d6f68d55a7577ea90fabeca1fa3fef6","observation_id":"1e6e4f8f-f1c6-4382-84b7-51b23a71b3e7","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.00504","last_updated":"2024-03-01T13:05:38Z","snapshot_observed_at":"2026-08-16T14:13:46.486630Z","submitted_at":"2024-03-01T13:05:38Z","title":"Learning and Leveraging World Models in Visual Representation Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.00504","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2403.00504 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2403.00504","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:b012732aa01aed43660a00e217d73590795b04c1c390147af0a57a16cb87ac18","observation_id":"2574f0cd-bd8d-4b30-aaf4-6750e568f37e","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.08544","last_updated":"2025-11-14T08:38:32Z","snapshot_observed_at":"2026-08-18T08:21:10.395771Z","submitted_at":"2025-11-11T18:21:55Z","title":"LeJEPA: Provable and Scalable Self-Supervised Learning Without the Heuristics","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.08544","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2511.08544 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2511.08544","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:8a436dda2c6d93260477c70e7da834d83c349bebaa86de2b67f89bcc03b03d4b","observation_id":"16367a4d-aa9b-49df-a479-003bd4bd73bb","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.13181","last_updated":"2025-04-28T18:01:39Z","snapshot_observed_at":"2026-08-11T19:32:15.215968Z","submitted_at":"2025-04-17T17:59:57Z","title":"Perception Encoder: The best visual embeddings are not at the output of the network","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.13181","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2504.13181 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2504.13181","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:6871431b4b5e97de4cf64733f32df533e51d46d7466a688de23514b082efbca8","observation_id":"44839e14-19b5-4577-a689-1ea2025bf5ff","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.15389","last_updated":"2023-03-27T17:02:21Z","snapshot_observed_at":"2026-08-20T09:51:02.717471Z","submitted_at":"2023-03-27T17:02:21Z","title":"EVA-CLIP: Improved Training Techniques for CLIP at Scale","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.15389","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2303.15389 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2303.15389","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:09029870f8f8f57e81e03a5fd605c4367d2c6c7530a5ec9ad78bdede9c9c46b8","observation_id":"e5d8635c-6d00-4c5c-9886-d8d9a71eeca5","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Proceedings of the IEEE/CVF conference on computer vision and pattern recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:8493f8b78526180b6258328608a9685602e5d8304dc76bb4bcd5559e84eacf76","observation_id":"3e42ed0a-fbbf-4cfb-abf2-0b802c7f26de","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1907.11692","last_updated":"2019-07-26T17:48:29Z","snapshot_observed_at":"2026-08-16T14:33:50.657682Z","submitted_at":"2019-07-26T17:48:29Z","title":"RoBERTa: A Robustly Optimized BERT Pretraining Approach","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1907.11692","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:1907.11692 , year=","venue":null,"work_id":null,"year":1907},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/1907.11692","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:06fe2c32fae9fe678406d495ced74081f6cc4f7d085b99fc6fb11efed885af64","observation_id":"08cc1f4a-46c8-4d2a-8827-2988d0a8adeb","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:73c1ce0ac5033c058d266da07ce66443b1114434f7f511fe2d6ee6ce3676276a","observation_id":"5c9fb3e8-4364-4fda-9516-886ce5659f61","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:f77a36c678de9c07334435305a0b694f16da765f9786e8d6ce962e481722d4ff","observation_id":"e2c902cc-df0e-41b8-8594-f4c2d509056f","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2104.08663","last_updated":"2021-10-21T01:18:28Z","snapshot_observed_at":"2026-08-09T23:22:21.200279Z","submitted_at":"2021-04-17T23:29:55Z","title":"BEIR: A Heterogenous Benchmark for Zero-shot Evaluation of Information Retrieval Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.08663","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2104.08663 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2104.08663","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:3370d7d0f9f2e4f998c50e12744af259a0bd59b44bc78b7b36859e72fc8e995c","observation_id":"5216b7e0-d826-4e48-8605-124bd966a756","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2001.08361","last_updated":"2020-01-23T03:59:20Z","snapshot_observed_at":"2026-08-13T17:41:53.092611Z","submitted_at":"2020-01-23T03:59:20Z","title":"Scaling Laws for Neural Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2001.08361","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv:2001.08361 , year=","venue":null,"work_id":null,"year":2001},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2001.08361","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:19021c3b588251cf6bf77d3ecbc00975aa361d68da0b250bc7254b4359e670c5","observation_id":"1a909210-018b-4196-98e4-82de565646d5","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"and Sifre, Laurent , title =","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:97f5c4c74f11159a6d1c4f90ea3b3bf04c53466c3e16ce83f88a9b2360599c46","observation_id":"f4da7a38-a7ed-47fe-a9f3-3c585fcb5f48","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.02954","last_updated":"2024-01-05T18:59:13Z","snapshot_observed_at":"2026-08-14T19:47:04.330126Z","submitted_at":"2024-01-05T18:59:13Z","title":"DeepSeek LLM: Scaling Open-Source Language Models with Longtermism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.02954","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv:2401.02954 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2401.02954","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:4e91a210e3300da4cd6b0ed5b64ddbd14136c4c85a706c9a5b4c4b39652954b0","observation_id":"e2d2898d-a5c9-4571-a03a-3a4bc30675f3","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08540","last_updated":"2024-06-14T20:21:05Z","snapshot_observed_at":"2026-08-16T14:10:11.362149Z","submitted_at":"2024-03-13T13:54:00Z","title":"Language models scale reliably with over-training and on downstream tasks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.08540","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv:2403.08540 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2403.08540","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:667fe6aef644e5c10a30d70bbc3acade3eeff4efbcad4497c5e4042f1a44db2c","observation_id":"5afe7ad2-b0fc-424d-afc7-d3a3a7d81422","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:1e01af25f5223353cdf8da46b8a909d00b3ebdda313551395f359593c1b4beb3","observation_id":"5020cd9c-3c71-4129-8527-b8b8db863df2","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Journal of machine learning research , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:f52d8dec557515bc045ebf0b661dcb659afee0b706aea56ad6e52c000b3516b1","observation_id":"88cd7eef-cdf1-4551-b21a-74539cf44f51","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2104.08821","last_updated":"2022-05-18T12:29:49Z","snapshot_observed_at":"2026-07-06T11:01:05.577957Z","submitted_at":"2021-04-18T11:27:08Z","title":"SimCSE: Simple Contrastive Learning of Sentence Embeddings","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.08821","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2104.08821 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2104.08821","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:0ad3a3438b0058c2212002c7fb5d5f454d55b28385b8d433ff848834a5e68761","observation_id":"5230b686-ce9a-4045-99ae-c79d157e59ce","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.03281","last_updated":"2023-08-07T03:52:59Z","snapshot_observed_at":"2026-08-14T21:56:20.223557Z","submitted_at":"2023-08-07T03:52:59Z","title":"Towards General Text Embeddings with Multi-stage Contrastive Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.03281","snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"arXiv preprint arXiv:2308.03281 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"cited_paper":"/paper/2308.03281","citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:66c496db70c16b9161369bfd1dc059665e6fe01f54f101acecdc51d47e0a5ac1","observation_id":"d74ece71-b0d8-4cef-9942-b4e02ae7c108","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"Proceedings of the 17th Conference of the European Chapter of the Association for Computational Linguistics , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:7021b4b09e9f098d7991bde78128f418715f5a4cea4d3ec8496191e4999a4e0a","observation_id":"93981c5f-898f-4d7b-be0f-97b0420b5cc2","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"2020 , eprint=","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:d8121fb75ec2caa5de7562d0deaf935b0d69b599d77556f73d31bdc2cd6cc737","observation_id":"4ab70ef9-b7ae-4f03-8ce0-fa9ce811bac7","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T22:20:52.422988Z","title":"2023 , eprint=","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:f6354ffadb44f39e1a5f5a9ff317210e90b193c3987381544846b8e37429b8e3","observation_id":"bcc744ac-c545-4b43-b3b9-4e68fe6ca1a8","resolution":{"observed_at":"2026-07-11T22:20:52.422988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.findings-emnlp.535","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Are ELECTRA ' s Sentence Embeddings Beyond Repair? The Case of Semantic Textual Similarity","venue":null,"work_id":"8500d5d5-a396-46b5-bbb3-1abff62c674a","year":2024},"citing_paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-07-11T22:20:52.422988Z"},"links":{"citing_paper":"/paper/2607.04011"},"observation_digest":"sha256:67cfd3c4579124397db1d7fb919c734e95ad2842b7dc5ddd04fc5b840dddc6db","observation_id":"046de455-fdae-4f44-b8df-ac7a9ef7fda9","resolution":{"observed_at":"2026-07-11T22:28:20.487537Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:23.012685+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:23.012685+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2607.04011","last_updated":"2026-07-04T20:13:38Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-19T07:44:13.811277Z","submitted_at":"2026-07-04T20:13:38Z","title":"Separating Representation from Reconstruction Enables Scalable Text Encoders"},"reference_resolution":{"displayed":41,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":39,"verified_exact":2,"verified_fuzzy":0},"total_outbound_references":41},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 41 of 41 outbound references and 0 inbound Pith citation observations for arXiv:2607.04011."}