{"as_of":"2026-08-19T23:57:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2405d07a06fba3f52de70ed77447e4a6ef61790e86677874bd2669d13c640a8b","coverage":[{"denominator":43,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":43,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T19:02:04.157451Z","state":"measured"},{"denominator":43,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":43,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.10709/citation-record","integrity":"/paper/2608.10709/integrity","json":"/paper/2608.10709/citation-record.json","paper":"/paper/2608.10709"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:03.779686Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.779686Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:ef2a129bca359ca3fb4af9368e281cec98742fafe4bcea10821522732e2d623d","observation_id":"05f9c82f-760c-453c-9f2f-cb7bff92e45e","resolution":{"observed_at":"2026-08-12T19:02:03.779686Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:03.789566Z","title":"International conference on machine learning , pages=","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.789566Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:c2893f7079584c67117d13990319f9b61b5c73568c55c55f31e765cb64bd2c34","observation_id":"5cc5eac1-0880-42cb-91e1-d47f04d64288","resolution":{"observed_at":"2026-08-12T19:02:03.789566Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.08295","last_updated":"2021-06-15T17:12:42Z","snapshot_observed_at":"2026-08-19T05:42:08.537901Z","submitted_at":"2021-06-15T17:12:42Z","title":"A White Paper on Neural Network Quantization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.08295","snapshot_observed_at":"2026-08-12T19:02:03.795582Z","title":"arXiv preprint arXiv:2106.08295 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.795582Z"},"links":{"cited_paper":"/paper/2106.08295","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:e23719598ec8c40d3ef62a3527be2db22689480ef56613e82182bd284bee38d3","observation_id":"02da2389-2404-4898-a6ad-c5cf9e263dce","resolution":{"observed_at":"2026-08-12T19:02:03.795582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:03.802665Z","title":"Proceedings of the IEEE conference on computer vision and pattern recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.802665Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:5aebbe776f754bc28b831fdae7a032a457e7fc4ebbd70acfd4960eea05259807","observation_id":"9f13cc49-eb46-4e7f-b71e-068ccba370a5","resolution":{"observed_at":"2026-08-12T19:02:03.802665Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1503.02531","last_updated":"2015-03-09T15:44:49Z","snapshot_observed_at":"2026-08-16T18:00:58.008096Z","submitted_at":"2015-03-09T15:44:49Z","title":"Distilling the Knowledge in a Neural Network","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1503.02531","snapshot_observed_at":"2026-08-12T19:02:03.809356Z","title":"arXiv preprint arXiv:1503.02531 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.809356Z"},"links":{"cited_paper":"/paper/1503.02531","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:c036f78425e167cd7e1866d5646e38e411d40eaad199a946933a099767a8079a","observation_id":"3a619c7b-6027-41ce-ac7c-daa578dd60c6","resolution":{"observed_at":"2026-08-12T19:02:03.809356Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:05.397826Z","title":"International Conference on Artificial Intelligence and Statistics , pages=","venue":null,"work_id":"34ce1d86-71a5-4b6b-a9d3-4a98cb36f391","year":2024},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.816336Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:d3a068729b5c30cb6766133906792983fce03114a22fd83d9ad5fabbb1374d11","observation_id":"2ef3cbcc-3431-454b-8ac8-c1410a0996bd","resolution":{"observed_at":"2026-08-12T19:02:05.408373Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:05.362492Z","title":"Proceedings of the IEEE/CVF conference on computer vision and pattern recognition , pages=","venue":null,"work_id":"e094f8d4-9024-461f-b8b4-3a0382530c30","year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.827800Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:d8176c688d8024b2fae410c00c28e2513499ccc10231dfe4358e84cb87b3d4be","observation_id":"7c0306eb-e480-4f5e-b071-6321390ac07e","resolution":{"observed_at":"2026-08-12T19:02:05.374735Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1308.3432","last_updated":"2013-08-15T15:19:34Z","snapshot_observed_at":"2026-08-14T04:51:04.817737Z","submitted_at":"2013-08-15T15:19:34Z","title":"Estimating or Propagating Gradients Through Stochastic Neurons for Conditional Computation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1308.3432","snapshot_observed_at":"2026-08-12T19:02:03.838117Z","title":"arXiv preprint arXiv:1308.3432 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.838117Z"},"links":{"cited_paper":"/paper/1308.3432","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:b440e0d95117162090491d89df308496dcfee2fa22f1a9e33d19d6a829399663","observation_id":"487fa85f-e14a-4f93-845b-8d6849485000","resolution":{"observed_at":"2026-08-12T19:02:03.838117Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1810.00861","last_updated":"2019-03-05T00:28:48Z","snapshot_observed_at":"2026-08-15T19:52:05.516451Z","submitted_at":"2018-10-01T17:57:02Z","title":"ProxQuant: Quantized Neural Networks via Proximal Operators","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1810.00861","snapshot_observed_at":"2026-08-12T19:02:03.848114Z","title":"arXiv preprint arXiv:1810.00861 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.848114Z"},"links":{"cited_paper":"/paper/1810.00861","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:596b3d2629ccd8747a3ed3cede13cacaf849745bffda6aca1174a2db0c310445","observation_id":"eaecadf8-21c8-4bb9-abb6-2d5e7954c4f5","resolution":{"observed_at":"2026-08-12T19:02:03.848114Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:05.333151Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"42aada4e-df67-41d1-bb9c-fce4451b8d40","year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.860291Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:745f0c15c26b214c223aea7f94d0969a20daf1fd6f77642530fb0403e3456a06","observation_id":"1199dce6-df90-41f6-a5eb-ce023b12aa0b","resolution":{"observed_at":"2026-08-12T19:02:05.344120Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:05.307776Z","title":"Proceedings of the IEEE/CVF international conference on computer vision , pages=","venue":null,"work_id":"af131b7e-8e0e-4805-a793-85eadddd2097","year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.868956Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:fac5fd2dba178d716203cc1ab596ba06ddd48eba44ef0f6958077b1d78f821d6","observation_id":"1a61a792-ebe5-48c0-9a81-32681caae9e2","resolution":{"observed_at":"2026-08-12T19:02:05.315780Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1911.12491","last_updated":"2019-11-28T02:27:27Z","snapshot_observed_at":"2026-08-18T23:09:40.418319Z","submitted_at":"2019-11-28T02:27:27Z","title":"QKD: Quantization-aware Knowledge Distillation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1911.12491","snapshot_observed_at":"2026-08-12T19:02:03.877994Z","title":"arXiv preprint arXiv:1911.12491 , year=","venue":null,"work_id":null,"year":1911},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.877994Z"},"links":{"cited_paper":"/paper/1911.12491","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:a14360d0bffe9332c284002c2ce1e4643c3e03c3f1bdf5de501679f837f4d85c","observation_id":"493b5986-99e2-4f2b-aab4-c1ce33d2d640","resolution":{"observed_at":"2026-08-12T19:02:03.877994Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6550","last_updated":"2015-03-27T11:52:28Z","snapshot_observed_at":"2026-08-17T23:25:45.947983Z","submitted_at":"2014-12-19T22:40:51Z","title":"FitNets: Hints for Thin Deep Nets","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.6550","snapshot_observed_at":"2026-08-12T19:02:03.886995Z","title":"arXiv 2014 , author=","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.886995Z"},"links":{"cited_paper":"/paper/1412.6550","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:7b7c27f1d22471d3ea95d00a27b4a7c327d00fa44af407f58037a004157be7ff","observation_id":"6d2fe63e-86d1-4679-b701-648910dd4904","resolution":{"observed_at":"2026-08-12T19:02:03.886995Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1612.03928","last_updated":"2017-02-12T22:05:47Z","snapshot_observed_at":"2026-08-17T23:26:36.824318Z","submitted_at":"2016-12-12T21:15:57Z","title":"Paying More Attention to Attention: Improving the Performance of Convolutional Neural Networks via Attention Transfer","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1612.03928","snapshot_observed_at":"2026-08-12T19:02:03.895017Z","title":"arXiv preprint arXiv:1612.03928 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.895017Z"},"links":{"cited_paper":"/paper/1612.03928","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:bc9207b656d919df141d6a3eca42ab37a4cf9b30cc8ddd122f794b8762ba19ac","observation_id":"20f22dc9-6160-44c5-bb24-951b023124fa","resolution":{"observed_at":"2026-08-12T19:02:03.895017Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:03.901150Z","title":"Proceedings of the IEEE/CVF international conference on computer vision , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.901150Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:42891718ca387ecc59ede39dac68b539a7c45a8e2af8787f4b07309b58a67cc2","observation_id":"f0dd0058-68eb-4286-9c2e-32f7977e4bcc","resolution":{"observed_at":"2026-08-12T19:02:03.901150Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:03.912721Z","title":"Proceedings of the IEEE conference on computer vision and pattern recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.912721Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:636b0280d0a23cda9489bd18f74d0b766f669b3e7fa91a7763d291b212e648ad","observation_id":"36352b8c-c2e7-4016-9cea-fc29cebec145","resolution":{"observed_at":"2026-08-12T19:02:03.912721Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:05.242340Z","title":"Proceedings of the AAAI conference on artificial intelligence , volume=","venue":null,"work_id":"c34b10de-bfd6-4f1e-bf25-c25f477e8d9c","year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.919912Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:ff34713dfd611894eb1026f43f4d2e56f63dc99ff6a97fbf9c53d049bac7b6bb","observation_id":"04b91cd1-f9d0-4668-a81e-76c943241e58","resolution":{"observed_at":"2026-08-12T19:02:05.250596Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1902.08153","last_updated":"2020-05-07T03:30:49Z","snapshot_observed_at":"2026-08-18T09:17:55.370547Z","submitted_at":"2019-02-21T17:31:32Z","title":"Learned Step Size Quantization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1902.08153","snapshot_observed_at":"2026-08-12T19:02:03.928368Z","title":"arXiv preprint arXiv:1902.08153 , year=","venue":null,"work_id":null,"year":1902},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.928368Z"},"links":{"cited_paper":"/paper/1902.08153","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:5022b0726db69c1ee1c4cfb002e668a8b423c938d86ea5087054793270a83409","observation_id":"7ecdeb7f-7892-40fa-8731-7914f2bfd5cc","resolution":{"observed_at":"2026-08-12T19:02:03.928368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:05.214919Z","title":"Proceedings of the IEEE/CVF conference on computer vision and pattern recognition , pages=","venue":null,"work_id":"6cb34edb-5547-4f63-8660-0d7ca6527038","year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.937285Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:3fac58858fdf3518a064b579b7309e584806859874cb92398270ff093bd05bc4","observation_id":"7b3d5278-ac2c-4901-be4c-049174a6fbc5","resolution":{"observed_at":"2026-08-12T19:02:05.221839Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1805.06085","last_updated":"2018-07-17T07:33:19Z","snapshot_observed_at":"2026-08-17T06:12:39.940205Z","submitted_at":"2018-05-16T01:19:43Z","title":"PACT: Parameterized Clipping Activation for Quantized Neural Networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.06085","snapshot_observed_at":"2026-08-12T19:02:03.943335Z","title":"arXiv preprint arXiv:1805.06085 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.943335Z"},"links":{"cited_paper":"/paper/1805.06085","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:74574660aac4150f6c7f891201564de15c65fbf12db0d78e07f3751a6f40b8ad","observation_id":"52703f8b-4737-4b52-8abb-6676a4d4f896","resolution":{"observed_at":"2026-08-12T19:02:03.943335Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1905.11452","last_updated":"2020-05-22T17:02:41Z","snapshot_observed_at":"2026-08-15T23:55:48.282700Z","submitted_at":"2019-05-27T19:03:40Z","title":"Mixed Precision DNNs: All you need is a good parametrization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1905.11452","snapshot_observed_at":"2026-08-12T19:02:03.952100Z","title":"arXiv preprint arXiv:1905.11452 , volume=","venue":null,"work_id":null,"year":1905},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.952100Z"},"links":{"cited_paper":"/paper/1905.11452","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:b491966c8fc8ae888fa45ae1986bd6720fd876a15a67677631928da71bf57823","observation_id":"edb3abc0-cfcd-4a32-84af-41b9b5e7bfe1","resolution":{"observed_at":"2026-08-12T19:02:03.952100Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1909.13144","last_updated":"2020-02-02T14:15:07Z","snapshot_observed_at":"2026-08-19T02:03:55.357515Z","submitted_at":"2019-09-28T20:14:11Z","title":"Additive Powers-of-Two Quantization: An Efficient Non-uniform Discretization for Neural Networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1909.13144","snapshot_observed_at":"2026-08-12T19:02:03.959326Z","title":"arXiv preprint arXiv:1909.13144 , year=","venue":null,"work_id":null,"year":1909},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.959326Z"},"links":{"cited_paper":"/paper/1909.13144","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:4a95f0cb42b4bef7d74ef6c4e947e03c74432aa917866c50b3a4eb3187947a25","observation_id":"14685629-7af2-4cdc-b190-eb23444b5aff","resolution":{"observed_at":"2026-08-12T19:02:03.959326Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:03.967095Z","title":null,"venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.967095Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:da81e272b7de771e47852b30380918cc24b79a2365e8018fd410da79c3dc6bfc","observation_id":"44a0cb39-4bb6-4fd2-a695-9b1c8a9f399e","resolution":{"observed_at":"2026-08-12T19:02:03.967095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:03.976296Z","title":"2009 IEEE conference on computer vision and pattern recognition , pages=","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.976296Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:a79686df38a3a5a1c36a66ad898907ebeaf83c1507b2235d7aad9dc4ddb2e207","observation_id":"e915b8d9-a5e0-4aac-af46-3cad119ea857","resolution":{"observed_at":"2026-08-12T19:02:03.976296Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:03.989187Z","title":"CS 231N , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.989187Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:8f96bedbd3d8e0e29836426c1b1510c1cc3d81ce9f6161186bdecca67a36b4bb","observation_id":"8a80469f-8354-4032-8866-d47c0cacc854","resolution":{"observed_at":"2026-08-12T19:02:03.989187Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:03.997947Z","title":"Proceedings of the IEEE conference on computer vision and pattern recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:03.997947Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:8539197684d3d5b887d66450ac50df3026540c4b7f5e5f4b542dc78aa88f38a5","observation_id":"9f955794-cbc3-4952-ade7-2f07b3a7a419","resolution":{"observed_at":"2026-08-12T19:02:03.997947Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1409.1556","last_updated":"2015-04-10T16:25:04Z","snapshot_observed_at":"2026-08-17T19:17:06.411141Z","submitted_at":"2014-09-04T19:48:04Z","title":"Very Deep Convolutional Networks for Large-Scale Image Recognition","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1409.1556","snapshot_observed_at":"2026-08-12T19:02:04.006401Z","title":"arXiv preprint arXiv:1409.1556 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.006401Z"},"links":{"cited_paper":"/paper/1409.1556","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:c2a0520eb5ea97e748e70bd7aebb183e76e126baf1e1a1ebf881f9cc9b4ea0e8","observation_id":"0c37cd6b-bb24-4db9-b9bc-9fa29d7517ba","resolution":{"observed_at":"2026-08-12T19:02:04.006401Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:04.013329Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.013329Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:8cd115a80b51f8601012bfad4f008f9662e65f076bf0a6fe1390ed559634ca19","observation_id":"e8974964-e5bf-4cea-a0fe-59d6a6a5307d","resolution":{"observed_at":"2026-08-12T19:02:04.013329Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:04.021205Z","title":"The annals of mathematical statistics , pages=","venue":null,"work_id":null,"year":1951},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.021205Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:dfd919f64cdabe261c7d9c422722dcb3c053e8e3b9455c0ccd8dab8b699289cd","observation_id":"6124ae9e-9756-4328-81db-cf26c1be8f74","resolution":{"observed_at":"2026-08-12T19:02:04.021205Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6980","last_updated":"2017-01-30T01:27:54Z","snapshot_observed_at":"2026-08-17T19:26:44.032537Z","submitted_at":"2014-12-22T13:54:29Z","title":"Adam: A Method for Stochastic Optimization","version":9},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.6980","snapshot_observed_at":"2026-08-12T19:02:04.036659Z","title":"arXiv preprint arXiv:1412.6980 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.036659Z"},"links":{"cited_paper":"/paper/1412.6980","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:20b510d76686b375ae40ad13508e1488510bbe31a3d79c09c5c58624adc8eb6b","observation_id":"f4171cb8-44bf-4b14-a276-825ab8c299f0","resolution":{"observed_at":"2026-08-12T19:02:04.036659Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1711.05101","last_updated":"2019-01-04T21:01:49Z","snapshot_observed_at":"2026-08-14T20:13:52.872565Z","submitted_at":"2017-11-14T14:24:06Z","title":"Decoupled Weight Decay Regularization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.05101","snapshot_observed_at":"2026-08-12T19:02:04.045373Z","title":"arXiv preprint arXiv:1711.05101 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.045373Z"},"links":{"cited_paper":"/paper/1711.05101","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:758ea19478e05f9c54ae5b373db3c1f17cffe950498d775e8bb6f809a8f0aa59","observation_id":"37f0b4ff-6c97-46d7-a7dc-b2da33c66a2e","resolution":{"observed_at":"2026-08-12T19:02:04.045373Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1608.03983","last_updated":"2017-05-03T16:28:09Z","snapshot_observed_at":"2026-07-06T05:06:55.589962Z","submitted_at":"2016-08-13T13:46:05Z","title":"SGDR: Stochastic Gradient Descent with Warm Restarts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1608.03983","snapshot_observed_at":"2026-08-12T19:02:04.057443Z","title":"arXiv preprint arXiv:1608.03983 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.057443Z"},"links":{"cited_paper":"/paper/1608.03983","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:183ee6869197fa261dc1a59925771dffd3e7d3d325c0921a628c01ed548747e6","observation_id":"8f9d1296-99f1-4e71-b2f0-aa8dd6761785","resolution":{"observed_at":"2026-08-12T19:02:04.057443Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1510.00149","last_updated":"2016-02-15T06:25:40Z","snapshot_observed_at":"2026-08-04T16:59:47.843960Z","submitted_at":"2015-10-01T09:03:44Z","title":"Deep Compression: Compressing Deep Neural Networks with Pruning, Trained Quantization and Huffman Coding","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1510.00149","snapshot_observed_at":"2026-08-12T19:02:04.069324Z","title":"arXiv preprint arXiv:1510.00149 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.069324Z"},"links":{"cited_paper":"/paper/1510.00149","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:cb686c2ebf5664ce6e66c1f456e943624e66da5120f71f02eaee26f4c644c759","observation_id":"7a0ee266-767d-4722-ae6c-f62f90c296cb","resolution":{"observed_at":"2026-08-12T19:02:04.069324Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:05.049610Z","title":"Expert Systems with Applications , volume=","venue":null,"work_id":"6f620700-996e-495f-b7e2-b5e8e6340898","year":2025},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.079618Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:dbce6247b22ee6e0dbd0967629c6cf0a3e5b6498f89ade298b0ecb21725aef16","observation_id":"3939cdeb-8f41-4d5b-874d-7e6995e298bc","resolution":{"observed_at":"2026-08-12T19:02:05.057986Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:05.009277Z","title":"Proceedings of the 31st ACM International Conference on Multimedia , pages=","venue":null,"work_id":"6c632e91-28ba-4153-8c98-38a03b05d93e","year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.088042Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:ef3e5eedabeb5f8e81902b17d1d8acba2fc29fdf346bd86a0cedb6d15742ff28","observation_id":"7d38b349-8bde-4501-b409-16332a3cc458","resolution":{"observed_at":"2026-08-12T19:02:05.018032Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:04.982852Z","title":"Proceedings of the AAAI conference on artificial intelligence , volume=","venue":null,"work_id":"f1afbbf0-f48a-4af1-aec6-507ca5c99652","year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.098389Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:8b91fe3f9af9660a45c8bdf04122fdf49334c2200096443fb916e409d8492529","observation_id":"bc30f054-fce2-4046-a88e-04dd60147aef","resolution":{"observed_at":"2026-08-12T19:02:04.991117Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-16T09:25:53.087782Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-12T19:02:04.107849Z","title":"arXiv preprint arXiv:2010.11929 , year=","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.107849Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:435b097d20a32b621bce847ae712fc26e9c63da4dd04db6fd535b96fccaf0cb2","observation_id":"ab2a1f24-8266-4a97-8434-7adfbc735158","resolution":{"observed_at":"2026-08-12T19:02:04.107849Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:04.117283Z","title":"International conference on machine learning , pages=","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.117283Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:c331cdc59277418de41be65669503334afae1d7a0a447da203254c399c206137","observation_id":"fd099657-2ebb-4d92-92a6-d3dec0f97cf8","resolution":{"observed_at":"2026-08-12T19:02:04.117283Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:04.126749Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.126749Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:69ff0c16e0bc08886262f3d92bb67eddb31b3c29ece0e85340a38aa10d911fea","observation_id":"c8aac5f8-1b59-40b2-860a-c8bda43eb3ac","resolution":{"observed_at":"2026-08-12T19:02:04.126749Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:04.133609Z","title":"Proceedings of the 2018 EMNLP workshop BlackboxNLP: Analyzing and interpreting neural networks for NLP , pages=","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.133609Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:7753edf642c82cc7abf3de8dde50e3c121f9b6a811809f5c252b414ca7c4272a","observation_id":"5ad38c3c-070e-47a9-b457-b347b9358aa1","resolution":{"observed_at":"2026-08-12T19:02:04.133609Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:04.900993Z","title":"International conference on machine learning , pages=","venue":null,"work_id":"006980fb-57ff-4d2d-a7a4-fc7852ed804b","year":2023},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.142531Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:48f7b3a549f4caf3157a06c3992f993105505bf6ac16a9a9c2aa308dd5023922","observation_id":"5725ee66-dbd9-468a-86e5-940746a1c776","resolution":{"observed_at":"2026-08-12T19:02:04.906622Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1605.04711","last_updated":"2022-11-20T14:21:55Z","snapshot_observed_at":"2026-08-14T21:57:14.951937Z","submitted_at":"2016-05-16T10:21:25Z","title":"Ternary Weight Networks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1605.04711","snapshot_observed_at":"2026-08-12T19:02:04.149749Z","title":"arXiv preprint arXiv:1605.04711 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.149749Z"},"links":{"cited_paper":"/paper/1605.04711","citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:136e855155976bdb4e26e0447a280bfc177d822943cc6449062ba8b4c78fc843","observation_id":"35ce7163-7816-4fa3-9e6f-74cc7b6784ad","resolution":{"observed_at":"2026-08-12T19:02:04.149749Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:02:04.858613Z","title":"Proceedings of the AAAI conference on artificial intelligence , volume=","venue":null,"work_id":"f80bb725-ce50-4026-85e4-6fd82094003e","year":null},"citing_paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-12T19:02:04.157451Z"},"links":{"citing_paper":"/paper/2608.10709"},"observation_digest":"sha256:eabc50c5a4de4381cb02cc5ac68d564f51acd3c5d893bb96b48f2afc840e8f9c","observation_id":"be9eecf2-5ba9-43dc-a34a-12e71745e0d7","resolution":{"observed_at":"2026-08-12T19:02:04.878681Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2608.10709","last_updated":"2026-08-11T09:29:30Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-19T22:32:56.373060Z","submitted_at":"2026-08-11T09:29:30Z","title":"SQuaT: Self-Supervised Knowledge Distillation via Student-Aware Quantized Teacher Features"},"reference_resolution":{"displayed":43,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":32,"verified_exact":0,"verified_fuzzy":11},"total_outbound_references":43},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 43 of 43 outbound references and 0 inbound Pith citation observations for arXiv:2608.10709."}