{"as_of":"2026-08-08T02:30:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:bd92e2e44de5bf5454864726ea513487adf9724f2de612a813792bbd000c8615","coverage":[{"denominator":81,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":81,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:06:30.520235Z","state":"measured"},{"denominator":81,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":81,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.16346/citation-record","integrity":"/paper/2505.16346/integrity","json":"/paper/2505.16346/citation-record.json","paper":"/paper/2505.16346"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:29.398472Z","title":"Visualizing size of large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.398472Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:d1aaeb417f9473b37768fe88ee9c9ff842a8151acc658a397b0dc8b8f69a55e4","observation_id":"c44066fd-fa82-459e-ac9a-e92d9805bfd9","resolution":{"observed_at":"2026-08-07T15:06:29.398472Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:32.133099Z","title":"Trends in deep learning hardware,","venue":null,"work_id":"d1045393-f907-4728-bd05-4eb40bd39f8f","year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.403476Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:5a755d0ae5ac08c9dc54f632de6a4e732a72d4e5c8239b597654b52516dfd435","observation_id":"14d25844-d249-4f4c-977e-e25c5c79aa5c","resolution":{"observed_at":"2026-08-07T15:06:32.138137Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:32.109175Z","title":"Roofline: an insightful visual performance model for multicore architectures,","venue":null,"work_id":"85576604-dc2d-45f8-8ebb-db60f7b42c93","year":2009},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.408467Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:a0fbe52400c6776629cf46b1ef9fa3c7eeb1220498bb42c8c82cf640b35bfbf5","observation_id":"7086185c-8b13-4bc9-9546-2fd5afc2a679","resolution":{"observed_at":"2026-08-07T15:06:32.113461Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:32.091560Z","title":"Eyeriss: A apatial architecture for energy-efficient dataflow for convolutional neural networks,","venue":null,"work_id":"89ba89f3-bc71-42a7-a58e-241a0f766ad9","year":2016},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.412848Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:f00850526b0ceccfa4aa4f35b16cb907dec760ec21feaf1fb27580cafbe8d1e7","observation_id":"83c26e29-7cd7-4049-b5c9-4c43790a0f65","resolution":{"observed_at":"2026-08-07T15:06:32.096760Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:32.072440Z","title":"Roofline performance analysis of dnn architectures on cpu and gpu systems,","venue":null,"work_id":"cee3300d-8405-478b-96b3-a57f02dc1b1b","year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.419412Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:1fc5227567d7c0aeaf10231ace7530e0df419e3135e16bd47c298f3c5982776a","observation_id":"dba73b76-278c-45fe-8064-49a7104e97cc","resolution":{"observed_at":"2026-08-07T15:06:32.079775Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:32.040261Z","title":null,"venue":null,"work_id":"71b893ec-1a9b-4357-90d7-3a53a48693c7","year":2010},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.424820Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:545c5a78b466a6d7a9d81c135b174da3b4b92fc8571a172c1a1f3ee5c2e5d30d","observation_id":"3bc34efd-f60a-415e-b5a8-0faad797ffe0","resolution":{"observed_at":"2026-08-07T15:06:32.048841Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:32.023615Z","title":"Understanding reuse, performance, and hardware cost of dnn dataflow: A data-centric approach,","venue":null,"work_id":"56db4f99-a288-4372-9767-818fe8f05615","year":2019},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.438721Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:e8090499fcbc702f2a9a856ab9a200a6db32886cb700a522d05c4b454d406fda","observation_id":"5ebbde3f-f00f-4f35-a464-fe03ea9f0e38","resolution":{"observed_at":"2026-08-07T15:06:32.028318Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:32.002273Z","title":"A roofline model of energy,","venue":null,"work_id":"5e3f6132-7bfb-4b95-bea3-47e91e7e6f3a","year":2013},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.443914Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:8e2c761c0ef2720495c698bea04168944c78dfe65d355d001942869b44fd7cb7","observation_id":"191aa578-ee30-4289-a75c-2836b4ec5a45","resolution":{"observed_at":"2026-08-07T15:06:32.008663Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.986646Z","title":"Symphony: Orchestrating sparse and dense tensors with hierarchical heterogeneous processing,","venue":null,"work_id":"87eb65ca-e48c-4644-93b4-b663adaa55a7","year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.449801Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:ecb3a2f0ac221134b3a6b5282af509a79dab435b5e6b40731637583566f6d82c","observation_id":"2db46215-a0bb-4c48-84f7-8a8d5ad0049d","resolution":{"observed_at":"2026-08-07T15:06:31.991855Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.970967Z","title":"Lots of questions on Google’s “Trillium","venue":null,"work_id":"70e02035-106e-44f1-886e-1ea03e3a8760","year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.455217Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:58f2fe4a0d7e663631c22531b13b0ac365d72b738ad7f3408e8fcfe839f0027d","observation_id":"79867109-ebbf-4d1d-8b49-b0c9761378ec","resolution":{"observed_at":"2026-08-07T15:06:31.976054Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.954391Z","title":"Envi- sion: A 0.26-to-10tops/w subword-parallel dynamic-voltage-accuracy- frequency-scalable convolutional neural network processor in 28nm fdsoi,","venue":null,"work_id":"903fc061-463e-4031-9948-548ab85f0a54","year":2017},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.461094Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:ab2a505db13642e6bdc650026e925ab5519df77da7919d8fa1cb00fab84dea17","observation_id":"fcbb83b6-24a5-426b-b612-a27c077aa0f7","resolution":{"observed_at":"2026-08-07T15:06:31.959854Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.933535Z","title":"9.5 a 6k-mac feature-map-sparsity-aware neural processing unit in 5nm flagship mobile soc,","venue":null,"work_id":"111c151b-1fbe-47fe-ac7d-f0a3f986c572","year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.468026Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:798a2e38c9c77d3598f075c2373139813f40139e633f9fca694883b29748c1b6","observation_id":"03919dc5-ba5d-4c84-a46c-846a3f0ba0cc","resolution":{"observed_at":"2026-08-07T15:06:31.939247Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.916665Z","title":"Compute solution for tesla’s full self-driving computer,","venue":null,"work_id":"b8032d20-e006-40f4-a9af-0cbc59778518","year":2020},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.472900Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:db5a501f347399c9f0ab8ae7c70c28b2f3338700f3541f3d36ac42046cd1c7e2","observation_id":"563c3ace-a04a-40e0-b028-5fc467f91985","resolution":{"observed_at":"2026-08-07T15:06:31.921995Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.897477Z","title":"7.2 a 12nm programmable convolution-efficient neural- processing-unit chip achieving 825tops,","venue":null,"work_id":"38e43981-b595-42b3-8c7f-7e5472fbb4cd","year":2020},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.477473Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:50e7bea61e8949706ff1efed511ac4bc3b2523aa014497afd35e97f63b05f851","observation_id":"a7c94c0d-2efb-4268-8c84-19ebc1307c38","resolution":{"observed_at":"2026-08-07T15:06:31.902241Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.871981Z","title":"Groq rocks neural networks,","venue":null,"work_id":"d7a6453a-771d-4f5c-af87-7e6dbbaa226f","year":2020},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.484255Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:c3f2c408642f7c6d7121ef87422104b956e2f7b900d7219f266569ee8f0144e0","observation_id":"d74fe782-4893-4413-82b4-4950a2811a22","resolution":{"observed_at":"2026-08-07T15:06:31.883652Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.854499Z","title":"9.1 a 7nm 4-core ai chip with 25.6tflops hybrid fp8 training, 102.4tops int4 inference and workload-aware throttling,","venue":null,"work_id":"446b218c-5962-4632-8b83-e4006f353de6","year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.491187Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:32b1c68dbb6228b77260e4e3d226a32bf8f4cbf7037230e8101bbaac4d9872e1","observation_id":"db943f5f-3a0c-4ed3-b6d0-bdf3d7e97261","resolution":{"observed_at":"2026-08-07T15:06:31.859940Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.834726Z","title":"16.7 a 40-310tops/w sram-based all-digital up to 4b in-memory computing multi-tiled nn accelerator in fd-soi 18nm for deep-learning edge applications,","venue":null,"work_id":"0391d17e-67a8-4c49-a690-1196e2903b12","year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.501017Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:c95d5a75652cc11228ecf2bafc8c625812cbbe6b3cf233f5a27b88f177856990","observation_id":"21688fde-3472-4885-9133-5e67b0f4bbc7","resolution":{"observed_at":"2026-08-07T15:06:31.841139Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:29.512744Z","title":"Charm: Composing heterogeneous accelerators for matrix multiply on versal acap architecture,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.512744Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:77a39cedc1049b2fe7e5dd87c90519fbe764a5aa142c24379eda4356d4eeb938","observation_id":"64e534e9-7d1c-4f72-aba6-e2a83c65919d","resolution":{"observed_at":"2026-08-07T15:06:29.512744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.815928Z","title":"Davinci: A scalable architecture for neural network computing,","venue":null,"work_id":"c9e28205-c424-4f84-b5c8-e7022bbf2355","year":2019},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.523587Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:5217a178203d84e0e2e88fc5f349034ee560fca51103398beee59fbe587257c9","observation_id":"074e189b-9dd1-4e3d-bd31-b2e6bb4362eb","resolution":{"observed_at":"2026-08-07T15:06:31.821177Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:29.538744Z","title":"Nvidia tensor core programmability, performance & precision,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.538744Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:b97fe986421887386fd2ba0adf6ab068280187dfea2859c6cc9f33817312dd9f","observation_id":"70b5df40-95fa-47b8-b99e-193403ae7e95","resolution":{"observed_at":"2026-08-07T15:06:29.538744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:29.557111Z","title":"A charge domain sram compute-in-memory macro with c-2c ladder- based 8-bit mac unit in 22-nm finfet process for edge inference,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.557111Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:f9ab4be02a163aeef17bd8f8524468d598fe0813dbc844b6fd5947b2549894a6","observation_id":"a7876e8c-e647-4b9b-b903-617d5ecc35ae","resolution":{"observed_at":"2026-08-07T15:06:29.557111Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.767622Z","title":"A 22 nm, 1540 top/s/w, 12.1 top/s/mm 2 in-memory analog matrix-vector-multiplier for dnn acceleration,","venue":null,"work_id":"eb51f8e1-81fc-4949-9767-1cfdccb189d5","year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.584583Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:d6fcdbfb2f0f70f87cbc7945398bcc213c82a4a2becc1e89492529e67982a222","observation_id":"8aa39835-d1be-4fdd-968a-736471b5bba6","resolution":{"observed_at":"2026-08-07T15:06:31.775674Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.746417Z","title":"A 64-tile 2.4- mb in-memory-computing cnn accelerator employing charge-domain compute,","venue":null,"work_id":"87e9056e-9ba0-4a21-bd0a-b0f9cb4568fc","year":2019},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.613354Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:6fc1500465920b154a827108e69eba1317d1daf7f0c56ad5b348cb9bb24c0c81","observation_id":"060dfa3f-6df2-41e0-8d36-efaafad155e2","resolution":{"observed_at":"2026-08-07T15:06:31.752647Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.727534Z","title":"Compute Solution for Tesla’s Full Self-Driving Computer,","venue":null,"work_id":"ad14f055-f5a0-4da3-88ed-d079ce6fadc1","year":2020},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.643554Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:2951d70f67d4f0ba2e2abad28ede475aaaedebb90698cbc7804fde9790718f47","observation_id":"0b0afbc5-c8bf-4721-b4fb-82278323b147","resolution":{"observed_at":"2026-08-07T15:06:31.733368Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.711944Z","title":"Hardware for deep learning,","venue":null,"work_id":"f39df72f-99aa-4aad-a142-c4ce67ac4b03","year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.669402Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:290b61c1a73c12b247bb07f48ec78652102c506cb529933be579eca95fa2dfef","observation_id":"fae8ddfc-2b2d-4e3a-9f0c-497ee270e85d","resolution":{"observed_at":"2026-08-07T15:06:31.715995Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.695889Z","title":"Lincoln ai computing survey (laics) update,","venue":null,"work_id":"eefb6414-44d8-43a6-9f89-86cd35c77982","year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.691612Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:6e9c84117941d68d7c991715fab008736ac94e6e91fa1b73353b054a6fe8ac55","observation_id":"da16470b-2d3e-4fe0-84a9-b071dd26d719","resolution":{"observed_at":"2026-08-07T15:06:31.701147Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.677417Z","title":"Neural network accelerator comparison","venue":null,"work_id":"7d00c15b-dc64-4389-a17f-69b24d5ba344","year":null},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.702506Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:35279625a16b1462dde044256a1c82633253e6d3859caa11d84f0ce67bbe60b3","observation_id":"901acd9a-d4f1-4ea0-ac01-b5ca945a06b6","resolution":{"observed_at":"2026-08-07T15:06:31.682603Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.16363","last_updated":"2024-05-01T20:42:28Z","snapshot_observed_at":"2026-08-01T22:52:35.092898Z","submitted_at":"2024-02-26T07:33:05Z","title":"LLM Inference Unveiled: Survey and Roofline Model Insights","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.16363","snapshot_observed_at":"2026-08-07T15:06:29.707975Z","title":"Llm inference unveiled: Survey and roofline model insights,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.707975Z"},"links":{"cited_paper":"/paper/2402.16363","citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:8ed192211bda0fae624b67a29f7f4ad2cec1b21bc26aa528253043c867232390","observation_id":"97d55eaf-e1ea-4233-aa9f-aea58d0b400a","resolution":{"observed_at":"2026-08-07T15:06:29.707975Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.662490Z","title":"Minifloats on risc-v cores: Isa extensions with mixed- precision short dot products,","venue":null,"work_id":"78cb8c93-f529-42f6-aada-72ee83e1edf1","year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.712718Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:672d0dea520c1491c9e7471908ad2b27b0449fca5d0e09c72af0585c50339e52","observation_id":"08353473-84cd-4cb4-bd49-537fa186db6a","resolution":{"observed_at":"2026-08-07T15:06:31.666883Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.645937Z","title":"Cutie: Beyond petaop/s/w ternary dnn inference acceleration with better-than-binary energy efficiency,","venue":null,"work_id":"6295f29e-8332-4758-ac5d-4dbe5b4ac1f5","year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.719228Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:37ad6c47acba7aff9572d63db803d5f70683d05cb296806d23f3c215a3c12d88","observation_id":"8555f8ba-09e2-4b4a-af27-409f6d015b63","resolution":{"observed_at":"2026-08-07T15:06:31.651276Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.629906Z","title":"Binareye: An always-on energy-accuracy-scalable binary cnn processor with all memory on chip in 28nm cmos,","venue":null,"work_id":"c403e61f-59ba-46d1-b153-b8e5444ef82e","year":2018},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.727130Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:8cdb939cac79e4203aeafb302e6cca361353b8a7bfc393c1e08dfe74c34871ac","observation_id":"0e670c07-2879-40bf-b013-1f9a3a083047","resolution":{"observed_at":"2026-08-07T15:06:31.635016Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.11453","last_updated":"2023-10-17T17:59:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-17T17:59:15Z","title":"BitNet: Scaling 1-bit Transformers for Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.11453","snapshot_observed_at":"2026-08-07T15:06:29.752659Z","title":"Bitnet: Scaling 1-bit transformers for large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.752659Z"},"links":{"cited_paper":"/paper/2310.11453","citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:a889d5add1dea437f8ca4e6f7459a6abc64a845918b6b7513ac5382f5a4783fd","observation_id":"0a3287c9-08c5-4d5d-a841-4c814fd7ab59","resolution":{"observed_at":"2026-08-07T15:06:29.752659Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.611362Z","title":"A 3 tops/w risc-v parallel cluster for inference of fine-grain mixed-precision quantized neural networks,","venue":null,"work_id":"6cf30ccf-19ab-4220-8fc8-5dc429f2d5e6","year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:29.907854Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:4bab3e398a635b5e63e17cc92b4a3d16f01c0437a903e8fc72e280122c741eec","observation_id":"1519180e-97b3-4911-bfc5-9e5d87d183b3","resolution":{"observed_at":"2026-08-07T15:06:31.615689Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.596589Z","title":"Marsellus: A heterogeneous risc-v ai-iot end-node soc with 2–8 b dnn acceleration and 30%-boost adaptive body biasing,","venue":null,"work_id":"9a4d1053-0c3f-478f-8614-bb44901a175b","year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.045702Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:2be00b78c6967a5a015e2642acbaf39db470fff64c8f1e274ec6662798368166","observation_id":"0ba24f90-6cda-4cec-98bc-13dd8c9fa289","resolution":{"observed_at":"2026-08-07T15:06:31.601169Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.10537","last_updated":"2023-10-19T16:38:33Z","snapshot_observed_at":"2026-08-08T01:04:32.134805Z","submitted_at":"2023-10-16T16:07:41Z","title":"Microscaling Data Formats for Deep Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.10537","snapshot_observed_at":"2026-08-07T15:06:30.167879Z","title":"Microscaling data formats for deep learning,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.167879Z"},"links":{"cited_paper":"/paper/2310.10537","citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:a37046ecc765e4554356dfe853713f7ea22dc1d0d334b952619534d42c5bce31","observation_id":"99a00f12-8207-439c-b6f2-350163c31afb","resolution":{"observed_at":"2026-08-07T15:06:30.167879Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.581513Z","title":"Nvidia blackwell platform: Advancing generative ai and accelerated computing,","venue":null,"work_id":"83c8a462-256e-4b8f-b5d0-0001580b04ad","year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.260534Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:c8fdb5142be6372cff47c10f808d93d9ef944f46b2e209258f24ccb10db0c2ff","observation_id":"2cf05c9b-d6e0-4f88-98ba-393fd718c14d","resolution":{"observed_at":"2026-08-07T15:06:31.585901Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.566755Z","title":"Siracusa: A 16 nm heterogenous risc-v soc for extended reality with at-mram neural engine,","venue":null,"work_id":"7b7976e5-6b7a-42f8-854c-ce6e5b5a34db","year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.295340Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:403db8c2898abc94e0a6fd130984d21de004cddcc0a50235d03dbaf7d65eb7d5","observation_id":"1680c9ec-81fd-4344-8309-0f0dbff0fc78","resolution":{"observed_at":"2026-08-07T15:06:31.571402Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.552604Z","title":"Onyx: A 12nm 756 gops/w coarse-grained reconfigurable array for accelerating dense and sparse applications,","venue":null,"work_id":"0db87afa-a160-4c69-b6f7-402281de7955","year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.311228Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:5a964da42ee7556b37f6d39317a76ccb096f109aa0ecc0e5ef86338c9a27b558","observation_id":"3fe2bff3-6cba-4d4f-9df5-377dcb967997","resolution":{"observed_at":"2026-08-07T15:06:31.557095Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.04010","last_updated":"2021-04-18T10:18:00Z","snapshot_observed_at":"2026-08-04T18:04:45.224844Z","submitted_at":"2021-02-08T05:55:47Z","title":"Learning N:M Fine-grained Structured Sparse Neural Networks From Scratch","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.04010","snapshot_observed_at":"2026-08-07T15:06:30.316475Z","title":"Learning n: m fine-grained structured sparse neural networks from scratch,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.316475Z"},"links":{"cited_paper":"/paper/2102.04010","citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:4ac536ed3d1ed9748854ad69b236da337b4e6e82c8516ffee5345bae72700517","observation_id":"811baf5b-cbc8-40a8-b6b8-58e74ad81a01","resolution":{"observed_at":"2026-08-07T15:06:30.316475Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.538617Z","title":"3.2 the a100 datacenter gpu and ampere architecture,","venue":null,"work_id":"b94b18cd-d48f-49a3-8168-e57cef4b23f2","year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.321479Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:8dc86611a0ed7e0c02e847bc43999f0154b0de3233c67f97b6af59c2d785eec1","observation_id":"fd4078a9-32f3-41e9-9fcb-d49e353faaa0","resolution":{"observed_at":"2026-08-07T15:06:31.543279Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.522188Z","title":"Venom: A vectorized n: M format for unleashing the power of sparse tensor cores,","venue":null,"work_id":"0a87cbb6-f541-481b-9ee1-085706be9494","year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.327263Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:0fe123b6afc30ec04c2d122d717b62ef5f75521e9e8d1592895acfef79181374","observation_id":"1b9ff632-a182-4763-8a18-bfe3f114406f","resolution":{"observed_at":"2026-08-07T15:06:31.527977Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.506919Z","title":"Occamy: A 432-core dual-chiplet dual-hbm2e 768-dp-gflop/s risc-v system for 8- to-64-bit dense and sparse computing in 12-nm finfet,","venue":null,"work_id":"a07343aa-3cfc-4e88-8ca6-e0efbc5ef8f1","year":2025},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.332379Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:f856527f8ee2a17c83359d7d5474d98e86741475ee0e01170937da01c24bff59","observation_id":"442b8fd4-7f49-4dfd-bff1-240ef09e13e4","resolution":{"observed_at":"2026-08-07T15:06:31.511828Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.338151Z","title":"Neupims: Npu-pim heterogeneous acceleration for batched llm inferencing,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.338151Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:2387dddd07ea4eef6d062f739535e507ec1f2acb3a01873293da84ce1bfa732d","observation_id":"a916b1ac-aac4-455c-b223-d85dd214cb28","resolution":{"observed_at":"2026-08-07T15:06:30.338151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.07984","last_updated":"2024-01-17T16:00:27Z","snapshot_observed_at":"2026-07-06T16:18:41.478772Z","submitted_at":"2023-09-14T18:42:29Z","title":"Inclusive-PIM: Hardware-Software Co-design for Broad Acceleration on Commercial PIM Architectures","version":3},"cited_work":{"arxiv_id":"2309.07984","doi":null,"metadata_source":"pith","pith_arxiv_id":"2309.07984","snapshot_observed_at":"2026-08-07T15:06:30.768225Z","title":"Inclusive-PIM: Hardware-Software Co-design for Broad Acceleration on Commercial PIM Architectures","venue":"cs.AR","work_id":"9b3ec8fc-5a9b-4a20-9553-f0681db36ed8","year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.344978Z"},"links":{"cited_paper":"/paper/2309.07984","citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:dd577e218990526a1ba5893a11c7c882783d889723f2dd40f1aef1cc9b534abd","observation_id":"17058cdc-295b-4200-86ac-6371834add2f","resolution":{"observed_at":"2026-08-07T15:06:30.773144Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.490662Z","title":"In-memory computation of a machine-learning classifier in a standard 6t sram array,","venue":null,"work_id":"2c32c6df-7c81-4d80-b767-ce4ffb522d90","year":2017},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.350099Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:a833373edfb8394e108140f7f224ed4c2b2cb536cf63d6f38f1ab46248fc5b26","observation_id":"c504234d-02c1-4508-a794-fdaf31d7348c","resolution":{"observed_at":"2026-08-07T15:06:31.496919Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.474383Z","title":"An energy-efficient memory-based high-throughput vlsi architecture for convolutional networks,","venue":null,"work_id":"185713ea-d50a-489f-b077-83e731402768","year":2015},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.354542Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:7e856421700a466d7a5b1915247c37be8f2c62d7219e6d1d9f619a7e28173aec","observation_id":"eabae999-f60c-4fe5-bc52-73e6730bdc6a","resolution":{"observed_at":"2026-08-07T15:06:31.479682Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.452145Z","title":"Fast, energy-efficient, robust, and reproducible mixed-signal neuromorphic classifier based on embedded nor flash memory technology,","venue":null,"work_id":"f79bdc82-6316-4515-a13b-08ef5595256a","year":2017},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.359436Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:f4f59341248bfc587f701554bf484446762e53f569d8c3336ca46444c0f88d8a","observation_id":"dc024c9b-6911-468e-a603-4c7ed3dba6b1","resolution":{"observed_at":"2026-08-07T15:06:31.457417Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.436387Z","title":"Analog in-memory subthreshold deep neural network accelerator,","venue":null,"work_id":"34538f09-bfed-4505-abc9-16b3d8cedc48","year":2017},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.364825Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:638f42482a771bfebe74efc88a27ef94680491c52130f7a7d57ba65aaf393a94","observation_id":"1983e91c-efc9-4b7e-9918-6b05794fdb74","resolution":{"observed_at":"2026-08-07T15:06:31.441396Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.420073Z","title":"A 5-nm 254-tops/w 221-tops/mm2 fully-digital computing-in-memory macro supporting wide-range dynamic-voltage- frequency scaling and simultaneous mac and write operations,","venue":null,"work_id":"2b20b492-751d-4890-bf48-efc3ee9b862d","year":2022},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.369284Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:ccffa55088df9a01ddca328b62b5582548a9fada0335385f30f4964425431ca3","observation_id":"a97abc10-7d55-49d8-8883-98a46308033d","resolution":{"observed_at":"2026-08-07T15:06:31.425267Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.405423Z","title":"16.4 an 89tops/w and 16.3tops/mm2 all-digital sram-based full-precision compute-in memory macro in 22nm for machine-learning edge applications,","venue":null,"work_id":"e0cc53a5-f23d-4e3d-8098-c3574966938f","year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.373554Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:2a3c869ff717766d9d69304107abba66d6cd9ce91dc2fa05878ebe52ea5dcd4b","observation_id":"dc0712c1-5f9f-4572-8f47-94f1c6f54a44","resolution":{"observed_at":"2026-08-07T15:06:31.410366Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.389835Z","title":"A maximally row- parallel mram in-memory-computing macro addressing readout circuit sensitivity and area,","venue":null,"work_id":"21f36dcc-86ef-4b03-b3ca-b231b21a4613","year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.378539Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:c8332ef5a92706f71125f512f22210a1261d26342106c10073a8bfbcfe520f90","observation_id":"db9547c9-63ff-48aa-9301-808349548048","resolution":{"observed_at":"2026-08-07T15:06:31.395141Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.383254Z","title":"A programmable heterogeneous microprocessor based on bit-scalable in-memory comput- ing,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.383254Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:fa16c1c094784df29ad9356bd15cc7bfce4dad256dcdd008b43c44df104882b9","observation_id":"fd9d4dde-1f9e-452b-994a-62a1a717dc71","resolution":{"observed_at":"2026-08-07T15:06:30.383254Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.387652Z","title":"A crossbar array of magnetoresistive memory devices for in-memory computing,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.387652Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:19e4d8141904b96443720e01ca81e7eba5037e4d58b19aa4ff590c900e80c690","observation_id":"6393fdbe-e573-4a73-bd6a-fc229cb2005d","resolution":{"observed_at":"2026-08-07T15:06:30.387652Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.392346Z","title":"In-memory computing: Advances and prospects,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.392346Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:8627e7b4a7125574061784540fe16d3f769c771b832ab9fbf38474693e7e5f92","observation_id":"7cd66dde-bc47-49df-a0eb-f1c9e5eba12e","resolution":{"observed_at":"2026-08-07T15:06:30.392346Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.348748Z","title":"14.2 a compute sram with bit-serial integer/floating-point operations for programmable in-memory vector acceleration,","venue":null,"work_id":"3bcf5754-83e7-4915-933e-1bcc86ce4db8","year":2019},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.397269Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:0d659c727c04626f2a53331b0bd3ea2cf842391e828f09d0065ac45b518e4df0","observation_id":"37ddb158-edd6-4020-8d8a-5accaacf7207","resolution":{"observed_at":"2026-08-07T15:06:31.354015Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.331705Z","title":"A 40nm 64kb 26.56tops/w 2.37mb/mm2rram binary/compute-in-memory macro with 4.23x im- provement in density and > 75% use of sensing dynamic range,","venue":null,"work_id":"647217d1-76e6-4bb7-bdb8-3f116cba8cd3","year":2022},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.402092Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:798e0fa9a3e4ca09f44059fe3f05aeee5426a8f5be24a0933d130f190a4bb565","observation_id":"c4466db2-cdb4-4a65-b9c0-9ca6dfc39acd","resolution":{"observed_at":"2026-08-07T15:06:31.337684Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.313148Z","title":"Funda- mental limits on energy-delay-accuracy of in-memory architectures in inference applications,","venue":null,"work_id":"4605a889-120f-4178-8329-f6fb2c81a7dc","year":2022},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.406524Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:7db085a3d1924b13b8abbb7649b4bc2aa05cad04957b1acc8437f2e4afc8eccf","observation_id":"04b325b6-8759-454e-9434-474b467266b1","resolution":{"observed_at":"2026-08-07T15:06:31.318288Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.298493Z","title":"11.3 metis aipu: A 12nm 15tops/w 209.6tops soc for cost- and energy-efficient inference at the edge,","venue":null,"work_id":"6a726b6b-35d8-4435-958c-02382c93209e","year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.410836Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:df1ecb325ab540e104d864747c033a27d669543075c22ad1d0d8f11f667003a9","observation_id":"64a7278b-04c8-4782-a5a6-d54aab5e3df7","resolution":{"observed_at":"2026-08-07T15:06:31.302813Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.415453Z","title":"Benchmarking in-memory computing architectures,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.415453Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:8cef1ed96972e9f2937f679e16a7a43f8ef463cba3ed66ebf38c53f33838f63b","observation_id":"1f3e82c2-4fd0-45b6-be87-3f584cd7db5a","resolution":{"observed_at":"2026-08-07T15:06:30.415453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.272534Z","title":"A 22nm 128-kb mram row/column-parallel in-memory computing macro with memory- resistance boosting and multi-column adc readout,","venue":null,"work_id":"7b99787e-c789-4e2d-9034-75e477c8c9a0","year":2022},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.420737Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:64e4a2c3842fffe005f32a53ae68e39fcb51203fde4b36d326507802b254bab5","observation_id":"c2fb4c15-99b8-4560-b5be-7029727d1d2e","resolution":{"observed_at":"2026-08-07T15:06:31.278613Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.425731Z","title":"A 64-core mixed-signal in-memory compute chip based on phase-change memory for deep neural network inference,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.425731Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:beaa01db22f20a5355142ee3c9d0098eb527875786f0cc6f970aa352131fabcc","observation_id":"8af9691b-3084-4a53-9b91-dc7aaad9beaf","resolution":{"observed_at":"2026-08-07T15:06:30.425731Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.244984Z","title":"An n40 256k×44 embedded rram macro with sl-precharge sa and low-voltage current limiter to improve read and write performance,","venue":null,"work_id":"40390da2-0096-4ed6-b23d-3d5c087097b6","year":2018},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.430911Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:9eee6470707f856a4e1bf96018004c559dafc1eabb61df39bb818b79f5a42b78","observation_id":"e721fac0-ee09-4f4e-b3ae-d4b9c7e2f8d7","resolution":{"observed_at":"2026-08-07T15:06:31.250233Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.230478Z","title":"Cmos- embedded stt-mram arrays in 2x nm nodes for gp-mcu applications,","venue":null,"work_id":"6c64404a-070d-4101-aa41-7ae3318bdacf","year":2017},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.435960Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:9708c0644695629ff99dcc533e880386ec93b44ddd217ee0e5d3ae707911bfdc","observation_id":"57ef3e98-bfbe-45af-9fe0-48b39b3e2f75","resolution":{"observed_at":"2026-08-07T15:06:31.234950Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.215631Z","title":"A switched-capacitor sram in-memory computing macro with high-precision, high-efficiency differential archi- tecture,","venue":null,"work_id":"978a6910-72de-4949-a3d8-69399bd198a8","year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.441328Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:3a43c3a756db6dc69bcdc5cc5cedcfea38be781ca81661939dfcfbcf07bde10c","observation_id":"847322d6-5d32-4d10-b9f9-6b0dceb45256","resolution":{"observed_at":"2026-08-07T15:06:31.220998Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.201911Z","title":"Scalable and Programmable Neural Network Inference Accelerator Based on In-Memory Computing,","venue":null,"work_id":"a76c3ce7-e741-45bc-af06-aa9b1d78ec51","year":2022},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.445701Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:9d86b5b3a34f82db84648c0baba838dd09b9c6d871b6c928f8199d46051861a3","observation_id":"499a3c28-1233-4957-b69f-4fa88b284e65","resolution":{"observed_at":"2026-08-07T15:06:31.206563Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.450465Z","title":"Interstellar: Using Halide’s Scheduling Language to Analyze DNN Accelerators,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.450465Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:baa85ba4c7c61b3b6a488a4d1c764716fba03dad6c573459d260a6bbb5d42992","observation_id":"a78fcd0f-2f3d-4bea-892a-66ee8ff32747","resolution":{"observed_at":"2026-08-07T15:06:30.450465Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.188459Z","title":"MAESTRO: A Data-Centric Approach to Understand Reuse, Performance, and Hardware Cost of DNN Mappings,","venue":null,"work_id":"88de8691-d96c-4dcd-9500-dfe06388b318","year":2020},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.454805Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:0fc6bfd9509b130fed412f73de029d54e1755e550fd15d49abb2b770db064b37","observation_id":"d5106f5f-7b45-4e8f-8825-d07f2c8fa3a4","resolution":{"observed_at":"2026-08-07T15:06:31.192491Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.174162Z","title":"Timeloop: A Systematic Approach to DNN Accelerator Evaluation,","venue":null,"work_id":"af03fa35-4f56-49de-9354-2a5566d0b7a8","year":2019},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.459381Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:677792c60d998d0b2c9c3577009baa312eadab6040859c73c4a75a924a525596","observation_id":"abdc799a-4fe5-46f0-9623-1de4d6bb730f","resolution":{"observed_at":"2026-08-07T15:06:31.179380Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.158362Z","title":"ZigZag: Enlarging Joint Architecture-Mapping Design Space Exploration for DNN Accelerators,","venue":null,"work_id":"6c5172ee-6863-4bc5-bebb-e61d1b23c54d","year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.463852Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:98cede17ff67ee76a628fe2e593e54e4070e2df5fa233985191e6f32995e858e","observation_id":"b8576801-923c-4513-acf6-f513d0ae90d5","resolution":{"observed_at":"2026-08-07T15:06:31.163852Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.142493Z","title":"CoSA: Scheduling by constrained op- timization for spatial accelerators,","venue":null,"work_id":"3a0ec4de-bd50-4bff-8672-583fdb6244c2","year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.471518Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:cf9ebd0c86f01cb32c30eeb708b27c7adfaf3fc16538fc6100387f814c1ff10f","observation_id":"2cc2e601-e74a-4511-9e51-3ef07c37610c","resolution":{"observed_at":"2026-08-07T15:06:31.147252Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.476277Z","title":"Mind Mappings: Enabling Efficient Algorithm-Accelerator Mapping Space Search,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.476277Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:01e6f1abbe0ffd0ee17245779b02ac6b86463c9e774a7768251f761f768e7f99","observation_id":"b6ac7aae-c0da-4d9c-a770-63efc2c79417","resolution":{"observed_at":"2026-08-07T15:06:30.476277Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1145/3400302","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.559600Z","title":"GAMMA: Automating the HW Mapping of DNN Models on Accelerators via Genetic Algorithm,","venue":null,"work_id":"633c9725-544b-4fb5-8fb5-d692e3aeefc0","year":2020},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.480952Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:bca585a84ef68e18a572d2e63088ce6f4e6374d6acc78db08559513c0c2a6e30","observation_id":"8ce0af98-da9c-489c-9b97-929ae3a283eb","resolution":{"observed_at":"2026-08-07T15:06:30.566889Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.127388Z","title":"Stream: Design space exploration of layer-fused dnns on hetero- geneous dataflow accelerators,","venue":null,"work_id":"314e3aaf-c6a6-42f3-a07b-4873858a8786","year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.485112Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:6aa646a9393260a57412e10cc2953621f33c3cc0b3d66a01442772a23c568156","observation_id":"2c87e16f-c313-489e-91ac-527b88f5fe22","resolution":{"observed_at":"2026-08-07T15:06:31.131537Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.112922Z","title":"The groq software-defined scale-out tensor streaming multiprocessor : From chips-to-systems architectural overview,","venue":null,"work_id":"75edca07-7922-4d78-8ad7-c9834e702712","year":2022},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.488902Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:2313f43be7857dcec5a281d7d6800604f596366578a7906cc2d2287a4b133f12","observation_id":"8a588eea-947f-4fc9-90b8-ecc9ac56877d","resolution":{"observed_at":"2026-08-07T15:06:31.117852Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.099203Z","title":"Application specific instruction processor based implementation of a gnss receiver on an fpga,","venue":null,"work_id":"d99a85f6-cbd3-409d-b44f-929b0e265e05","year":2006},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.492972Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:4aaaa73a0938e53d82361c6d2729f7214b5f2e563a6093b06168b7790cf00d5d","observation_id":"17a54891-ec19-44ba-bd33-97e9311fa29c","resolution":{"observed_at":"2026-08-07T15:06:31.103230Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.084003Z","title":"How flexible is your com- puting system?","venue":null,"work_id":"b051c7da-c699-42b9-a00e-44597915040f","year":2022},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.497613Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:29816cf4c3d376ef3cd27974bd0127953ac88cf3cbacc085a3db1675a5a46e28","observation_id":"1d680a15-8dd8-4544-a543-a003d2402db8","resolution":{"observed_at":"2026-08-07T15:06:31.088800Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.501738Z","title":"Tandem processor: Grappling with emerging operators in neural networks,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.501738Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:0ed0c88714ea7d4a9c8a66324d8f51ce65d39b469aa39c8d350420415830d9d9","observation_id":"a9dedbe3-43ab-407c-aeec-1b8d2a4311e5","resolution":{"observed_at":"2026-08-07T15:06:30.501738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.058978Z","title":"Mec: memory-efficient convolution for deep neural network,","venue":null,"work_id":"ed8ccdad-0d5a-4fb3-a1f8-6e5769c3b8a4","year":2017},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.505928Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:cc4524dcf7b7a2acbf8b50e3ea0e6c68bd28433c6e9bb83d8eb89a1fe5652b5b","observation_id":"cc7cdc0a-dd28-4d26-9de6-1a32daf336d9","resolution":{"observed_at":"2026-08-07T15:06:31.063778Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.043860Z","title":"A formalism of dnn accelerator flexibility,","venue":null,"work_id":"c5a060cd-dca6-4079-ab95-c6dfa5cfa9d0","year":2022},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.510290Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:1ffc0aa7e60bce38f9f06304e96171901e9c3639f2225dabc61ce679489d8e23","observation_id":"196c9181-8b97-4077-97e0-4ae7c239699b","resolution":{"observed_at":"2026-08-07T15:06:31.048331Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:30.515891Z","title":"Mlir: Scaling compiler infrastructure for domain specific computation,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.515891Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:b83137315ce572ff2bccd1c9e7f0f701eb8f33e620b434e1ea1435acde752aca","observation_id":"84c6dd27-560d-4be8-a390-6d1be9a0e729","resolution":{"observed_at":"2026-08-07T15:06:30.515891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:06:31.018753Z","title":"The hardware lottery,","venue":null,"work_id":"e1780fd1-bf85-4e88-8f91-d753ce5de127","year":2021},"citing_paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!","version":2},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:30.520235Z"},"links":{"citing_paper":"/paper/2505.16346"},"observation_digest":"sha256:6d166926b68b66c2bc5f0e2028bea503633286bf1f6a1ec24d8cec79e59968cd","observation_id":"d484f9f9-f15a-4eb2-988f-a8ac376c7f42","resolution":{"observed_at":"2026-08-07T15:06:31.023084Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.16346","last_updated":"2025-05-23T07:21:27Z","latest_version":2,"primary_category":"cs.AR","snapshot_observed_at":"2026-08-08T01:05:13.992626Z","submitted_at":"2025-05-22T07:59:26Z","title":"How to keep pushing ML accelerator performance? Know your rooflines!"},"reference_resolution":{"displayed":81,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":19,"verified_exact":2,"verified_fuzzy":60},"total_outbound_references":81},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 81 of 81 outbound references and 0 inbound Pith citation observations for arXiv:2505.16346."}