{"as_of":"2026-08-07T14:08:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:cf0af8fd2b52d73151f045967f141d24de0f0a7160958ec8444f8a85c8680b30","coverage":[{"denominator":98,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":98,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T23:35:06.211001Z","state":"measured"},{"denominator":98,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":98,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.17644/citation-record","integrity":"/paper/2506.17644/integrity","json":"/paper/2506.17644/citation-record.json","paper":"/paper/2506.17644"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:57.251099Z","title":"Cybercrime To Cost The World $10.5 Trillion Annually By","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:57.251099Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:e5291f009279f66d514eb1a4b43ad2e554f0ce0a50c91bc1b09beb3865da6d6c","observation_id":"9cb6e339-fc42-43e8-9cf1-a0c944bf33a1","resolution":{"observed_at":"2026-08-06T23:34:57.251099Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:57.504753Z","title":"AI Cyber Challenge Opens Registration, Adds $4 Million in Prizes, Shows Scoring Algorithm and Challenge Exemplar","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:57.504753Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:337d16a888dfde878f38024739cdf6b6e5cd72788054d2b84c87795a08e1daa2","observation_id":"ba61912d-7a22-4695-9516-9d12112eb63e","resolution":{"observed_at":"2026-08-06T23:34:57.504753Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:57.662253Z","title":"DEF CON®27 Hacking Conference Contests & Events","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:57.662253Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:016e6d66540703843ab35e2b5a3c3b8082e61d77cb2f421f9444b6a111fd2619","observation_id":"f12cf11d-1a26-4238-a8e1-d2de158ea976","resolution":{"observed_at":"2026-08-06T23:34:57.662253Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:57.854753Z","title":"0CTF 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:57.854753Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:8d8d4b7fb3da8821158852e31ae10ca55710266a1c50f23542aa0fc06ca06d5c","observation_id":"260f2c39-988b-4727-85b5-939ad76ef832","resolution":{"observed_at":"2026-08-06T23:34:57.854753Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.014833Z","title":"All about CTF","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.014833Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:bdb1d6d8cb89f2f3327046043080c793af9009be9c8ac42d7e3037eb6bbc5d5b","observation_id":"19d0120f-3d35-42f7-baf5-04eb405c69bc","resolution":{"observed_at":"2026-08-06T23:34:58.014833Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.064927Z","title":"Assistants API Overview","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.064927Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f72eb6540b200da812d291411146f986bbfb420a09644bcf6ca80fe74caabf49","observation_id":"eb70e278-d662-4f74-882f-2477daf33024","resolution":{"observed_at":"2026-08-06T23:34:58.064927Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.184386Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.184386Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:0421c8c6bf5ed332596c1ae9d0540c150c775749206a1857a02f12fc360d6141","observation_id":"3b083d52-2508-49a3-b0de-b6a6a6f3e9d3","resolution":{"observed_at":"2026-08-06T23:34:58.184386Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.301890Z","title":"Capture the Flag","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.301890Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:962693c4224c263afeacdb4f8714673cbfb6156a43bea3a7f9fecb29d2feed27","observation_id":"650dc15d-d2a8-41e6-af98-bbe8750a80c2","resolution":{"observed_at":"2026-08-06T23:34:58.301890Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.476834Z","title":"Capture the Flag for Empowered Cybersecurity Training","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.476834Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:a97d75b5d80053dd6b020210e4c62a84a1776409e5e0ab71370aaa21ab549657","observation_id":"6a65862f-cbc7-458c-9a5d-dbcf4066deb7","resolution":{"observed_at":"2026-08-06T23:34:58.476834Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.591413Z","title":"CGC: Cyber Grand Challenge","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.591413Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d6b448ced051f730736e961aa1ecd6dc79e818a80ea8538cb3e29cb5dbb23f74","observation_id":"e77e7ea6-cd15-434c-ab16-51c38dfd96f3","resolution":{"observed_at":"2026-08-06T23:34:58.591413Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.724750Z","title":"Claude 3.5 Sonnet","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.724750Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:5582b81818bbe78a0c58b34b53b4ccf5c3bd278946815a2027c29fa6f4b1ea63","observation_id":"a79cd1ce-01cf-4ae7-bfe8-bdecf9432220","resolution":{"observed_at":"2026-08-06T23:34:58.724750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.866687Z","title":"DeepSeek","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.866687Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:61ef24681beb186890dabcc631a3972fda6d0b82c0ddefb756cbaa44b523bef5","observation_id":"5273ab2a-785d-4aa5-b7fd-91bd215ac2af","resolution":{"observed_at":"2026-08-06T23:34:58.866687Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:59.000038Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.000038Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:1667509a5879dbc4a9b6f13850aeb905316f2bb43edb988213938778641347f8","observation_id":"8805a6da-5bd5-40b8-86fa-27a117f7be0f","resolution":{"observed_at":"2026-08-06T23:34:59.000038Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:59.114155Z","title":"Function calling","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.114155Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:588d419307e675e3546f6ca7ea5c608f639d671721f8a14e9d327651c6feb10a","observation_id":"5c07c62d-178a-4c67-ae97-44f8938b5f72","resolution":{"observed_at":"2026-08-06T23:34:59.114155Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:59.248142Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.248142Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:598c2b0494d15a13afe5cf711267394cd4ae671242a544045ad3718866467422","observation_id":"2e37072a-33d4-4543-92f1-777a2f2ff43e","resolution":{"observed_at":"2026-08-06T23:34:59.248142Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:59.347005Z","title":"Google CTF","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.347005Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:2a4c6cb4a4dbb63a5018f86fcee088276ee3fa22651c38a241210436daeae177","observation_id":"1c1b7b75-ee7f-4ec1-9a59-35bbc5a2d63b","resolution":{"observed_at":"2026-08-06T23:34:59.347005Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:59.451941Z","title":"gpt-3-5-turbo","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.451941Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:7f97b062fc2507e2a5c890cd7f70cfe744dd8dc0278e54469976ae5c32cf3949","observation_id":"79cc5202-a3f4-4b7b-be62-a79c829e95da","resolution":{"observed_at":"2026-08-06T23:34:59.451941Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.676492Z","title":null,"venue":null,"work_id":"da1f0611-f508-47fb-97a0-14353b598398","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.593056Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:b54b3e84eef48ca34fa1ccbb85f2170054fbcb19c7e3d12796cebaa65d2f569f","observation_id":"5bc6a0a4-3c28-4eed-9afd-a378298a3947","resolution":{"observed_at":"2026-08-06T23:35:09.697153Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.655777Z","title":null,"venue":null,"work_id":"d4eb6bb9-0dfb-46bd-8499-a78d9196b8b2","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.729431Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:de6301ad94c43bf11fde825c6e7373186f23f2458524cc63690f9eb31d898207","observation_id":"d5d2d54e-f4bf-4255-a601-a8b5a6798d2a","resolution":{"observed_at":"2026-08-06T23:35:09.659284Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.612871Z","title":null,"venue":null,"work_id":"44e0ca21-805e-4132-a426-e8041a2fbda1","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.875473Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f6910303ce9a330ea017a6db650f36344ae3285269b312fb0a33509dd1fcf61f","observation_id":"0a1f6a03-2465-4a06-82ba-b7270204573a","resolution":{"observed_at":"2026-08-06T23:35:09.645896Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.559733Z","title":null,"venue":null,"work_id":"5c45fd35-bb68-4a64-bacf-003dbea1dcb1","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.014973Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:80ff149dbb58b0d265dd792deb4b8203c9754209a06764da8f70cdc2e7afee82","observation_id":"b68f27cb-b5d1-42a4-a925-ec0c07d4b4de","resolution":{"observed_at":"2026-08-06T23:35:09.588364Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.542639Z","title":"Learning to Reason with LLMs | OpenAI","venue":null,"work_id":"8cf0950c-197b-4e0e-8b50-1704cc92ae2f","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.112630Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:ed5788965f786b133208fa2256292cae3daf55f88c1ab45ae509e7afbc96d812","observation_id":"1a3b9d3c-cb7b-48b1-a577-74f6f5aaa6c9","resolution":{"observed_at":"2026-08-06T23:35:09.546629Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.493438Z","title":"Meet Llama 3.1","venue":null,"work_id":"d9b38911-715b-404d-8655-8717d3c87640","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.226552Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:022471c1207c15a0b82a27461cdac5c67298888fd902009b968231878f84fb7a","observation_id":"2b86e511-6854-4305-bf91-9d42170b1dc6","resolution":{"observed_at":"2026-08-06T23:35:09.514750Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.430441Z","title":"Mixtral of experts | Mistral AI | Frontier AI in your hands","venue":null,"work_id":"fac336c1-648b-496f-9265-5c9313ed993c","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.273456Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:a4e11851f673c40b0f55f8f57c5645a8b23fbcc3f5264a111d08ebf47d70736f","observation_id":"464592a9-cb53-4b38-8bf2-39527696ec49","resolution":{"observed_at":"2026-08-06T23:35:09.446817Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.346944Z","title":"picoCTF - CMU Cybersecurity Competition","venue":null,"work_id":"727d7fa2-5a56-4325-ace9-0004a2d626eb","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.491515Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:221e561da211af322eefb69374ef842d44a170096a270b0ea11ad0c88f0ae487","observation_id":"19c2bea4-f1c4-43f8-8886-11d25172a832","resolution":{"observed_at":"2026-08-06T23:35:09.374754Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.328547Z","title":"picoCTF2024","venue":null,"work_id":"cf390a1f-3559-4abe-8a75-c61a2ba081ae","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.653650Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d6fa5df2e79b7f92b1bd66b716651640978127273b965c96eead6b4ff1a61aa3","observation_id":"8896b226-bfba-4220-b710-40b80568e4fa","resolution":{"observed_at":"2026-08-06T23:35:09.333649Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.288503Z","title":"Top 10 Cyber Hacking Competitions - Capture the Flag (CTF)","venue":null,"work_id":"63533bac-dd21-4c93-8029-122a8b2c914b","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.761007Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:34984c5cf7f757bd9a230858089155041d81cb778bf9e1f8a6d3ae9bc38b9286","observation_id":"37147e23-699d-4064-acb6-337f0f9dfada","resolution":{"observed_at":"2026-08-06T23:35:09.306715Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.260365Z","title":"UIUCTF 2024","venue":null,"work_id":"66b2028a-a6c8-497f-bc4f-ebc15553f390","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.886621Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:dd77f001a51bb9ae71892ddab7f88ff92c8c9dbe955246f7809a44b49e706c0b","observation_id":"7cf06032-3ce7-46fa-bb71-bc4b413962b1","resolution":{"observed_at":"2026-08-06T23:35:09.266291Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.242739Z","title":"VicOne & Block Harbor Spearhead Biggest Automotive Cap- ture the Flag Competition for Cybersecurity Enthusiasts World- wide","venue":null,"work_id":"c65f6413-c878-4ebe-a720-f7a932e8b344","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.033315Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:9607ea4a44464ff7ff86c86009089e3d211dbddad47da6cb732bd4318c91efa2","observation_id":"5bff1ced-39b7-49e3-ac64-2a1d862a56e1","resolution":{"observed_at":"2026-08-06T23:35:09.246250Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.217536Z","title":"Burp Suite - Application Security Testing Software","venue":null,"work_id":"ce584f39-f6db-4c0b-8f1f-73719f03ac89","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.118662Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:056d3c5130f3500b2fea1a3e8db31b37f3b8c7cc9c654c18b829cade6996efc3","observation_id":"31de45ff-2047-459e-bb53-ff29266a8b9a","resolution":{"observed_at":"2026-08-06T23:35:09.222551Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"kql-and-b/4390932","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:07.938631Z","title":"Microsoft Security Copilot Blog","venue":null,"work_id":"e035d54f-1729-40d2-8f2c-12d33508ad12","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.231335Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:66f880317593a8110aa2001b441407b5d6bc4363631cccbf49b02d67985f59dd","observation_id":"062a5e70-6c47-46c7-9bff-31b258fbe3c3","resolution":{"observed_at":"2026-08-06T23:35:07.980525Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.184445Z","title":"Proactive Defense: The Role of Offensive Security in Cybersecurity","venue":null,"work_id":"5a584eff-0baf-473e-b9c3-50911b3637bb","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.325726Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f41109ae4725d39a872054df772bf56b608180b937811b63c23c3eebb00a9162","observation_id":"a8e0dc8b-8b0b-45b5-bf40-a2c63f1d6465","resolution":{"observed_at":"2026-08-06T23:35:09.197784Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.149674Z","title":"Using AI for Offensive Security","venue":null,"work_id":"2df61746-3392-4ff8-ad98-f18964352f65","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.437033Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:3c4a46e67027e23d2495c37eea9601f92d6f357c28e00e431a549d0903878a13","observation_id":"6f9759e6-fe31-4284-95ce-e4e9de9dcdbc","resolution":{"observed_at":"2026-08-06T23:35:09.167578Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.139841Z","title":"What is Automated Vulnerability Remediation? https: //www.sentinelone.com/cybersecurity-101/cybersecurity/what-is-automated- vulnerability-remediation/","venue":null,"work_id":"1c666f32-c9ea-41bd-9160-f6bbf203d1f1","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.552239Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d6f0936d803db7faf35e20b65584863dd24dcba9e58f7eb7753cae1048b1c4db","observation_id":"b82a640e-60d4-45f1-abc8-35f5617842ca","resolution":{"observed_at":"2026-08-06T23:35:09.143171Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-06T23:35:01.653994Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.653994Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:20bd19092bd146757279ec89b4c736ed9879ccb113eb613ce27f5058235d2d0e","observation_id":"58b32c76-db2a-4c7c-bf97-38495c49041b","resolution":{"observed_at":"2026-08-06T23:35:01.653994Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.123478Z","title":null,"venue":null,"work_id":"a0a0fc87-6bf1-45f9-9b65-1390bc727cec","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.764451Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:3811090199150fa89c612c789c19d3262900f43d1c89e138994cbb69c9eb7644","observation_id":"f8604f7e-7da3-4be3-8670-960caaf3f1b2","resolution":{"observed_at":"2026-08-06T23:35:09.127689Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.13161","last_updated":"2024-04-19T20:11:12Z","snapshot_observed_at":"2026-08-07T13:01:36.699787Z","submitted_at":"2024-04-19T20:11:12Z","title":"CyberSecEval 2: A Wide-Ranging Cybersecurity Evaluation Suite for Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.13161","snapshot_observed_at":"2026-08-06T23:35:01.878295Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.878295Z"},"links":{"cited_paper":"/paper/2404.13161","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:89b1e8dd5391c0be7196d07252489f0746228d012726ed77ebb373ea4d30d3d4","observation_id":"0cbcfb12-1d95-4c42-a6f6-0988bfc48c0e","resolution":{"observed_at":"2026-08-06T23:35:01.878295Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.103479Z","title":null,"venue":null,"work_id":"87bb854f-1186-4341-9e83-c60c3028e5e5","year":2018},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.008375Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:3a8ede61acd7c5bca024bd2a48f260e30cb89290b413cce46fe8563ea382cfb9","observation_id":"4859e62d-cb3b-49a2-91b7-b3416159ae50","resolution":{"observed_at":"2026-08-06T23:35:09.106683Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.092331Z","title":null,"venue":null,"work_id":"4ba4a262-c7d5-4114-84d3-ee8cae9051b4","year":2017},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.138798Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:3224d532c8b2fcc3716506b78717830ca017f5588b9d5e541a3b8efedbebe6fa","observation_id":"2b9a2b1c-2705-48bf-9dec-bef8021c1d37","resolution":{"observed_at":"2026-08-06T23:35:09.096009Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.06573","last_updated":"2024-07-09T06:07:45Z","snapshot_observed_at":"2026-08-01T19:24:09.434113Z","submitted_at":"2024-07-09T06:07:45Z","title":"LLM for Mobile: An Initial Roadmap","version":1},"cited_work":{"arxiv_id":"2407.06573","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.06573","snapshot_observed_at":"2026-08-06T23:35:07.646513Z","title":"LLM for Mobile: An Initial Roadmap","venue":"cs.SE","work_id":"ba8a34ae-e6a1-4a24-97cf-c6fd93ba7aed","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.325227Z"},"links":{"cited_paper":"/paper/2407.06573","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:6fde15449a815648d8d77ceebe9293eb74f8bd6ba1e748ec2754a02945d54cae","observation_id":"30f5794f-077c-488f-a7f0-5d7da8de41aa","resolution":{"observed_at":"2026-08-06T23:35:07.675542Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.19470","last_updated":"2025-09-23T03:45:42Z","snapshot_observed_at":"2026-07-06T20:58:19.305457Z","submitted_at":"2025-03-25T09:00:58Z","title":"ReSearch: Learning to Reason with Search for LLMs via Reinforcement Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.19470","snapshot_observed_at":"2026-08-06T23:35:02.407420Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.407420Z"},"links":{"cited_paper":"/paper/2503.19470","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:2759ecbd41cef4b80408013b996a22191835bb92ffdf3ce772e6490be3674441","observation_id":"243b50d8-72a2-4196-b8dc-409c83a4a933","resolution":{"observed_at":"2026-08-06T23:35:02.407420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.072403Z","title":null,"venue":null,"work_id":"5a472660-539a-4b31-9398-02557a3abed7","year":2014},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.488945Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:044dcbb61c6b9861389f4bd490678829813eb032b64ae5db3d5de2885c7ae418","observation_id":"24fd8075-1986-4e33-a9e0-0f41bff2310d","resolution":{"observed_at":"2026-08-06T23:35:09.084729Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.061883Z","title":null,"venue":null,"work_id":"f435c73f-fdd2-4418-a243-d07b8282fee6","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.559690Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:3703aac63eb5330651013a34592f57c8439a55dcfbbf1583ecb29681042dcb11","observation_id":"bf25768b-f992-4bdd-bf90-a5a351051373","resolution":{"observed_at":"2026-08-06T23:35:09.066019Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.039717Z","title":null,"venue":null,"work_id":"2b9d6a1a-9b1b-4293-bcd0-a7075442872e","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.650297Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f47904fbcd4149b1dbb68dc82399094cd955d78e437452cbf800581a4115ec80","observation_id":"fb07652c-ed80-4ed3-962e-7e00ff3ae398","resolution":{"observed_at":"2026-08-06T23:35:09.047651Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:02.788629Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.788629Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:b23b0c95215a58fe14624585dbac968333c76331d234caffe31d354a2fe639d2","observation_id":"19ccf574-abf8-469d-9c9b-810527772f15","resolution":{"observed_at":"2026-08-06T23:35:02.788629Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.984755Z","title":null,"venue":null,"work_id":"6baaebdb-5f85-4a20-ad3f-760fb78cc723","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.909045Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:86734d0b5b01dc9637c7afea959da67f1bcd64896665b01592df86735258c692","observation_id":"34fb4ab3-b49e-404c-a45a-cff2061fd3cb","resolution":{"observed_at":"2026-08-06T23:35:08.989382Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01637","last_updated":"2025-03-30T00:26:48Z","snapshot_observed_at":"2026-08-04T01:48:56.038521Z","submitted_at":"2024-06-02T16:25:26Z","title":"Teams of LLM Agents can Exploit Zero-Day Vulnerabilities","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01637","snapshot_observed_at":"2026-08-06T23:35:02.972447Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.972447Z"},"links":{"cited_paper":"/paper/2406.01637","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:df868a99be4502ec35e79993616370ab38a5e42d0fc663475acde95bf19ca574","observation_id":"52263a2f-b4c7-4518-acc0-075488ea41d0","resolution":{"observed_at":"2026-08-06T23:35:02.972447Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.949958Z","title":null,"venue":null,"work_id":"b869abfd-431b-4826-8c30-d23d09639b01","year":2007},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.105682Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:3cfd8b0f9df7dd401ae35870c947812730d464e1dfbd09c8264eb2c17b98c0fc","observation_id":"c854178b-4f06-428b-90eb-9b4760c2e97c","resolution":{"observed_at":"2026-08-06T23:35:08.953936Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.10997","last_updated":"2024-03-27T09:16:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-18T07:47:33Z","title":"Retrieval-Augmented Generation for Large Language Models: A Survey","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.10997","snapshot_observed_at":"2026-08-06T23:35:03.279545Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.279545Z"},"links":{"cited_paper":"/paper/2312.10997","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c660c289bb50074c5dfc2501f5003def9a7ee2fdc44c9d7a892939bebd5d4f21","observation_id":"e4b21f42-151b-42ca-8c60-d815a75b1025","resolution":{"observed_at":"2026-08-06T23:35:03.279545Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12420","last_updated":"2023-12-22T14:07:16Z","snapshot_observed_at":"2026-07-06T16:50:26.807430Z","submitted_at":"2023-11-21T08:20:39Z","title":"How Far Have We Gone in Vulnerability Detection Using Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.12420","snapshot_observed_at":"2026-08-06T23:35:03.441508Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.441508Z"},"links":{"cited_paper":"/paper/2311.12420","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d7a2cbf997fb347066cb094c3e7525840b958744f25dafa237bcc5aebe9a68e3","observation_id":"8e153b3b-cf27-45dd-9f3f-4cf8fdfec779","resolution":{"observed_at":"2026-08-06T23:35:03.441508Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.910692Z","title":null,"venue":null,"work_id":"00092062-f5b4-4583-9d15-3f640165c96f","year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.634822Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f82c4c4dcdff34a516b26d3b850e42e16f2819ae617c132d1c38f2d72a1ba996","observation_id":"0786019f-5f60-4474-8a98-0f6d8a838c17","resolution":{"observed_at":"2026-08-06T23:35:08.934745Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2009.03300","last_updated":"2021-01-12T18:57:11Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-09-07T17:59:25Z","title":"Measuring Massive Multitask Language Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.03300","snapshot_observed_at":"2026-08-06T23:35:03.753687Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.753687Z"},"links":{"cited_paper":"/paper/2009.03300","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:9b0c2047ccb2c28e57773ba862e872745cf634462309b4e00c8aaaa3a25e1ba8","observation_id":"b6265da3-c39e-4b0c-97d4-da05c695a7b8","resolution":{"observed_at":"2026-08-06T23:35:03.753687Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17788","last_updated":"2024-07-25T05:42:14Z","snapshot_observed_at":"2026-08-05T00:09:41.708678Z","submitted_at":"2024-07-25T05:42:14Z","title":"PenHeal: A Two-Stage LLM Framework for Automated Pentesting and Optimal Remediation","version":1},"cited_work":{"arxiv_id":"2407.17788","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.17788","snapshot_observed_at":"2026-08-06T23:35:07.025550Z","title":"PenHeal: A Two-Stage LLM Framework for Automated Pentesting and Optimal Remediation","venue":"cs.CR","work_id":"9b52dbe8-e248-4ffc-80de-473b1400c14d","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.833812Z"},"links":{"cited_paper":"/paper/2407.17788","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:8a64c70c7d0cb3803300809b2bf811de3b16696e9e15ade16cdcda13c3950a5b","observation_id":"d326bf29-32f0-4432-a82b-45699265daed","resolution":{"observed_at":"2026-08-06T23:35:07.054889Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.887049Z","title":null,"venue":null,"work_id":"49a665fd-f105-4e8e-90fd-789f0d4fea84","year":2017},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.017101Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:80481c297de3a9c555debf497eccf0e83f2a3fbce0378a61ec29afdf303ce0ad","observation_id":"ef47d230-57d8-47be-b044-1cb5a37ceb36","resolution":{"observed_at":"2026-08-06T23:35:08.890346Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.20787","last_updated":"2025-01-06T07:22:50Z","snapshot_observed_at":"2026-08-05T03:53:59.997985Z","submitted_at":"2024-12-30T08:11:54Z","title":"SecBench: A Comprehensive Multi-Dimensional Benchmarking Dataset for LLMs in Cybersecurity","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.20787","snapshot_observed_at":"2026-08-06T23:35:04.049872Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.049872Z"},"links":{"cited_paper":"/paper/2412.20787","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c42a82db71ca57c848c018735174671f5e0d40dae4800e6af98cbfd4868ab20d","observation_id":"086b98e7-cc7d-49df-9a5d-1232a30abb70","resolution":{"observed_at":"2026-08-06T23:35:04.049872Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.16169","last_updated":"2024-10-23T07:32:15Z","snapshot_observed_at":"2026-07-06T16:53:22.999597Z","submitted_at":"2023-11-16T13:17:20Z","title":"Understanding the Effectiveness of Large Language Models in Detecting Security Vulnerabilities","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.16169","snapshot_observed_at":"2026-08-06T23:35:04.178937Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.178937Z"},"links":{"cited_paper":"/paper/2311.16169","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:a76393121463478c5491f0617f307fef954fa7f12db4d666497ac85437a90d4f","observation_id":"89402b97-58a7-49fc-a238-510ff92fd30d","resolution":{"observed_at":"2026-08-06T23:35:04.178937Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.856949Z","title":null,"venue":null,"work_id":"1596f5bc-e1b4-45f4-83fc-53434346bdef","year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.282967Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:a9bebe2841116b7bbb2c293649b639e0b96bb82ddf40aeb9f024ce9b34d41ca9","observation_id":"d131b1f6-b22d-423e-a734-d38acffb53b0","resolution":{"observed_at":"2026-08-06T23:35:08.865416Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.830065Z","title":null,"venue":null,"work_id":"5fd4a2d0-6520-4bc5-ac10-c69a5a722c1d","year":2020},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.420784Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:0a8a5fe0796c0da2d79b962e2f7422d461f2b0e47ce74e3e6a2b1734f8f07226","observation_id":"1dc2312a-cb0a-4bf8-85a4-7f859317c0c1","resolution":{"observed_at":"2026-08-06T23:35:08.840467Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23383","last_updated":"2025-03-30T10:16:25Z","snapshot_observed_at":"2026-08-07T06:56:33.119385Z","submitted_at":"2025-03-30T10:16:25Z","title":"ToRL: Scaling Tool-Integrated RL","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.23383","snapshot_observed_at":"2026-08-06T23:35:04.518537Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.518537Z"},"links":{"cited_paper":"/paper/2503.23383","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:62224cbe1bad533d4a9fd17c51322e3c6907716837e5b730df3c1df455579d5b","observation_id":"7c9e36bc-fb98-4b7f-90c5-0b5834acc2fa","resolution":{"observed_at":"2026-08-06T23:35:04.518537Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.05459","last_updated":"2024-05-08T06:16:23Z","snapshot_observed_at":"2026-08-02T13:57:57.119489Z","submitted_at":"2024-01-10T09:25:45Z","title":"Personal LLM Agents: Insights and Survey about the Capability, Efficiency and Security","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.05459","snapshot_observed_at":"2026-08-06T23:35:04.607502Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.607502Z"},"links":{"cited_paper":"/paper/2401.05459","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:381689b2d8a3dc1f179cab97770850bab5d0c0b699bd7d53b9365e1c059a7950","observation_id":"099806aa-a002-40b8-92bd-97eaf28c1dd4","resolution":{"observed_at":"2026-08-06T23:35:04.607502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.774746Z","title":null,"venue":null,"work_id":"aff8d476-c663-40fd-946b-6ac981806faa","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.734768Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:e28d9e5bd71684611eab9227f6261791ef564054b46614090b0f62fcc921bac6","observation_id":"8e124330-059f-4511-b28e-b21ab49eebd4","resolution":{"observed_at":"2026-08-06T23:35:08.794746Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.17419","last_updated":"2025-06-25T02:24:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-24T18:50:52Z","title":"From System 1 to System 2: A Survey of Reasoning Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.17419","snapshot_observed_at":"2026-08-06T23:35:04.775152Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.775152Z"},"links":{"cited_paper":"/paper/2502.17419","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:daba77b68cf41d0d0ff50ee5dc0834f73389d50d2fba794c85148df6412c4d8d","observation_id":"18e460b7-36cd-41a8-9b2d-765927d27ab7","resolution":{"observed_at":"2026-08-06T23:35:04.775152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.01990","last_updated":"2025-08-02T12:44:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-31T18:00:29Z","title":"Advances and Challenges in Foundation Agents: From Brain-Inspired Intelligence to Evolutionary, Collaborative, and Safe Systems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.01990","snapshot_observed_at":"2026-08-06T23:35:04.862557Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.862557Z"},"links":{"cited_paper":"/paper/2504.01990","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:8ed109af93f086f6c453ce095aed941ec092d18e9e8d739b6b91fff24113ea8e","observation_id":"26d4119a-f7a0-4190-98f9-31a588c7a61f","resolution":{"observed_at":"2026-08-06T23:35:04.862557Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:04.956694Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.956694Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d35473d0018f9fa3f667b43bbd3788c3105263c7fa64994145927bf9fd143ead","observation_id":"c57679e1-33b8-465b-97d4-1f70a032c624","resolution":{"observed_at":"2026-08-06T23:35:04.956694Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:05.076984Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.076984Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:0d3f72163f7a582997850dc3c42d2b1cf5ede3801bf07a5a3b60e0a196086049","observation_id":"134025a9-b61d-4fdf-9a4f-bbfc9c00a3b7","resolution":{"observed_at":"2026-08-06T23:35:05.076984Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.596990Z","title":null,"venue":null,"work_id":"168e7a47-2a29-414f-90a7-3d9bc187f071","year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.236867Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:da73930c2a0224da1b9c371b71416d99035389a3e421126897c32eb2d8486886","observation_id":"5832686b-7ff9-4d8e-8815-ad7d9eba0abe","resolution":{"observed_at":"2026-08-06T23:35:08.607318Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.627892Z","title":null,"venue":null,"work_id":"41572718-e355-4dc3-8747-833a420baa14","year":null},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.104820Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:589181e4a47bb54c97eeab332933fd3cdb4a3e68f9acebc723f13d4cf423c569","observation_id":"061e2c6b-20ed-4e55-8616-a217004c1bc6","resolution":{"observed_at":"2026-08-06T23:35:08.634278Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.557612Z","title":null,"venue":null,"work_id":"0bbc35e4-d496-4278-8ab2-e487d2ace1fb","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.435569Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:0b407348eae35e182db20536484941dc395703f77dd933f70498e19ff3b18362","observation_id":"28b389a7-918c-48a5-b7d5-85a811bdf43e","resolution":{"observed_at":"2026-08-06T23:35:08.564431Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.583120Z","title":null,"venue":null,"work_id":"d08a32d8-843c-4273-85c7-a3a6069872b5","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.320121Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:5805ff174807ef760c0caa52224cfa537cdc7da89a45a5107d6a7917ec7c16b1","observation_id":"595f317a-4f9d-4b7b-bf9d-6156dc4092d4","resolution":{"observed_at":"2026-08-06T23:35:08.586523Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11814","last_updated":"2024-02-19T04:08:44Z","snapshot_observed_at":"2026-07-06T17:31:58.268105Z","submitted_at":"2024-02-19T04:08:44Z","title":"An Empirical Evaluation of LLMs for Solving Offensive Security Challenges","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.11814","snapshot_observed_at":"2026-08-06T23:35:05.704955Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.704955Z"},"links":{"cited_paper":"/paper/2402.11814","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:794773494f56ef74953b8ad0318583ee4b974e150ca448ed57629527447de0a1","observation_id":"531bb9b8-a3ff-4145-b1e8-b789c3bd6c80","resolution":{"observed_at":"2026-08-06T23:35:05.704955Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:05.584511Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.584511Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:367b58fb3ceaa402a15980113e645916cb8629c466dfc20460d98b90ead327ad","observation_id":"c9f8a3be-43ea-4c3d-811f-12bf00b57524","resolution":{"observed_at":"2026-08-06T23:35:05.584511Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.516212Z","title":null,"venue":null,"work_id":"cdb64cab-0f6b-415d-bfd2-a7891101e89a","year":2021},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.896211Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:337ed7176af7b89ee436cc535e6eef22baf27a284145b3cd4b55230a293b275b","observation_id":"d30d15d8-7bc0-4777-9a8f-bd27cbc5d928","resolution":{"observed_at":"2026-08-06T23:35:08.524858Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05590","last_updated":"2025-02-18T12:26:33Z","snapshot_observed_at":"2026-07-06T18:27:36.903877Z","submitted_at":"2024-06-08T22:21:42Z","title":"NYU CTF Bench: A Scalable Open-Source Benchmark Dataset for Evaluating LLMs in Offensive Security","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05590","snapshot_observed_at":"2026-08-06T23:35:05.819135Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.819135Z"},"links":{"cited_paper":"/paper/2406.05590","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:073df5f4fe9a18914682e5fb41d09ea94aaa87e2a01417ef4448fc2a9f241ba8","observation_id":"425b3577-9f55-4a6f-9635-ccaf8025aa32","resolution":{"observed_at":"2026-08-06T23:35:05.819135Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.16185","last_updated":"2025-06-07T17:03:11Z","snapshot_observed_at":"2026-08-04T22:35:18.426928Z","submitted_at":"2024-01-29T14:32:27Z","title":"LLM4Vuln: A Unified Evaluation Framework for Decoupling and Enhancing LLMs' Vulnerability Reasoning","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.16185","snapshot_observed_at":"2026-08-06T23:35:05.906614Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.906614Z"},"links":{"cited_paper":"/paper/2401.16185","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:86b621f7489f75a2fe2e58e1bf383f61b215a34309df07a0debe8d2ac3e43d36","observation_id":"f0d47db9-8e72-46b0-b5ac-3c7addd97718","resolution":{"observed_at":"2026-08-06T23:35:05.906614Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12566","last_updated":"2021-07-27T03:01:47Z","snapshot_observed_at":"2026-07-06T11:32:47.487435Z","submitted_at":"2021-07-27T03:01:47Z","title":"Thunder CTF: Learning Cloud Security on a Dime","version":1},"cited_work":{"arxiv_id":"2107.12566","doi":null,"metadata_source":"pith","pith_arxiv_id":"2107.12566","snapshot_observed_at":"2026-08-06T23:35:06.652953Z","title":"Thunder CTF: Learning Cloud Security on a Dime","venue":"cs.CR","work_id":"04142027-bf19-44d4-a589-c1a5e5ed429b","year":2021},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.900130Z"},"links":{"cited_paper":"/paper/2107.12566","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f9a0a9bfb244dbbb8a2489b00c666bdf0b18c9cd900f16e13ee29e6c8f6c14b1","observation_id":"3ed92649-b2e0-40ac-b220-6790d434d41c","resolution":{"observed_at":"2026-08-06T23:35:06.658565Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.10443","last_updated":"2023-08-21T03:30:21Z","snapshot_observed_at":"2026-08-05T16:13:40.630996Z","submitted_at":"2023-08-21T03:30:21Z","title":"Using Large Language Models for Cybersecurity Capture-The-Flag Challenges and Certification Questions","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.10443","snapshot_observed_at":"2026-08-06T23:35:05.945362Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.945362Z"},"links":{"cited_paper":"/paper/2308.10443","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:2b7ca425ce0a8a188e034ae424af9d6f371828bb7faa05ffb362c96b0ec7d877","observation_id":"53081299-40d5-450f-ba3a-f98341845241","resolution":{"observed_at":"2026-08-06T23:35:05.945362Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.447158Z","title":null,"venue":null,"work_id":"da95c98e-ad08-42cd-849f-ea6e5c289fe7","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.926998Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f139437c9e0d41e5674a59c63d837949442ea2e77ea12e0409cdf07557f9313e","observation_id":"061eca14-bf72-4530-9a4a-3b9e8c672612","resolution":{"observed_at":"2026-08-06T23:35:08.485010Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.376344Z","title":null,"venue":null,"work_id":"ddc553d6-4733-44c6-b039-1b953ad304fb","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.955326Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:701f660f7d7e8de4e38cd4b8f0fa1d5ec87d39edbb0f1c40c2ad9147c04b31ce","observation_id":"e7e2d59b-4718-4bb0-b950-824fad8fc02c","resolution":{"observed_at":"2026-08-06T23:35:08.386518Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.409853Z","title":null,"venue":null,"work_id":"8bdb3ecc-5966-4364-b46c-fac2da5eea6c","year":2022},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.949603Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:18732a2004f7df07f312e90b04fd192391f80357bc76630a6f5bc81efacaf066","observation_id":"ecb831ee-ce95-4cf9-8b6a-aea9b41a7e58","resolution":{"observed_at":"2026-08-06T23:35:08.415128Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.12575","last_updated":"2024-07-24T07:49:14Z","snapshot_observed_at":"2026-07-06T17:05:35.500263Z","submitted_at":"2023-12-19T20:19:43Z","title":"LLMs Cannot Reliably Identify and Reason About Security Vulnerabilities (Yet?): A Comprehensive Evaluation, Framework, and Benchmarks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.12575","snapshot_observed_at":"2026-08-06T23:35:05.983497Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.983497Z"},"links":{"cited_paper":"/paper/2312.12575","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:0983eff8177a42945af5389114bcf0e773ea9242b450bfe86771395757aa31a8","observation_id":"f450ee49-39b8-4ee6-82a8-6824a1947f21","resolution":{"observed_at":"2026-08-06T23:35:05.983497Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.353257Z","title":null,"venue":null,"work_id":"31bbef40-0815-4926-9a68-96000e41fe40","year":null},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.971759Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:4d03650a7cdc006d178d64c61d89913c24accae742ab9fa85c5a8f042b66f9e6","observation_id":"03dbb136-9d8e-43a4-8088-bfcf897e4419","resolution":{"observed_at":"2026-08-06T23:35:08.361894Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:05.991880Z","title":null,"venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.991880Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:27555fd8633763c28f489158eb0910119c3c88e826c17abbc469ed00ed41e473","observation_id":"5a9b9dbd-8469-481d-b93b-c746889611af","resolution":{"observed_at":"2026-08-06T23:35:05.991880Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.239120Z","title":null,"venue":null,"work_id":"1248cb14-c093-480b-99c0-455480968d02","year":2020},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.005780Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:a8d96e2df7aa9b9814b982f3fd285f6b500304020086cec03793787cdd72b8d7","observation_id":"9bae381d-cd8a-47af-887b-cefb8338e859","resolution":{"observed_at":"2026-08-06T23:35:08.245162Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.278904Z","title":"Department of Health and Human Services","venue":null,"work_id":"5b230e85-cc4e-40ec-b0ec-3608946895fe","year":2018},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.988781Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:cd7973928542aec43353129c22018180b9e09669254a58bba15f776b7e9c3e9c","observation_id":"0c0021e4-5a44-40ac-91e8-a2a920c61026","resolution":{"observed_at":"2026-08-06T23:35:08.288849Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.224753Z","title":null,"venue":null,"work_id":"61b0bbc4-ecd0-4701-b15c-af0f9336cbd8","year":2018},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.057837Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:6000cdcce605fa70b03c523d9f508f0e264d6296dedfac6d7cdba02d3e83afcf","observation_id":"600f77bb-75a4-4f1d-ac20-77f11af97fdd","resolution":{"observed_at":"2026-08-06T23:35:08.230058Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.174516Z","title":null,"venue":null,"work_id":"02fda5b3-d423-4d1e-a4d6-b32f683d484f","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.105715Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:4c7eac467ed08563f3bac061fb855393567779cc293d2d1a16af20b73b48013a","observation_id":"ad7539cb-ae77-48be-ae1d-009ccee3cd18","resolution":{"observed_at":"2026-08-06T23:35:08.191001Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.13945","last_updated":"2024-12-21T15:21:27Z","snapshot_observed_at":"2026-07-06T18:03:33.406287Z","submitted_at":"2024-04-22T07:41:41Z","title":"How Multi-Modal LLMs Reshape Visual Deep Learning Testing? A Comprehensive Study Through the Lens of Image Mutation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.13945","snapshot_observed_at":"2026-08-06T23:35:06.031857Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.031857Z"},"links":{"cited_paper":"/paper/2404.13945","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c00e9c2a247f0dc3f6a5c7536734b9bf88f5506f2b879a1da0a5d5fbb0089dc9","observation_id":"882b0845-7082-451a-a65d-acbe4c3413df","resolution":{"observed_at":"2026-08-06T23:35:06.031857Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.052045Z","title":null,"venue":null,"work_id":"403c3c34-d3fd-44f1-a673-78bed36ae2d9","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.174475Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:6662554675c5c85c3beae213899f3881f0bce8c0134d31ff5c30376b641da7bc","observation_id":"454d7eff-83d5-4f20-83fb-b4e15faf8d39","resolution":{"observed_at":"2026-08-06T23:35:08.064873Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:06.178897Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.178897Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:33d5fcef45f03a374f46a63789a88356f74bad4239b74795279d8da7909b4aed","observation_id":"038fb3f8-62e9-42ac-9fe8-9d1b5928b1bf","resolution":{"observed_at":"2026-08-06T23:35:06.178897Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.109218Z","title":null,"venue":null,"work_id":"84a5ba28-24f3-4aca-937d-36d7b4f9a616","year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.132701Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:59753df41a104984d471ff932b1dde2264194009ef132c4e21aec6819e39ca0e","observation_id":"bb0a00d6-6ee1-47c9-a1ff-e900aef2e2d9","resolution":{"observed_at":"2026-08-06T23:35:08.114044Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.03629","last_updated":"2023-03-10T01:00:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-10-06T01:00:32Z","title":"ReAct: Synergizing Reasoning and Acting in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.03629","snapshot_observed_at":"2026-08-06T23:35:06.199616Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.199616Z"},"links":{"cited_paper":"/paper/2210.03629","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c12561a253bebb57a09c6d3752772557f27a455952a5bb389c292c224519bb8c","observation_id":"1d8fa8c4-2b69-43ef-84e7-60eb88c035c4","resolution":{"observed_at":"2026-08-06T23:35:06.199616Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.09210","last_updated":"2024-10-03T04:35:39Z","snapshot_observed_at":"2026-07-06T16:48:05.271182Z","submitted_at":"2023-11-15T18:54:53Z","title":"Chain-of-Note: Enhancing Robustness in Retrieval-Augmented Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.09210","snapshot_observed_at":"2026-08-06T23:35:06.202672Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.202672Z"},"links":{"cited_paper":"/paper/2311.09210","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f69714cb62a5489839d3d5ed89591aadfa6127eb6bebf66f8ce358923e701d67","observation_id":"fe207e1e-b66f-417f-9ada-cbb87495f96f","resolution":{"observed_at":"2026-08-06T23:35:06.202672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.06838","last_updated":"2025-07-21T05:24:59Z","snapshot_observed_at":"2026-08-04T23:33:44.230238Z","submitted_at":"2024-03-11T15:59:59Z","title":"ACFIX: Guiding LLMs with Mined Common RBAC Practices for Context-Aware Repair of Access Control Vulnerabilities in Smart Contracts","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.06838","snapshot_observed_at":"2026-08-06T23:35:06.207079Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.207079Z"},"links":{"cited_paper":"/paper/2403.06838","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:6e62c08ced49926c1c49cc0b2753e7b15ffe6f613f89f470d3fef95c703f3ed0","observation_id":"6927f7f0-3050-432d-aeb3-f45efa2e8c02","resolution":{"observed_at":"2026-08-06T23:35:06.207079Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:07.991614Z","title":null,"venue":null,"work_id":"3a00464f-b58a-4237-af7e-77b47055141f","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":94,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.196569Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:b0ac346a6e89aea9591dc59ff115ee1d03e1aa4236341540b76c4bcc5df1fff5","observation_id":"056695c2-c176-4ce6-9295-9da74768aaf7","resolution":{"observed_at":"2026-08-06T23:35:07.996021Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.16906","last_updated":"2024-06-06T06:01:41Z","snapshot_observed_at":"2026-08-03T14:14:25.841503Z","submitted_at":"2024-02-25T00:56:27Z","title":"Debug like a Human: A Large Language Model Debugger via Verifying Runtime Execution Step-by-step","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.16906","snapshot_observed_at":"2026-08-06T23:35:06.211001Z","title":"The knowledge fully matches the write- up and accurately reflects its content","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.211001Z"},"links":{"cited_paper":"/paper/2402.16906","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:0844538c89ec50be109530ce8f99557504dc4a0f5dc264dfa08b236bc1f06490","observation_id":"234bcb10-d3a6-47f2-bc4d-fe64c78c9cb2","resolution":{"observed_at":"2026-08-06T23:35:06.211001Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.014852Z","title":"InMulti-Agent Security Workshop@ NeurIPS’23","venue":null,"work_id":"eda81d94-b4fc-40ed-9193-b07069752afa","year":null},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.182493Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:41c177b038bdad78f39ef2a8ef10dfd5f0971bf591c121b5f13a0f982d9a94c9","observation_id":"89e2fd64-527f-4049-a1b2-eaf4aa002179","resolution":{"observed_at":"2026-08-06T23:35:08.019250Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.07688","last_updated":"2024-06-03T08:14:45Z","snapshot_observed_at":"2026-07-06T17:28:56.092370Z","submitted_at":"2024-02-12T14:53:28Z","title":"CyberMetric: A Benchmark Dataset based on Retrieval-Augmented Generation for Evaluating LLMs in Cybersecurity Knowledge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.07688","snapshot_observed_at":"2026-08-06T23:35:05.975106Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.975106Z"},"links":{"cited_paper":"/paper/2402.07688","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:47d01e7d7d4fe6047f244ea64353f22d49cdce4bfe77b6c5c08d95fb4da21b24","observation_id":"7a79103d-b8de-4a37-8e5e-96ceec3b28e6","resolution":{"observed_at":"2026-08-06T23:35:05.975106Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:57.394753Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:57.394753Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:797b7d209d5a0111f095122659810c1030b303a0fc98b6765bb982025283f22d","observation_id":"1b5fe0dc-2320-4683-80a4-40beeb135602","resolution":{"observed_at":"2026-08-06T23:34:57.394753Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-07T12:06:20.410800Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges"},"reference_resolution":{"displayed":98,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":79,"verified_exact":4,"verified_fuzzy":14},"total_outbound_references":98},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 98 of 98 outbound references and 0 inbound Pith citation observations for arXiv:2506.17644."}