{"as_of":"2026-08-08T17:55:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:bb97b31e7b7a5ccc05f73d62d1057e461bb0d5447ed36505f2eb1d0c7206cd85","coverage":[{"denominator":53,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":53,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:35:41.214905Z","state":"measured"},{"denominator":56,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":56,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-26T16:24:25.357338Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"cited_work":{"arxiv_id":"2506.07594","doi":"10.48550/arxiv.2506.07594","metadata_source":"arxiv_reference","pith_arxiv_id":"2506.07594","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Evaluating LLMs Ef- fectiveness in Detecting and Correcting Test Smells","venue":"ArXiv.org","work_id":"ab5a0407-c572-41af-9eee-fb647d01276d","year":2025},"citing_paper":{"arxiv_id":"2604.23361","last_updated":"2026-04-25T16:05:30Z","snapshot_observed_at":"2026-07-06T23:09:34.300163Z","submitted_at":"2026-04-25T16:05:30Z","title":"An Empirical Evaluation of Locally Deployed LLMs for Bug Detection in Python Code","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-08T07:59:25.198736Z"},"links":{"cited_paper":"/paper/2506.07594","citing_paper":"/paper/2604.23361"},"observation_digest":"sha256:d65e43e01b43ac4d9a49f9da806af61ac170677c8ade114b3f43719aacfc3ab0","observation_id":"a8bbc461-08d7-43e7-a15d-eac1263a4d91","resolution":{"observed_at":"2026-05-11T20:46:14.700137Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"cited_work":{"arxiv_id":"2506.07594","doi":"10.48550/arxiv.2506.07594","metadata_source":"arxiv_reference","pith_arxiv_id":"2506.07594","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Evaluating LLMs Ef- fectiveness in Detecting and Correcting Test Smells","venue":"ArXiv.org","work_id":"ab5a0407-c572-41af-9eee-fb647d01276d","year":2025},"citing_paper":{"arxiv_id":"2605.02091","last_updated":"2026-05-03T23:21:13Z","snapshot_observed_at":"2026-07-06T23:15:11.601042Z","submitted_at":"2026-05-03T23:21:13Z","title":"How Compliant Are GitHub Actions Workflows? A Checklist-Based Study with LLM-Assisted Auditing","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-08T19:19:18.667967Z"},"links":{"cited_paper":"/paper/2506.07594","citing_paper":"/paper/2605.02091"},"observation_digest":"sha256:da74ac7ac5eebbf42822a358355b2c6cdd6905620b30747cff1ba81ea549efcf","observation_id":"8902ec67-7c8b-4674-894b-3fbf0d010e3c","resolution":{"observed_at":"2026-05-09T05:55:31.517733Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"cited_work":{"arxiv_id":"2506.07594","doi":"10.48550/arxiv.2506.07594","metadata_source":"arxiv_reference","pith_arxiv_id":"2506.07594","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Evaluating LLMs Ef- fectiveness in Detecting and Correcting Test Smells","venue":"ArXiv.org","work_id":"ab5a0407-c572-41af-9eee-fb647d01276d","year":2025},"citing_paper":{"arxiv_id":"2606.20173","last_updated":"2026-06-18T12:40:43Z","snapshot_observed_at":"2026-07-06T23:55:20.134506Z","submitted_at":"2026-06-18T12:40:43Z","title":"Qiskit Code Migration with LLMs","version":1},"reference_index":155,"source":"arxiv_source","source_observed_at":"2026-06-26T16:24:25.357338Z"},"links":{"cited_paper":"/paper/2506.07594","citing_paper":"/paper/2606.20173"},"observation_digest":"sha256:bb00de6f5d8c404143cf42292f4d291fabd3b2975d8a8a2304e43c981232b11f","observation_id":"1fdc2dcf-d38d-430e-9dd7-f7c350880795","resolution":{"observed_at":"2026-06-26T16:29:35.581303Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.07594/citation-record","integrity":"/paper/2506.07594/integrity","json":"/paper/2506.07594/citation-record.json","paper":"/paper/2506.07594"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.751923Z","title":"On the relation of test smells to software code quality,","venue":null,"work_id":"acbf8be4-3b65-4e0e-ad41-63910e2d6e3e","year":2018},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.943703Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:473fe39e8fdbe35594be1a851f12feb98589ff025dd395c08ba65b5f90db1939","observation_id":"23247134-ec7a-462b-b92b-b0da15036b50","resolution":{"observed_at":"2026-08-07T05:35:42.756482Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"7010.28970","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.275041Z","title":"On the diffusion of test smells in automatically generated test code: An empirical study,","venue":null,"work_id":"2aa83e0a-4201-424d-b9c9-0e9252b2d39e","year":2016},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.949345Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:ea04737390a79882ae23e932ab2440a1c6a88d779b0699c22457e80396d765df","observation_id":"8a415b7c-ea53-4703-807f-b2a336d20b4f","resolution":{"observed_at":"2026-08-07T05:35:42.283784Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.736861Z","title":"When and why your code starts to smell bad,","venue":null,"work_id":"b14253c2-59e2-4acd-a015-ccb05293e184","year":2015},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.954889Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:f67039f1e07284e7ef7621a66f792650d4178a8ef337134c3f5d2460297ad247","observation_id":"db152661-eaa6-4851-905d-3fa62b9f4660","resolution":{"observed_at":"2026-08-07T05:35:42.741730Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.721327Z","title":"An empirical analysis of the distribution of unit test smells and their impact on software maintenance,","venue":null,"work_id":"66f0efef-8446-464f-abd6-f5c58cfba008","year":2012},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.960193Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:937c940107ce224791a26637726326f556b56f4ad9472df36502edfbb0f4e024","observation_id":"688e61d6-4cdb-49a5-a2de-f687eabf87dc","resolution":{"observed_at":"2026-08-07T05:35:42.726187Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:40.966381Z","title":"Just-in-time test smell detection and refactoring: The darts project,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.966381Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:542c5c60197dfe33456dd8d91ec7d34c33887fa535eb76f1ccb6eaddb905dfe5","observation_id":"568dc405-c62c-401c-b339-7d00346a4871","resolution":{"observed_at":"2026-08-07T05:35:40.966381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-07T05:35:40.971440Z","title":"Gpt-4 technical report,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.971440Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:160926956ec2fa651ebc9c3f43a1c6a297bb83c17da8357cc06a894087dbdf0a","observation_id":"1d6ec493-2c87-4666-9895-45d04d52dc64","resolution":{"observed_at":"2026-08-07T05:35:40.971440Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-07T05:35:40.977343Z","title":"The llama 3 herd of models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.977343Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:908510c6635c37389542ec6eb4a72a49633c9a4959cf609109c4d8cc7652b614","observation_id":"748b762e-2e35-4b29-a32f-bc2578dbe139","resolution":{"observed_at":"2026-08-07T05:35:40.977343Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05530","last_updated":"2024-12-16T17:39:39Z","snapshot_observed_at":"2026-07-06T17:41:42.995949Z","submitted_at":"2024-03-08T18:54:20Z","title":"Gemini 1.5: Unlocking multimodal understanding across millions of tokens of context","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.05530","snapshot_observed_at":"2026-08-07T05:35:40.982398Z","title":"Gemini 1.5: Unlocking multimodal understanding across millions of tokens of context,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.982398Z"},"links":{"cited_paper":"/paper/2403.05530","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:d77e7dd636f296f62086a0e1d2d6f70047658e9b357f837b4fd6ebc0e8fdd4e2","observation_id":"aea9bf46-fe6a-4274-9c69-d0eb6c0509fe","resolution":{"observed_at":"2026-08-07T05:35:40.982398Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.704367Z","title":"Codebert: A pre-trained model for programming and natural languages,","venue":null,"work_id":"7b7b2b4d-fde3-48fb-8405-d5fa2f0327cb","year":2020},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.987859Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:36f22e92d87f0b68f1a605c1613a7f4ac9f40f65b59477c5f814233ff56fe1e4","observation_id":"08a1ec8d-69e3-4d06-b149-1248e6a9c634","resolution":{"observed_at":"2026-08-07T05:35:42.710348Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.688004Z","title":"Codexglue: A machine learning benchmark dataset for code understanding and generation,","venue":null,"work_id":"ebe78fcc-f5b1-46f9-a3fc-c4e59b8d41b6","year":2021},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.992839Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:f5ec3f4602c17be86f44c4eaf7111a745b6f16fd72602e455c4d8ddcb7200e3a","observation_id":"37b2e45a-2e62-4589-a29e-30195993cc57","resolution":{"observed_at":"2026-08-07T05:35:42.692702Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.672772Z","title":"Top programming languages - the state of the octoverse 2022,","venue":null,"work_id":"39904449-2cf2-4b09-be7c-13367e13c9e7","year":2022},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.998380Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:86cc43db5252a04fad9dee33df4cb46889a5852ea3f5a03a72d22b2d333f6116","observation_id":"84c2c9f9-5d92-4b0f-858b-61b0e668b247","resolution":{"observed_at":"2026-08-07T05:35:42.677548Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.656396Z","title":"Utilization of pre-trained language model for adapter-based knowledge transfer in software engineering,","venue":null,"work_id":"58b8343e-c181-49c1-99ce-fb2978cae1ed","year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.003516Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:d17c12456829c3621408f2228a69c068385cfdf37a17dd4e8c05634cce51cca3","observation_id":"fe91d1d0-56e7-4e29-8824-7615b7467435","resolution":{"observed_at":"2026-08-07T05:35:42.660795Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.640561Z","title":"To- wards efficient fine-tuning of pre-trained code models: An experimental study and beyond,","venue":null,"work_id":"bcaaa806-8019-44c4-9294-e664ed36da74","year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.008583Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:c151ba6b85172d6c9fd0071eae704cf736d91e7651b2dbc02b23ecc7ce105934","observation_id":"11ba22fa-d781-4eb2-b61f-2334d7482892","resolution":{"observed_at":"2026-08-07T05:35:42.645541Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.013423Z","title":"An empirical comparison of pre-trained models of source code,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.013423Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:5732dbddd560eba49c1f50ef7f395a550e8ce24b5213a084dd0e2305f9b6cced","observation_id":"d2584707-6a9a-40f0-b8f8-2338e8f740ac","resolution":{"observed_at":"2026-08-07T05:35:41.013423Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.624026Z","title":"(2024) Testsmellsrefactoringbyllms","venue":null,"work_id":"57530050-2fe4-47a8-bad4-dcf79f18c3bf","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.018130Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:313934bc9f0446ebe0fd952f8a89413a81a337334efe775aeca72eaad19da07c","observation_id":"ab33be99-7a39-4325-9b52-4b05844df9d9","resolution":{"observed_at":"2026-08-07T05:35:42.629396Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.023442Z","title":"Large language models for software engineering: A systematic literature review,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.023442Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:9ad1a07a7cb6ca4893c6a4571d9a7ae8884014986b14a69e627614e8b0b86ba2","observation_id":"abc5f614-3046-456b-a24d-e7157091aecb","resolution":{"observed_at":"2026-08-07T05:35:41.023442Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.028343Z","title":"Software testing with large language models: Survey, landscape, and vision,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.028343Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:27a3195ce1c88bedfb39e0ba972da1bad361eb462547aea86cea3dcfc497d815","observation_id":"9007f747-60b9-4a80-9dc6-1896b7c41e46","resolution":{"observed_at":"2026-08-07T05:35:41.028343Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.034055Z","title":"An empirical evaluation of using large language models for automated unit test generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.034055Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:c7158dfa6f26ecf92c560a14631ab34c2f61a91ef0e11832e4f5fd1e618a238a","observation_id":"604f551a-1ac6-4e77-876c-0e1f253f2d38","resolution":{"observed_at":"2026-08-07T05:35:41.034055Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.587596Z","title":"Automated test case repair using language models,","venue":null,"work_id":"3161ac86-05e3-4705-828f-82ad1f67d457","year":2025},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.038870Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:42f3b827db1c952d1c54f1b7d81a1d401fa80532571b91dbdb5c3ec7cb7ee75c","observation_id":"16a0c54f-72b7-452f-ad91-f24e0a7b3b19","resolution":{"observed_at":"2026-08-07T05:35:42.592740Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.572020Z","title":"Chatunitest: a chatgpt- based automated unit test generation tool,","venue":null,"work_id":"ea797eaa-df5b-4bce-98a8-0f3b692c94ed","year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.044486Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:4813db3b472db088a4485302aaa093c5f2ce9d56b673734306ab40e98cee41e8","observation_id":"a607be27-a5ec-42e0-9f50-e90ee04bff8a","resolution":{"observed_at":"2026-08-07T05:35:42.577039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.555832Z","title":"An empirical study of using large language models for unit test generation,","venue":null,"work_id":"34856d8c-79f5-48cd-9d6a-4d88fb063fd4","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.049310Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:6d4ca2a73a10d8d32eaf68a97308a975de60f40ab039a4515635cd3a18e1534f","observation_id":"7501baa6-8233-41ba-922e-bc3d59d8bb78","resolution":{"observed_at":"2026-08-07T05:35:42.561428Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.539874Z","title":"Towards an understanding of large language models in software engineering tasks,","venue":null,"work_id":"bc1ebd87-a344-4a92-96f5-3361af2684c2","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.055046Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:21e61482523b2b5fa69de5ac0ed96c90360b9b60a9c98f7f25442c46c22caa7f","observation_id":"d4444700-dfc5-48e5-8936-21cab8d709c6","resolution":{"observed_at":"2026-08-07T05:35:42.545326Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.060347Z","title":"Pynose: a test smell detector for python,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.060347Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:3a47808c2ec7e55e20eea252d4de84bd6e9cb40a5bfd7dcbce0eba4a3f0659bf","observation_id":"01643684-ce9e-417e-b4ce-ee1b3ecd2056","resolution":{"observed_at":"2026-08-07T05:35:41.060347Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.065639Z","title":"Tempy: Test smell detector for python,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.065639Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:1c356406d361f91bab24c8cf73219cca73e83f62634901cc9f7c5bfadc759a2f","observation_id":"37e1047a-9964-4933-ad72-31fc5fd62b2c","resolution":{"observed_at":"2026-08-07T05:35:41.065639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"4624.34770","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.823854Z","title":"Handling test smells in python: Results from a mixed-method study,","venue":null,"work_id":"2216b61e-08e8-4190-9f80-2c0f4fff5b63","year":2021},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.070373Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:79840c77f8417ca0a86cba9c11e7a2f81dcb2aabf6833dcbd443ae8c8fd14d5c","observation_id":"1808a497-898a-4853-89d7-46f6bb0b93c4","resolution":{"observed_at":"2026-08-07T05:35:41.834355Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.519336Z","title":"A trend analysis of test smells in python test code over commit history,","venue":null,"work_id":"b624f8ca-1bd5-42df-9381-7c422513291e","year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.075592Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:61106a6fb6f18afed94b2c14004177ca7da4e1bc7dd0635ec39f26ad0e8f7d5b","observation_id":"0b467529-eab7-4e27-910f-cd8f9cd3ffd2","resolution":{"observed_at":"2026-08-07T05:35:42.525234Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.080279Z","title":"Pytest-smell: A smell detection tool for python unit tests,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.080279Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:30e9c5744a4eba74a162d6f49d361396b903792c3018c0d351d7497fe26ab060","observation_id":"a7830eb5-2748-452a-9657-0adedbb77461","resolution":{"observed_at":"2026-08-07T05:35:41.080279Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.085461Z","title":"A trend analysis of test smells in python test code over commit history,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.085461Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:365c27816eae0c9e3f0484a25e9257b9eacceeb9cdf59518ffe55cd9c52dc6fc","observation_id":"15ad84c6-b314-495e-a34f-c35c84b732d8","resolution":{"observed_at":"2026-08-07T05:35:41.085461Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.090496Z","title":"Detecting test smells in python test code generated by LLM: an empirical study with github copilot,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.090496Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:68aaebdd93f4ca8c11351840b41a8db89165d16b3705b957a81f36c7f1f28fc7","observation_id":"6e9b218f-0a68-451f-b6e1-53b2cb1d204a","resolution":{"observed_at":"2026-08-07T05:35:41.090496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.095249Z","title":"Tsdetect: An open source test smells detection tool,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.095249Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:8e62f23e6cf4f27df8ea928b49eb7c8988a78ca7917a641c4fc2dcd84ccd4636","observation_id":"27036ac3-7f44-4dc0-90e4-a0c565ff228e","resolution":{"observed_at":"2026-08-07T05:35:41.095249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.502190Z","title":"The secret life of test smells - an empirical study on test smell evolution and maintenance,","venue":null,"work_id":"d2b81647-354b-42c8-94e7-4a523603a94e","year":2021},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.100530Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:9ceab6e42d112f77ac1978c7861a9a75fa752de5966f7677cf4cab992b8530ea","observation_id":"f37741b2-b5b4-4428-a81e-8f656088678b","resolution":{"observed_at":"2026-08-07T05:35:42.507134Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.485015Z","title":"An empirical investigation into the nature of test smells,","venue":null,"work_id":"32ddda6c-2fac-462c-92f5-a8c058793bb2","year":2016},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.105536Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:1ac670da47231aadef08597caa752ef33f7612495595d180adf9956b3ae80f62","observation_id":"c5a472aa-b24b-487a-893f-dded75dc24ff","resolution":{"observed_at":"2026-08-07T05:35:42.490301Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.468208Z","title":"An empirical evaluation of raide: A semi-automated approach for test smells detection and refactoring,","venue":null,"work_id":"eb90a6ba-daeb-4bfe-9d9e-1837a2825320","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.110958Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:47457e8442e8794207bf07d984556404b72ce2ae6d9ac4f4bcdc68f61c489ae3","observation_id":"55a94a4e-6461-4fba-9507-0c1234c75a4d","resolution":{"observed_at":"2026-08-07T05:35:42.473239Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.452324Z","title":"Machine learning-based test smell detection,","venue":null,"work_id":"7e3674d8-7773-4dbf-b52f-57c1051cff36","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.116048Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:91fbc01d55a6ee076bc1cce77cb6fc9e0456cf65ea5bfebea72366178663bcb3","observation_id":"4c6232e1-13ef-464f-9272-7c5832301035","resolution":{"observed_at":"2026-08-07T05:35:42.457599Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.436335Z","title":"Ml test smell detection - online appendix,","venue":null,"work_id":"0e8df838-81f2-401c-b4a4-2e383fcf1a2d","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.121328Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:aa8a810ad98c2c00f6b1544e05b8d0c1d68452eb106c5f8cdfc2a4a7c4775833","observation_id":"17333b22-3370-4470-ada3-de15e04eaf6e","resolution":{"observed_at":"2026-08-07T05:35:42.441334Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06608","last_updated":"2025-02-26T18:59:01Z","snapshot_observed_at":"2026-07-06T18:28:24.821865Z","submitted_at":"2024-06-06T18:10:11Z","title":"The Prompt Report: A Systematic Survey of Prompt Engineering Techniques","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06608","snapshot_observed_at":"2026-08-07T05:35:41.127695Z","title":"The prompt report: A systematic survey of prompt engineering techniques,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.127695Z"},"links":{"cited_paper":"/paper/2406.06608","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:8f9fa85851463253104efaf83281a38bf382d5a2e0607a7712af549e04789c0c","observation_id":"979a075e-875f-4c4e-817a-651e9631ee66","resolution":{"observed_at":"2026-08-07T05:35:41.127695Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.11382","last_updated":"2023-02-21T12:42:44Z","snapshot_observed_at":"2026-07-06T14:54:37.559648Z","submitted_at":"2023-02-21T12:42:44Z","title":"A Prompt Pattern Catalog to Enhance Prompt Engineering with ChatGPT","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.11382","snapshot_observed_at":"2026-08-07T05:35:41.133174Z","title":"A prompt pattern catalog to enhance prompt engineering with chatgpt,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.133174Z"},"links":{"cited_paper":"/paper/2302.11382","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:42c13aee707964d6f06bd782e60341962cac322559089729ed7bfed71ddd8293","observation_id":"9e1f8995-92c8-409a-9f5e-50e8760143c9","resolution":{"observed_at":"2026-08-07T05:35:41.133174Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2109.01652","last_updated":"2022-02-08T20:26:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-09-03T17:55:52Z","title":"Finetuned Language Models Are Zero-Shot Learners","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.01652","snapshot_observed_at":"2026-08-07T05:35:41.138681Z","title":"Finetuned language models are zero-shot learners,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.138681Z"},"links":{"cited_paper":"/paper/2109.01652","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:96cf4ead4ed0a31cd0aef1e5ba558f5322356414491e2cf5fc941e79ce7c706d","observation_id":"e92c0d03-5ce9-4438-8ffd-bad42eb089bf","resolution":{"observed_at":"2026-08-07T05:35:41.138681Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.143500Z","title":"Language models are few-shot learners,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.143500Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:68f1455c54c38e72a1cc0f126bbd8a0a5af2c84565a7e88a75cbfabbc65d146d","observation_id":"bcf2c254-a14a-4662-9088-e3d66ed00209","resolution":{"observed_at":"2026-08-07T05:35:41.143500Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2201.11903","last_updated":"2023-01-10T23:07:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-01-28T02:33:07Z","title":"Chain-of-Thought Prompting Elicits Reasoning in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2201.11903","snapshot_observed_at":"2026-08-07T05:35:41.154116Z","title":"Chain-of-thought prompting elicits reasoning in large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.154116Z"},"links":{"cited_paper":"/paper/2201.11903","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:ab190d5d57fdcef9e769102da99080098675bfdc1c4d18360fb2e6949a850137","observation_id":"e8bcc4d9-42d0-42c7-bd7c-674bfb4baf9a","resolution":{"observed_at":"2026-08-07T05:35:41.154116Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.409666Z","title":"Enhancing zero-shot chain-of-thought reasoning in large language models through logic,","venue":null,"work_id":"dacbe487-1248-4864-aab7-5bcdb43d2e6a","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.159476Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:3a70595f3570a29fb19f5dce2d4bc452ed08e3e22af8ad4ef95aa7fadcc14867","observation_id":"f6690eac-0686-4d46-8b51-481bc0203440","resolution":{"observed_at":"2026-08-07T05:35:42.414785Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.164376Z","title":"Wilcoxon, Individual Comparisons by Ranking Methods","venue":null,"work_id":null,"year":1992},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.164376Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:1f4ac7857ea22b6e035e99055ff3aa6aee2297cdba93136e4e24f0c3807add3c","observation_id":"c01ba966-08cf-47cb-af39-6179098f4152","resolution":{"observed_at":"2026-08-07T05:35:41.164376Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14261","last_updated":"2024-02-22T03:51:34Z","snapshot_observed_at":"2026-07-06T17:33:45.647306Z","submitted_at":"2024-02-22T03:51:34Z","title":"Copilot Evaluation Harness: Evaluating LLM-Guided Software Programming","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14261","snapshot_observed_at":"2026-08-07T05:35:41.169640Z","title":"Copilot evaluation harness: Evaluating llm-guided software programming,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.169640Z"},"links":{"cited_paper":"/paper/2402.14261","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:50eab0187efe50669c1f14ab49ae4926f4b11183e415563684e49e43e5983972","observation_id":"4b784f3d-ec10-4cef-ab7d-6c9848bebb76","resolution":{"observed_at":"2026-08-07T05:35:41.169640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.394234Z","title":"Towards effective validation and integration of llm-generated code,","venue":null,"work_id":"57ab2c31-5f7f-40d8-b7be-2324e9b730b3","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.174703Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:a1c858f36587c5b6a41c476f66d364da39b754559e27318347d352141ad07a57","observation_id":"203e52a2-3eb4-4dd5-93e2-9ec74aabd054","resolution":{"observed_at":"2026-08-07T05:35:42.399120Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.378728Z","title":"Challenges and opportunities in integrating llms into con- tinuous integration/continuous deployment (ci/cd) pipelines,","venue":null,"work_id":"f3a55412-31fc-4066-bc74-1cff34b3e926","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.179595Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:0366f40b9c4ffb1bbae61d9fcafe16b8338ca686d1835a4db15b80dc71e28f2d","observation_id":"4d3c0078-843f-4abe-a90e-af4cd3796809","resolution":{"observed_at":"2026-08-07T05:35:42.383731Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.362357Z","title":"Next-generation refactoring: Combining llm insights and ide capabilities for extract method,","venue":null,"work_id":"16f7f52f-a69a-415b-9bd0-17dbc3165d21","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.184895Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:e15fdbae07c790bcc2b2c2147f35bb228bc94b3a808459c95b132cfcf88ebabf","observation_id":"3866c2cb-de0c-4740-9325-25b67ecdb778","resolution":{"observed_at":"2026-08-07T05:35:42.367755Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.345946Z","title":"Llm-based multi-agent systems for software engineering: Literature review, vision and the road ahead,","venue":null,"work_id":"2bc82942-eb68-4592-a381-c780b41e26ad","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.189548Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:4ce5039deb0cfa53b3fae40dd9632a4db4ef8246b5730d5589e485f7fa44502a","observation_id":"bcf1d462-afeb-49f6-8b97-81e8497ba694","resolution":{"observed_at":"2026-08-07T05:35:42.350942Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.329809Z","title":"Autorefactoring: A platform to build refactoring agents,","venue":null,"work_id":"c1542b1b-8ea6-4cea-af0a-9ab9eded7205","year":2015},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.194572Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:722180eeb0d5c907a8e6e51d36128b73abcbc2b9dc79a731940e11f7bcdcc970","observation_id":"ba98147c-6e57-4c7a-86bd-39625fad115e","resolution":{"observed_at":"2026-08-07T05:35:42.335039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.19437","last_updated":"2025-02-18T17:26:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-27T04:03:16Z","title":"DeepSeek-V3 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.19437","snapshot_observed_at":"2026-08-07T05:35:41.200114Z","title":"Deepseek-v3 technical report,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.200114Z"},"links":{"cited_paper":"/paper/2412.19437","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:c46ec9de51924e38262072f484c3c27e75f736f9f707c72484b86b40bcaa9947","observation_id":"de96ca94-7c8f-43c9-8533-e95f24ffaad4","resolution":{"observed_at":"2026-08-07T05:35:41.200114Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-07T05:35:41.204978Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.204978Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:4af734ad0d41668232aaaa7e03e0bafa7f4607eb9879de19bacb772590f51483","observation_id":"d33ee6e0-9751-4ba5-9b5b-b0d791ab5ddc","resolution":{"observed_at":"2026-08-07T05:35:41.204978Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.313498Z","title":"Runeson, M","venue":null,"work_id":"8f6c9e2c-8a72-4f4b-94be-4e6d5daa2d1b","year":2012},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.210159Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:91e8fcc375160117978cbc6cfab259fc4dd48c20edbbfb34e791ee7a19e617b6","observation_id":"d9736758-fc93-45ee-819c-908bea04db40","resolution":{"observed_at":"2026-08-07T05:35:42.318400Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.297332Z","title":"Qualitative methods in empirical studies of software engineering,","venue":null,"work_id":"02caa27b-ba13-43bb-8a1a-f35cee23b65c","year":1999},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.214905Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:3b7be6247685a270d3841bc9ce026c4198a4aa89db7e27193513e3a5398cfce3","observation_id":"9a477ee6-450c-4e12-8886-ead2dd72cacc","resolution":{"observed_at":"2026-08-07T05:35:42.302511Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2005.14165","last_updated":"2020-07-22T19:47:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-05-28T17:29:03Z","title":"Language Models are Few-Shot Learners","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.14165","snapshot_observed_at":"2026-08-07T05:35:41.149159Z","title":"Available: https://arxiv.org/abs/2005.14165","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.149159Z"},"links":{"cited_paper":"/paper/2005.14165","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:572cf340e83d27a56a52b0ec27135d5feb0859ea44aa0ceff7ae450608334c01","observation_id":"a9133213-2008-441f-ab73-53fa2c8f35fb","resolution":{"observed_at":"2026-08-07T05:35:41.149159Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","latest_version":1,"primary_category":"cs.SE","snapshot_observed_at":"2026-08-07T05:27:42.571780Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study"},"reference_resolution":{"displayed":53,"state_counts":{"malformed_identifier":0,"metadata_mismatch":2,"parse_uncertain":0,"unresolved":24,"verified_exact":0,"verified_fuzzy":27},"total_outbound_references":53},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 53 of 53 outbound references and 3 inbound Pith citation observations for arXiv:2506.07594."}