{"as_of":"2026-08-18T22:35:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3bee82b98649094a6ab8f6d5c14e7989a243d1e4012cff0d6d1e08759a80983f","coverage":[{"denominator":45,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":45,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T01:04:03.644079Z","state":"measured"},{"denominator":49,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":49,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":4,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":4,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T11:17:35.374595Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-11T17:21:10.908503Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.18989","snapshot_observed_at":"2026-08-16T11:17:35.374595Z","title":"How propense are large language models at producing code smells? a benchmarking study,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2504.15989","last_updated":"2025-05-29T18:29:37Z","snapshot_observed_at":"2026-08-16T11:10:56.816544Z","submitted_at":"2025-04-22T15:51:00Z","title":"Optimizing Token Consumption in LLMs: A Nano Surge Approach for Code Reasoning Efficiency","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-16T11:17:35.374595Z"},"links":{"cited_paper":"/paper/2412.18989","citing_paper":"/paper/2504.15989"},"observation_digest":"sha256:c9961380b31189fe2faadd64980cdbde03c4050769551a94cd17027625a42a34","observation_id":"9e665133-d85c-4869-ae5b-c4ded1d91d73","resolution":{"observed_at":"2026-08-16T11:17:35.374595Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.18989","snapshot_observed_at":"2026-08-16T11:15:40.402491Z","title":"How propense are large language models at producing code smells? a benchmarking study","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2504.16027","last_updated":"2025-04-22T16:44:39Z","snapshot_observed_at":"2026-08-17T10:59:32.935243Z","submitted_at":"2025-04-22T16:44:39Z","title":"Benchmarking LLM for Code Smells Detection: OpenAI GPT-4.0 vs DeepSeek-V3","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-16T11:15:40.402491Z"},"links":{"cited_paper":"/paper/2412.18989","citing_paper":"/paper/2504.16027"},"observation_digest":"sha256:982f185f0eaa05164c182607d2407c330e8bf51853785196ee774c7f63102b1d","observation_id":"c8d06b7e-9688-49bb-bc04-b9eb0b08fdf5","resolution":{"observed_at":"2026-08-16T11:15:40.402491Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.18989","snapshot_observed_at":"2026-08-04T06:49:10.487496Z","title":"Palacio, and Denys Poshyvanyk","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.15817","last_updated":"2026-08-11T17:48:29Z","snapshot_observed_at":"2026-08-14T23:11:29.321629Z","submitted_at":"2025-11-19T19:18:28Z","title":"A Causal Perspective on Measuring, Explaining and Mitigating Smells in LLM-Generated Code","version":6},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-04T06:49:10.487496Z"},"links":{"cited_paper":"/paper/2412.18989","citing_paper":"/paper/2511.15817"},"observation_digest":"sha256:bbc3c7fc1eca75544fc88d29c3814f6e7b1d9daeda619f68d99c31f9444807f9","observation_id":"75937992-3c33-4279-abbe-2724dc63252a","resolution":{"observed_at":"2026-08-04T06:49:10.487496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"cited_work":{"arxiv_id":"2412.18989","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.18989","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Palacio, and Denys Poshyvanyk","venue":null,"work_id":"a7123f0c-7412-434f-8497-cb4d8dbae95e","year":2025},"citing_paper":{"arxiv_id":"2605.05267","last_updated":"2026-05-06T09:38:31Z","snapshot_observed_at":"2026-08-04T17:41:12.164736Z","submitted_at":"2026-05-06T09:38:31Z","title":"Bridging Generation and Training: A Systematic Review of Quality Issues in LLMs for Code","version":1},"reference_index":123,"source":"pdf_text","source_observed_at":"2026-05-08T17:37:51.790000Z"},"links":{"cited_paper":"/paper/2412.18989","citing_paper":"/paper/2605.05267"},"observation_digest":"sha256:f56237c75824d6b762b04c5ddb597d908b99124bbb62276feecb6197ff9fcbe4","observation_id":"c4be3898-66d8-4c11-9a73-d9fa90977d3d","resolution":{"observed_at":"2026-05-11T17:21:10.911009Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2412.18989/citation-record","integrity":"/paper/2412.18989/integrity","json":"/paper/2412.18989/citation-record.json","paper":"/paper/2412.18989"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.657090Z","title":"An empirical study on the usage of transformer models for code completion,","venue":null,"work_id":"46c80ba2-f568-4408-84b3-46f01111d27f","year":2022},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.403178Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:5a2630573854e3ebcbb650601b30e9544445cb1ee79e9cfa8067a1b2ed978f6a","observation_id":"8c8e02bf-26b7-4949-a787-60feb4091439","resolution":{"observed_at":"2026-08-11T01:04:04.663359Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15852","last_updated":"2024-10-31T14:43:58Z","snapshot_observed_at":"2026-08-16T14:07:05.751912Z","submitted_at":"2024-03-23T14:04:48Z","title":"SOEN-101: Code Generation by Emulating Software Process Models Using Large Language Model Agents","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.15852","snapshot_observed_at":"2026-08-11T01:04:03.409292Z","title":"When LLM-based Code Generation Meets the Software Development Process,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.409292Z"},"links":{"cited_paper":"/paper/2403.15852","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:4107ccadbdda1dd689a2a50e70b786f490652c474d15cf3695a26c3f0b66ee24","observation_id":"8e557a6a-917a-4907-9809-e78844a645d6","resolution":{"observed_at":"2026-08-11T01:04:03.409292Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.642219Z","title":"Toward deep learning software reposi- tories,","venue":null,"work_id":"d7d48263-9e86-46c0-9165-d029241574a4","year":2015},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.415672Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:a2c9075f4cc92eb4579823a3480a09717babae4b0ce26996138b42f1cad34005","observation_id":"3708045c-766f-4a85-9052-c70a950cc1c6","resolution":{"observed_at":"2026-08-11T01:04:04.646725Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.626045Z","title":"Few-shot training llms for project-specific code-summarization,","venue":null,"work_id":"edb2258f-ec33-497c-a17b-d412909b14fa","year":2023},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.420970Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:05f45ce15ddce1cbd0f283d1fa11603e44618e9b5e01d7efd02ed4811a0ccc32","observation_id":"3a94e3f5-60d1-4571-97a2-eb2b061cd0d6","resolution":{"observed_at":"2026-08-11T01:04:04.631605Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.604545Z","title":"Inferfix: End-to-end program repair with llms,","venue":null,"work_id":"5b3b71cd-effb-4656-8a53-e7a75d11a7e5","year":2023},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.425961Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:32bc2e8449ac8146be1410b106e34fbf01cc830fec560b54eed09afcaa5803b6","observation_id":"49064042-8472-40b9-98c2-5b66006880b9","resolution":{"observed_at":"2026-08-11T01:04:04.610995Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.584109Z","title":"Deep learning code fragments for code clone detection,","venue":null,"work_id":"be2f6892-a36e-481b-8947-5b75a0ea3bb6","year":2016},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.430951Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:655e8f141f5b2eced2cefceaa8f0ec021d7a4bbddf4556653ae6ef191ce799b0","observation_id":"1feb2ec2-6ea3-41c4-99e9-b77bd00cb82c","resolution":{"observed_at":"2026-08-11T01:04:04.591439Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.567822Z","title":"On learning meaningful assert statements for unit test cases,","venue":null,"work_id":"d65fc99b-91b5-4054-9201-d06a35d41466","year":2020},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.436960Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:19a46edd8ade3505e5292f678bf976611ba0eb0b485f8c160f7bd5f6bb2233f3","observation_id":"25084cca-18a8-4ac4-81ed-8bd501399e17","resolution":{"observed_at":"2026-08-11T01:04:04.572962Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.551478Z","title":"A systematic literature review on the use of deep learning in software engineering research,","venue":null,"work_id":"c127dbcc-1b55-4015-aa03-b17a4d4fbe27","year":2022},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.441520Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:fa819d2f870c0bc7dd8132fbc8673304d954ad4023295e5862d518ea0e7aca42","observation_id":"bf8022e6-97f5-4fd6-b529-bb16756d5eb4","resolution":{"observed_at":"2026-08-11T01:04:04.557003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.532274Z","title":"When and Why Your Code Starts to Smell Bad (and Whether the Smells Go Away),","venue":null,"work_id":"455ef518-0963-4c2a-8815-a4b451bb4f01","year":2017},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.445796Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:1bf71056f654ccbe94a8ec25dbf124b09f3f3ac76d51ea49f88a83be803d8983","observation_id":"d9daa52d-26e9-4ff6-969f-92aa962e5e51","resolution":{"observed_at":"2026-08-11T01:04:04.538147Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.512903Z","title":"On the diffuseness and the impact on maintainability of code smells: a large scale empirical investigation,","venue":null,"work_id":"0868fb29-105b-424a-b99b-e67f3691fc3b","year":2018},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.450187Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:c6a5eaca002f5d097a25ffcc9fa3c88bd0db57a08b4e3f1c1c62fa77292f5145","observation_id":"1d358ee4-d47a-49a8-8f7e-611103c185d6","resolution":{"observed_at":"2026-08-11T01:04:04.518728Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.08549","last_updated":"2024-08-16T06:31:44Z","snapshot_observed_at":"2026-08-16T13:26:05.627787Z","submitted_at":"2024-08-16T06:31:44Z","title":"Vulnerability Handling of AI-Generated Code -- Existing Solutions and Open Challenges","version":1},"cited_work":{"arxiv_id":"2408.08549","doi":null,"metadata_source":"pith","pith_arxiv_id":"2408.08549","snapshot_observed_at":"2026-08-11T01:04:04.019216Z","title":"Vulnerability Handling of AI-Generated Code -- Existing Solutions and Open Challenges","venue":"cs.SE","work_id":"4b21206a-d07a-4138-ae03-f7de32e42187","year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.454292Z"},"links":{"cited_paper":"/paper/2408.08549","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:1567fcede5d40d40d0fa37233ca8c05139f272a6f321c1fd5d8a8fc1bae2a3f9","observation_id":"08cafb2d-6691-48e3-8983-fca985485e3f","resolution":{"observed_at":"2026-08-11T01:04:04.024541Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.492908Z","title":"CodeLMSec benchmark: Systematically evaluating and finding security vulnerabilities in black-box code lan- guage models,","venue":null,"work_id":"7a952a16-2a2d-4ccb-acb3-df1661aa987f","year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.458850Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:1eb9e9584970158d5f77b76f1dd207348c71cfa4a7cfe112457d6d1445a853b5","observation_id":"9beb1c46-93eb-45af-ba9b-eb21d3080a0c","resolution":{"observed_at":"2026-08-11T01:04:04.499930Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.00889","last_updated":"2024-09-04T23:37:11Z","snapshot_observed_at":"2026-08-16T14:46:43.013864Z","submitted_at":"2023-11-01T22:46:31Z","title":"SALLM: Security Assessment of Generated Code","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.00889","snapshot_observed_at":"2026-08-11T01:04:03.463934Z","title":"SALLM: Security Assessment of Generated Code,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.463934Z"},"links":{"cited_paper":"/paper/2311.00889","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:5d611fac4f95f398c3f266b7c5b5c612e05ea38504c5cfb9d84c0a11ba7a8777","observation_id":"604cae6f-c46c-4e2e-ab6e-d00da5a12fac","resolution":{"observed_at":"2026-08-11T01:04:03.463934Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.01202","last_updated":"2024-05-02T11:44:52Z","snapshot_observed_at":"2026-08-18T21:20:18.384401Z","submitted_at":"2024-05-02T11:44:52Z","title":"DLAP: A Deep Learning Augmented Large Language Model Prompting Framework for Software Vulnerability Detection","version":1},"cited_work":{"arxiv_id":"2405.01202","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.01202","snapshot_observed_at":"2026-08-11T01:04:03.980211Z","title":"DLAP: A Deep Learning Augmented Large Language Model Prompting Framework for Software Vulnerability Detection","venue":"cs.SE","work_id":"44c2792d-4275-45cc-b6ba-0afa84b3d860","year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.468913Z"},"links":{"cited_paper":"/paper/2405.01202","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:7cf6c0f0f4c79a8059a5fe94e4b63ab2ef423f58882744541a2d6085645f4204","observation_id":"d75ffbc4-0343-4067-b2aa-57b852050647","resolution":{"observed_at":"2026-08-11T01:04:03.985672Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.19261","last_updated":"2024-07-30T12:16:54Z","snapshot_observed_at":"2026-08-18T18:03:27.278966Z","submitted_at":"2024-07-27T14:00:05Z","title":"Evaluating Large Language Models in Detecting Test Smells","version":2},"cited_work":{"arxiv_id":"2407.19261","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.19261","snapshot_observed_at":"2026-08-11T01:04:03.954758Z","title":"Evaluating Large Language Models in Detecting Test Smells","venue":"cs.SE","work_id":"f17a9a79-8eea-4832-be89-a1ab6f77e817","year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.474105Z"},"links":{"cited_paper":"/paper/2407.19261","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:3d13cc4defa8f7ea2de837e40086dffbc577cb8d371cac152f164dc4e4b06354","observation_id":"616c364a-45f6-4e45-a6e0-8660d05b530d","resolution":{"observed_at":"2026-08-11T01:04:03.960809Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.471865Z","title":"Code smell detection using hy- brid machine learning algorithms,","venue":null,"work_id":"a74bcfc6-6da9-4911-b04f-cbb218cd5d19","year":2023},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.479283Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:026cf8b96c00ecb307ce7b2d294971338e1790c55262ae2443badea71119b221","observation_id":"88b79e18-58ca-4138-9e1e-ea1cbd336ac1","resolution":{"observed_at":"2026-08-11T01:04:04.478649Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.452591Z","title":"Multi-Label Code Smell Detection with Hybrid Model based on Deep Learning,","venue":null,"work_id":"37141c92-ea67-432a-96b9-0a8a0b5f6480","year":2022},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.483872Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:e75e943c3289dd55d4ebe4283c2b50cac3a7389998b0bc55dd86682f9ce7f2d1","observation_id":"fcebad28-2922-4413-a8d3-c80377a9a3ee","resolution":{"observed_at":"2026-08-11T01:04:04.459092Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.431897Z","title":"iSMELL: Assembling LLMs with Expert Toolsets for Code Smell Detection and Refactoring,","venue":null,"work_id":"9d014f18-7579-4348-9971-c2933dd623fd","year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.489197Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:7ceef524bc762918a0a054043de25e5594c6a0a20449cca3ee6a589bd180649b","observation_id":"afc4a3f8-d920-4833-a499-935fc97b243e","resolution":{"observed_at":"2026-08-11T01:04:04.438375Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.412367Z","title":"BLEU: a method for automatic evaluation of machine translation,","venue":null,"work_id":"68f73940-7618-45f7-889c-2eb9e31a0b70","year":2002},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.493760Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:1dadf08266ccde32a37fd323d16f9b6eac7926c9f57dda244e7780349bc1cc9a","observation_id":"872550b3-bded-4f1f-8a79-fee044d9b43f","resolution":{"observed_at":"2026-08-11T01:04:04.417579Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2009.10297","last_updated":"2020-09-27T04:07:11Z","snapshot_observed_at":"2026-08-12T14:26:33.940158Z","submitted_at":"2020-09-22T03:10:49Z","title":"CodeBLEU: a Method for Automatic Evaluation of Code Synthesis","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.10297","snapshot_observed_at":"2026-08-11T01:04:03.498305Z","title":"CodeBLEU: a Method for Automatic Evaluation of Code Synthesis,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.498305Z"},"links":{"cited_paper":"/paper/2009.10297","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:44228c197c4c2f412b120f641087eb22f107c4fcae6dc6179c51175ed7d745c7","observation_id":"31fd5b27-340c-49fd-87e2-4660d7410cbd","resolution":{"observed_at":"2026-08-11T01:04:03.498305Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.393423Z","title":"ROUGE: A Package for Automatic Evaluation of Sum- maries,","venue":null,"work_id":"4e309bbe-885c-4c09-b141-1b41c2c323b1","year":2004},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.503459Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:fa5efce86b059a9a6957d8819fcb68b553209eda4783e5ba0506814f4c2188c0","observation_id":"535866fd-a9b0-41fc-807f-5ff433ee9696","resolution":{"observed_at":"2026-08-11T01:04:04.400047Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.373720Z","title":"Meteor: an automatic metric for mt evaluation with high levels of correlation with human judgments,","venue":null,"work_id":"061688a6-8738-4bae-8528-14c5fcb94b61","year":2007},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.508403Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:fc0927dee2ac722262039ffbfbca5758ab74ad07e07e0169a27da7324f5f53a6","observation_id":"e8b9eec1-90af-4a1c-b751-d4012a64dd18","resolution":{"observed_at":"2026-08-11T01:04:04.380061Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.354724Z","title":"A systematic literature review on the use of deep learning in software engineering research,","venue":null,"work_id":"36d94a30-aaab-40ef-8935-e4f99ae9cd05","year":2022},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.513445Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:c9e038c0703e49d8c91d3b65990eb31d07e5b358047e267b1ed134c1940f2f0f","observation_id":"d58e88c7-0c18-41e5-9ac3-51fec2e21e34","resolution":{"observed_at":"2026-08-11T01:04:04.360462Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:03.518465Z","title":"Towards More Trust- worthy and Interpretable LLMs for Code through Syntax-Grounded Explanations,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.518465Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:0cf1e5aecbe16fc1a083063942d3156f81bfc45fb6b089d5c25a6a7901c0f06a","observation_id":"f91b7c1d-0f8a-463d-b121-c915bd08c090","resolution":{"observed_at":"2026-08-11T01:04:03.518465Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.335054Z","title":"Which syntactic capabilities are statistically learned by masked language models for code?","venue":null,"work_id":"81650b08-e83a-473f-9681-1d78ece33c52","year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.523362Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:a4854be946d9924d0f49ff2fbeee612ab8a232d95e2514361781272538aaae6e","observation_id":"c6271d76-e3a4-4208-b32b-886d0b184242","resolution":{"observed_at":"2026-08-11T01:04:04.341291Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.316438Z","title":"Codesmells,","venue":null,"work_id":"ce04ebd1-cd3e-4c65-af95-41da25369a74","year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.528834Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:ebca29ba5dcef92d4a6a024bfcdd61ce0b701889de770fa75124a2541aeb0c7c","observation_id":"d843fc25-c085-40c5-a4bb-2e7e0b3b0dae","resolution":{"observed_at":"2026-08-11T01:04:04.322889Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1506.02078","last_updated":"2015-11-17T02:42:24Z","snapshot_observed_at":"2026-08-17T10:57:16.210995Z","submitted_at":"2015-06-05T22:33:04Z","title":"Visualizing and Understanding Recurrent Networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1506.02078","snapshot_observed_at":"2026-08-11T01:04:03.533227Z","title":"Visualizing and Understanding Recurrent Networks,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.533227Z"},"links":{"cited_paper":"/paper/1506.02078","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:1bc1abda558a6daa477a4ef374e747d39794267508819de29024c531e04e1333","observation_id":"8dcf6c74-f5d1-493f-9e93-defc121d301c","resolution":{"observed_at":"2026-08-11T01:04:03.533227Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.285878Z","title":null,"venue":null,"work_id":"8f64bbc2-f9aa-461c-8c93-c79cd0d5a997","year":2003},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.537843Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:81dfcae3d1e605271d08bdc24fbbdb7aed1073d2472e37f63d5badfd7057b562","observation_id":"53a9e426-9a90-445d-872d-308cadaa0d49","resolution":{"observed_at":"2026-08-11T01:04:04.291678Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.268900Z","title":"Benchmarking causal study to interpret large language models for source code,","venue":null,"work_id":"f209b763-0520-4331-8365-88c7e092f3fc","year":2023},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.544720Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:9f2a691bb0c53745587c46fd2567707e2819f201a7b0c0036e67e26dfbed9682","observation_id":"59a871b3-bfc3-46ff-82a7-213bc33082dd","resolution":{"observed_at":"2026-08-11T01:04:04.274862Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.12950","last_updated":"2024-01-31T19:47:26Z","snapshot_observed_at":"2026-08-18T03:59:39.242039Z","submitted_at":"2023-08-24T17:39:13Z","title":"Code Llama: Open Foundation Models for Code","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.12950","snapshot_observed_at":"2026-08-11T01:04:03.549799Z","title":"Code Llama: Open Foundation Models for Code,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.549799Z"},"links":{"cited_paper":"/paper/2308.12950","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:72018096788b169e26188714893fe9c2daca1166bc29bc33beccecafdf6b7e0f","observation_id":"75a4ba42-b2fc-46b5-b3e8-8a138d00d3b3","resolution":{"observed_at":"2026-08-11T01:04:03.549799Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06825","last_updated":"2023-10-10T17:54:58Z","snapshot_observed_at":"2026-08-17T20:30:34.016254Z","submitted_at":"2023-10-10T17:54:58Z","title":"Mistral 7B","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06825","snapshot_observed_at":"2026-08-11T01:04:03.556840Z","title":"Mistral 7B,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.556840Z"},"links":{"cited_paper":"/paper/2310.06825","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:374bfef9e87dd6476e60c8f7de958209c667d9c311131732af84bd16f99d748f","observation_id":"f8e76506-26e5-45db-bf6d-e992454b5e6d","resolution":{"observed_at":"2026-08-11T01:04:03.556840Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.251169Z","title":"Code smell detection using hy- brid machine learning algorithms,","venue":null,"work_id":"cb6b921e-e524-48a3-b75a-43febaddfcab","year":2023},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.562787Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:4c86d028ab4271fe0b3c6088efc65109ccc73f138322869462beeb8ef84d7306","observation_id":"7b08d064-358f-46ec-92d7-ce1252d348c2","resolution":{"observed_at":"2026-08-11T01:04:04.256271Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.234636Z","title":"Machine learning powered code smell de- tection as a business improvement tool,","venue":null,"work_id":"d2be138b-ec78-4d0e-89ee-6b28394db9fe","year":null},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.568264Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:a73d118dc865b33c2fb73a42057a0e818b9cf143f1a394dc6fceabd2d617fda9","observation_id":"6c9b0ac3-f9b6-4aa7-8564-e1f76c1ae49b","resolution":{"observed_at":"2026-08-11T01:04:04.239837Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12513","last_updated":"2024-06-18T11:29:34Z","snapshot_observed_at":"2026-08-18T09:28:27.823248Z","submitted_at":"2024-06-18T11:29:34Z","title":"Can We Trust Large Language Models Generated Code? A Framework for In-Context Learning, Security Patterns, and Code Evaluations Across Diverse LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12513","snapshot_observed_at":"2026-08-11T01:04:03.574583Z","title":"Can We Trust Large Language Mod- els Generated Code? A Framework for In-Context Learning, Security Patterns, and Code Evaluations Across Diverse LLMs,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.574583Z"},"links":{"cited_paper":"/paper/2406.12513","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:059e6eb531e03d7994b7324a01baef73bb4c49481192fff3ce36e456aa91ce9b","observation_id":"c1e00833-c93b-4aa2-b548-6e61aa40d5f3","resolution":{"observed_at":"2026-08-11T01:04:03.574583Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.216410Z","title":"A Systematic Literature Review on the Code Smells Datasets and Validation Mechanisms,","venue":null,"work_id":"25fe0dc1-acaa-4de1-b7f9-7bda2517e521","year":2023},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.580778Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:609434f9b412139394fc4c28531c51f237ec7386a6bb7b69f5ade5f727133c46","observation_id":"b746481d-dcd9-4cf4-8980-9f7b8922378e","resolution":{"observed_at":"2026-08-11T01:04:04.222541Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.199684Z","title":"ml-Codesmell: A code smell prediction dataset for machine learning approaches,","venue":null,"work_id":"eb2afaef-38df-4fd8-bf47-c8052656918d","year":2022},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.589571Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:031ae66d341ed3b477d7d2c5221bc43cf86f42e9b5ed5df3e3c2092fc990ea1b","observation_id":"66a55f9b-506f-41dd-a668-36d8ef0acdc9","resolution":{"observed_at":"2026-08-11T01:04:04.204730Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.182176Z","title":"Mlcq: Industry-relevant code smell data set,","venue":null,"work_id":"d524e475-50bf-433a-b96a-05d3d93b7d66","year":2020},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.596066Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:339c4b7aed7f9c1b48d0b89ac86f4c6ce905c31e64e14fc3fde9c0aa1a09f6e0","observation_id":"912c9ca7-7a9b-4cf9-b82d-db08d4c33db1","resolution":{"observed_at":"2026-08-11T01:04:04.187375Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.165991Z","title":"Evaluating the accuracy of machine learning algorithms on detecting code smells for different developers,","venue":null,"work_id":"210c3d10-7104-4314-b59d-1a53884f6e73","year":2017},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.601761Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:b4f5613461b6d5eb3cade80a57c199ef5a9c922121d7209f9d6c075758c082a4","observation_id":"2538c857-0072-43b6-91b9-45b9acff69f0","resolution":{"observed_at":"2026-08-11T01:04:04.171338Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.146907Z","title":"DACOS—a manually annotated dataset of code smells,","venue":null,"work_id":"ecae44d7-d853-413e-840c-02f11d8d025f","year":2023},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.606679Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:7a6e332289405b6c9a28e0ec64bc606d1fb30df28dd2a52dc98d58e593bff032","observation_id":"fcb823cf-1ffc-4338-a32d-292805d193a8","resolution":{"observed_at":"2026-08-11T01:04:04.153186Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.123815Z","title":"The Technical Debt Dataset,","venue":null,"work_id":"f6164d37-4962-428a-9f56-d644c6b236f0","year":2019},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.612873Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:093f58351f330eb91a1e94afe3fcb0bdd55a8c59d369e004e836dc611f63df47","observation_id":"912d416b-9b36-40d1-a474-3e8579c501c2","resolution":{"observed_at":"2026-08-11T01:04:04.129712Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.104895Z","title":"Using code evolution information to improve the quality of labels in code smell datasets,","venue":null,"work_id":"3f1bcbfe-06e2-49f2-b2d6-32b83b45cf33","year":2018},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.618462Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:e7aff3d6bffae3586b5210e7745436c2a6b08a6c185b20a6633481647b752650","observation_id":"69b60e96-dbfd-4169-8031-fa9fb4039542","resolution":{"observed_at":"2026-08-11T01:04:04.110103Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10398","last_updated":"2024-02-16T01:50:46Z","snapshot_observed_at":"2026-08-16T23:05:37.921126Z","submitted_at":"2024-02-16T01:50:46Z","title":"Prompt Learning for Multi-Label Code Smell Detection: A Promising Approach","version":1},"cited_work":{"arxiv_id":"2402.10398","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.10398","snapshot_observed_at":"2026-08-11T01:04:03.697423Z","title":"Prompt Learning for Multi-Label Code Smell Detection: A Promising Approach","venue":"cs.SE","work_id":"9cb3f2a2-e699-4ecc-824c-45deddc1eb79","year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.625710Z"},"links":{"cited_paper":"/paper/2402.10398","citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:2588cccdb204a00983e5291f7da4c09fd9186d4708aa505de8691b6adc13b15e","observation_id":"8b65d380-44ba-4d29-8582-1a4e1e7ec806","resolution":{"observed_at":"2026-08-11T01:04:03.704852Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.087725Z","title":"Toward a theory of causation for interpreting neural code models,","venue":null,"work_id":"e6f6d69e-b516-47ee-810d-dd47daf76a7c","year":2024},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.632997Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:60693e1272b4ac3fce4d0f161f27d00208b8f373144cab697c42f30a6c6e76af","observation_id":"944f827f-5e2a-4dfb-bb44-64a276c1f7d0","resolution":{"observed_at":"2026-08-11T01:04:04.093037Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.070469Z","title":"”why should i trust you?","venue":null,"work_id":"eb55d61a-c08d-4c26-a705-4b71a0cc4a9f","year":2016},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.638043Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:78398f10c37525b642dcb280ed20b9f54c08655852496340de28f4ff43c87b40","observation_id":"1ae8623c-9850-465a-ad51-b3071431b412","resolution":{"observed_at":"2026-08-11T01:04:04.076285Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T01:04:04.054153Z","title":"A unified approach to interpreting model predictions,","venue":null,"work_id":"0dc7dd29-0899-420f-af6f-fd18c5f21d31","year":2017},"citing_paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-11T01:04:03.644079Z"},"links":{"citing_paper":"/paper/2412.18989"},"observation_digest":"sha256:75a818c8b9b32de4e3ce0a074128191ac1d52210ef3db73553b46fd0a12c1285","observation_id":"0ca3de16-dcf5-4eba-9bfa-cba00155576b","resolution":{"observed_at":"2026-08-11T01:04:04.059394Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2412.18989","last_updated":"2025-01-18T20:14:21Z","latest_version":2,"primary_category":"cs.SE","snapshot_observed_at":"2026-08-15T19:25:12.024537Z","submitted_at":"2024-12-25T21:56:35Z","title":"How Propense Are Large Language Models at Producing Code Smells? A Benchmarking Study"},"reference_resolution":{"displayed":45,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":9,"verified_exact":4,"verified_fuzzy":32},"total_outbound_references":45},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 45 of 45 outbound references and 4 inbound Pith citation observations for arXiv:2412.18989."}