{"as_of":"2026-08-18T20:44:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:ab6b9becd4c7be43c5c0514d34f39bcaf840791ddd8bacff73ccde1c7b9f7330","coverage":[{"denominator":53,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":53,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:35:41.214905Z","state":"measured"},{"denominator":56,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":56,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-26T16:24:25.357338Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"cited_work":{"arxiv_id":"2506.07594","doi":"10.48550/arxiv.2506.07594","metadata_source":"arxiv_reference","pith_arxiv_id":"2506.07594","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Evaluating LLMs Ef- fectiveness in Detecting and Correcting Test Smells","venue":"ArXiv.org","work_id":"ab5a0407-c572-41af-9eee-fb647d01276d","year":2025},"citing_paper":{"arxiv_id":"2604.23361","last_updated":"2026-04-25T16:05:30Z","snapshot_observed_at":"2026-08-13T17:19:57.561446Z","submitted_at":"2026-04-25T16:05:30Z","title":"An Empirical Evaluation of Locally Deployed LLMs for Bug Detection in Python Code","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-08T07:59:25.198736Z"},"links":{"cited_paper":"/paper/2506.07594","citing_paper":"/paper/2604.23361"},"observation_digest":"sha256:317858591aa3bdf218675a390d277b49edad6dac894ad280410001a770b424fa","observation_id":"a8bbc461-08d7-43e7-a15d-eac1263a4d91","resolution":{"observed_at":"2026-05-11T20:46:14.700137Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"cited_work":{"arxiv_id":"2506.07594","doi":"10.48550/arxiv.2506.07594","metadata_source":"arxiv_reference","pith_arxiv_id":"2506.07594","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Evaluating LLMs Ef- fectiveness in Detecting and Correcting Test Smells","venue":"ArXiv.org","work_id":"ab5a0407-c572-41af-9eee-fb647d01276d","year":2025},"citing_paper":{"arxiv_id":"2605.02091","last_updated":"2026-05-03T23:21:13Z","snapshot_observed_at":"2026-08-15T06:30:42.383328Z","submitted_at":"2026-05-03T23:21:13Z","title":"How Compliant Are GitHub Actions Workflows? A Checklist-Based Study with LLM-Assisted Auditing","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-08T19:19:18.667967Z"},"links":{"cited_paper":"/paper/2506.07594","citing_paper":"/paper/2605.02091"},"observation_digest":"sha256:c46811bfc9126f3b4ba8efebd12c3169af8d4a6dbd3cdcdfddd6a619ca8432d9","observation_id":"8902ec67-7c8b-4674-894b-3fbf0d010e3c","resolution":{"observed_at":"2026-05-09T05:55:31.517733Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"cited_work":{"arxiv_id":"2506.07594","doi":"10.48550/arxiv.2506.07594","metadata_source":"arxiv_reference","pith_arxiv_id":"2506.07594","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Evaluating LLMs Ef- fectiveness in Detecting and Correcting Test Smells","venue":"ArXiv.org","work_id":"ab5a0407-c572-41af-9eee-fb647d01276d","year":2025},"citing_paper":{"arxiv_id":"2606.20173","last_updated":"2026-06-18T12:40:43Z","snapshot_observed_at":"2026-08-16T09:41:41.677616Z","submitted_at":"2026-06-18T12:40:43Z","title":"Qiskit Code Migration with LLMs","version":1},"reference_index":155,"source":"arxiv_source","source_observed_at":"2026-06-26T16:24:25.357338Z"},"links":{"cited_paper":"/paper/2506.07594","citing_paper":"/paper/2606.20173"},"observation_digest":"sha256:0ad1e7dccfacbb9c8b1776219dbc838f1f241d0ef14ff7a748e763fc47e6d03b","observation_id":"1fdc2dcf-d38d-430e-9dd7-f7c350880795","resolution":{"observed_at":"2026-06-26T16:29:35.581303Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.07594/citation-record","integrity":"/paper/2506.07594/integrity","json":"/paper/2506.07594/citation-record.json","paper":"/paper/2506.07594"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.751923Z","title":"On the relation of test smells to software code quality,","venue":null,"work_id":"acbf8be4-3b65-4e0e-ad41-63910e2d6e3e","year":2018},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.943703Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:810a40406cc148b36606f534ab4a0bc4f0b48cb4348c3f55c911045225473333","observation_id":"23247134-ec7a-462b-b92b-b0da15036b50","resolution":{"observed_at":"2026-08-07T05:35:42.756482Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"7010.28970","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.275041Z","title":"On the diffusion of test smells in automatically generated test code: An empirical study,","venue":null,"work_id":"2aa83e0a-4201-424d-b9c9-0e9252b2d39e","year":2016},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.949345Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:6a2b68e8bdcfd70ad9ea06133d704ea2a3bab8ab724f309ccc735d7ce249863a","observation_id":"8a415b7c-ea53-4703-807f-b2a336d20b4f","resolution":{"observed_at":"2026-08-07T05:35:42.283784Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.736861Z","title":"When and why your code starts to smell bad,","venue":null,"work_id":"b14253c2-59e2-4acd-a015-ccb05293e184","year":2015},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.954889Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:2b2cf684b34e42e226a341728444bf42fa8a043efde424ba6f28629b2440618c","observation_id":"db152661-eaa6-4851-905d-3fa62b9f4660","resolution":{"observed_at":"2026-08-07T05:35:42.741730Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.721327Z","title":"An empirical analysis of the distribution of unit test smells and their impact on software maintenance,","venue":null,"work_id":"66f0efef-8446-464f-abd6-f5c58cfba008","year":2012},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.960193Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:fff6feb0f69ded25f5b1e26511894f688c0ea5569045db2804fac480c4c8eeab","observation_id":"688e61d6-4cdb-49a5-a2de-f687eabf87dc","resolution":{"observed_at":"2026-08-07T05:35:42.726187Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:40.966381Z","title":"Just-in-time test smell detection and refactoring: The darts project,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.966381Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:031a59a9551d87628bf8b0e0e23e9b040e374bb66a92606c279a8b3bb81f1aae","observation_id":"568dc405-c62c-401c-b339-7d00346a4871","resolution":{"observed_at":"2026-08-07T05:35:40.966381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-17T09:58:46.058102Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-07T05:35:40.971440Z","title":"Gpt-4 technical report,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.971440Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:35a85afcaa46ff5aaadaf2b364b275bc39370d0a6d490c4d000ea9437c839e28","observation_id":"1d6ec493-2c87-4666-9895-45d04d52dc64","resolution":{"observed_at":"2026-08-07T05:35:40.971440Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-07T05:35:40.977343Z","title":"The llama 3 herd of models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.977343Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:d7dc5d7579f3652e28456ea1a9647fef60ba425e17905798c709298cfbaee8d4","observation_id":"748b762e-2e35-4b29-a32f-bc2578dbe139","resolution":{"observed_at":"2026-08-07T05:35:40.977343Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05530","last_updated":"2024-12-16T17:39:39Z","snapshot_observed_at":"2026-08-14T18:15:53.516440Z","submitted_at":"2024-03-08T18:54:20Z","title":"Gemini 1.5: Unlocking multimodal understanding across millions of tokens of context","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.05530","snapshot_observed_at":"2026-08-07T05:35:40.982398Z","title":"Gemini 1.5: Unlocking multimodal understanding across millions of tokens of context,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.982398Z"},"links":{"cited_paper":"/paper/2403.05530","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:045b8c154620388fb0c19178d3ff22731300a4b04e5fc2af8b8e72c98d4204f1","observation_id":"aea9bf46-fe6a-4274-9c69-d0eb6c0509fe","resolution":{"observed_at":"2026-08-07T05:35:40.982398Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.704367Z","title":"Codebert: A pre-trained model for programming and natural languages,","venue":null,"work_id":"7b7b2b4d-fde3-48fb-8405-d5fa2f0327cb","year":2020},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.987859Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:c8f23bf2cf737634b56ddd5a864f34d85da853666ec900cbc6224a232709a81c","observation_id":"08a1ec8d-69e3-4d06-b149-1248e6a9c634","resolution":{"observed_at":"2026-08-07T05:35:42.710348Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.688004Z","title":"Codexglue: A machine learning benchmark dataset for code understanding and generation,","venue":null,"work_id":"ebe78fcc-f5b1-46f9-a3fc-c4e59b8d41b6","year":2021},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.992839Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:7a54acfeb5bb5bbe0f6754e48c0f677540679e16661f3636ca7162f19bdf7dd0","observation_id":"37b2e45a-2e62-4589-a29e-30195993cc57","resolution":{"observed_at":"2026-08-07T05:35:42.692702Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.672772Z","title":"Top programming languages - the state of the octoverse 2022,","venue":null,"work_id":"39904449-2cf2-4b09-be7c-13367e13c9e7","year":2022},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:40.998380Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:046e489e2f923a82089607f7ca94609cc452895e9742fcc9121ab03e3c08714c","observation_id":"84c2c9f9-5d92-4b0f-858b-61b0e668b247","resolution":{"observed_at":"2026-08-07T05:35:42.677548Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.656396Z","title":"Utilization of pre-trained language model for adapter-based knowledge transfer in software engineering,","venue":null,"work_id":"58b8343e-c181-49c1-99ce-fb2978cae1ed","year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.003516Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:cd25a6f9fb00f892e6381da7af95421893eac8f57708107560d0808857805e5c","observation_id":"fe91d1d0-56e7-4e29-8824-7615b7467435","resolution":{"observed_at":"2026-08-07T05:35:42.660795Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.640561Z","title":"To- wards efficient fine-tuning of pre-trained code models: An experimental study and beyond,","venue":null,"work_id":"bcaaa806-8019-44c4-9294-e664ed36da74","year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.008583Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:adf9f9fddad7cb74df6121a86bd7202efd43a157587931ce241daa4744011b20","observation_id":"11ba22fa-d781-4eb2-b61f-2334d7482892","resolution":{"observed_at":"2026-08-07T05:35:42.645541Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.013423Z","title":"An empirical comparison of pre-trained models of source code,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.013423Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:7bd4172c20ba99972ce900f8fe93e6fcddbdd0eed2b21d740b75acf747d27e7c","observation_id":"d2584707-6a9a-40f0-b8f8-2338e8f740ac","resolution":{"observed_at":"2026-08-07T05:35:41.013423Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.624026Z","title":"(2024) Testsmellsrefactoringbyllms","venue":null,"work_id":"57530050-2fe4-47a8-bad4-dcf79f18c3bf","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.018130Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:515b0f3271e659ff76b4fbae4ed918a267d5883d4b57339154b3bd8253774288","observation_id":"ab33be99-7a39-4325-9b52-4b05844df9d9","resolution":{"observed_at":"2026-08-07T05:35:42.629396Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.023442Z","title":"Large language models for software engineering: A systematic literature review,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.023442Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:994c03b28b934a39de07ce92662287fcdfbe51a75f83fc03d89233ee9dfd0c8c","observation_id":"abc5f614-3046-456b-a24d-e7157091aecb","resolution":{"observed_at":"2026-08-07T05:35:41.023442Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.028343Z","title":"Software testing with large language models: Survey, landscape, and vision,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.028343Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:bafc0c97778d8661351314f49d96651f73ff26ff1f10145d8ae486a27ee15363","observation_id":"9007f747-60b9-4a80-9dc6-1896b7c41e46","resolution":{"observed_at":"2026-08-07T05:35:41.028343Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.034055Z","title":"An empirical evaluation of using large language models for automated unit test generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.034055Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:5ac306b8b8bd2f36dd3c6ae4b746d14af01b3c1e42890569bee3848c9da6614f","observation_id":"604f551a-1ac6-4e77-876c-0e1f253f2d38","resolution":{"observed_at":"2026-08-07T05:35:41.034055Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.587596Z","title":"Automated test case repair using language models,","venue":null,"work_id":"3161ac86-05e3-4705-828f-82ad1f67d457","year":2025},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.038870Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:f0cdd3abe0967a2a51439b8b7dba8a9a5a741a461e6c7f4d660c61e9b3a2eb0c","observation_id":"16a0c54f-72b7-452f-ad91-f24e0a7b3b19","resolution":{"observed_at":"2026-08-07T05:35:42.592740Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.572020Z","title":"Chatunitest: a chatgpt- based automated unit test generation tool,","venue":null,"work_id":"ea797eaa-df5b-4bce-98a8-0f3b692c94ed","year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.044486Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:8c0ee15bd79e2cf88964c83dcbf9e93d9e53df0b0b6dcfcfd0e121ee924f472f","observation_id":"a607be27-a5ec-42e0-9f50-e90ee04bff8a","resolution":{"observed_at":"2026-08-07T05:35:42.577039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.555832Z","title":"An empirical study of using large language models for unit test generation,","venue":null,"work_id":"34856d8c-79f5-48cd-9d6a-4d88fb063fd4","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.049310Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:06031962fc3f2268eb5c4a8ff9695d5201b82af086ba84e476c6407feb809c23","observation_id":"7501baa6-8233-41ba-922e-bc3d59d8bb78","resolution":{"observed_at":"2026-08-07T05:35:42.561428Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.539874Z","title":"Towards an understanding of large language models in software engineering tasks,","venue":null,"work_id":"bc1ebd87-a344-4a92-96f5-3361af2684c2","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.055046Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:1b9ef4a2db47182019d3e2aae2dd5bacf015dd5b5dd3175688b102e592857988","observation_id":"d4444700-dfc5-48e5-8936-21cab8d709c6","resolution":{"observed_at":"2026-08-07T05:35:42.545326Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.060347Z","title":"Pynose: a test smell detector for python,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.060347Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:01434753dcf80de95651a5d2c1f8165262f653f4c4c0d1c6b521fcd879fd89ac","observation_id":"01643684-ce9e-417e-b4ce-ee1b3ecd2056","resolution":{"observed_at":"2026-08-07T05:35:41.060347Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.065639Z","title":"Tempy: Test smell detector for python,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.065639Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:114f7275a2ef9ad132ee10c41f395d024decb9979065478c09ef802b1b15def8","observation_id":"37e1047a-9964-4933-ad72-31fc5fd62b2c","resolution":{"observed_at":"2026-08-07T05:35:41.065639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"4624.34770","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.823854Z","title":"Handling test smells in python: Results from a mixed-method study,","venue":null,"work_id":"2216b61e-08e8-4190-9f80-2c0f4fff5b63","year":2021},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.070373Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:5a2b9ece6a27f6dc8ffd00d18387849cd646125d0900992531c3ae9a5308a1ce","observation_id":"1808a497-898a-4853-89d7-46f6bb0b93c4","resolution":{"observed_at":"2026-08-07T05:35:41.834355Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.519336Z","title":"A trend analysis of test smells in python test code over commit history,","venue":null,"work_id":"b624f8ca-1bd5-42df-9381-7c422513291e","year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.075592Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:d45a279c0cf056c8957d6f97513f70947ed616031cff052bb2707c44d04f8b81","observation_id":"0b467529-eab7-4e27-910f-cd8f9cd3ffd2","resolution":{"observed_at":"2026-08-07T05:35:42.525234Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.080279Z","title":"Pytest-smell: A smell detection tool for python unit tests,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.080279Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:0a69a54c8d9a79f7febb74a8dcf0a56877c24d181959cbb90b05ce26253876e5","observation_id":"a7830eb5-2748-452a-9657-0adedbb77461","resolution":{"observed_at":"2026-08-07T05:35:41.080279Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.085461Z","title":"A trend analysis of test smells in python test code over commit history,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.085461Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:33c287ba5d081ca64be471b8dc9036c27bd58053a7799abe5deb809d8fef471a","observation_id":"15ad84c6-b314-495e-a34f-c35c84b732d8","resolution":{"observed_at":"2026-08-07T05:35:41.085461Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.090496Z","title":"Detecting test smells in python test code generated by LLM: an empirical study with github copilot,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.090496Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:8f6e532bb9ba8b0cbc83b9a1f9d803a5f40e2361b443050154ec12659a77ed88","observation_id":"6e9b218f-0a68-451f-b6e1-53b2cb1d204a","resolution":{"observed_at":"2026-08-07T05:35:41.090496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.095249Z","title":"Tsdetect: An open source test smells detection tool,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.095249Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:a171c7e52f88b7407641262e493e4f2b65b625e736a6bcb517b8cdb5c7dcc046","observation_id":"27036ac3-7f44-4dc0-90e4-a0c565ff228e","resolution":{"observed_at":"2026-08-07T05:35:41.095249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.502190Z","title":"The secret life of test smells - an empirical study on test smell evolution and maintenance,","venue":null,"work_id":"d2b81647-354b-42c8-94e7-4a523603a94e","year":2021},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.100530Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:7432517888b4c0c8c55e08e91466df2a1cdad446e0f8d74df0920d4ca94794d4","observation_id":"f37741b2-b5b4-4428-a81e-8f656088678b","resolution":{"observed_at":"2026-08-07T05:35:42.507134Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.485015Z","title":"An empirical investigation into the nature of test smells,","venue":null,"work_id":"32ddda6c-2fac-462c-92f5-a8c058793bb2","year":2016},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.105536Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:9f3910d072458bdecca159209d148210cd9a52e6765787490a1a932573631c97","observation_id":"c5a472aa-b24b-487a-893f-dded75dc24ff","resolution":{"observed_at":"2026-08-07T05:35:42.490301Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.468208Z","title":"An empirical evaluation of raide: A semi-automated approach for test smells detection and refactoring,","venue":null,"work_id":"eb90a6ba-daeb-4bfe-9d9e-1837a2825320","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.110958Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:cd91df11465bcdc68f40d5fda6740f7f868ec9a026b1313942bcfd15d9bca255","observation_id":"55a94a4e-6461-4fba-9507-0c1234c75a4d","resolution":{"observed_at":"2026-08-07T05:35:42.473239Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.452324Z","title":"Machine learning-based test smell detection,","venue":null,"work_id":"7e3674d8-7773-4dbf-b52f-57c1051cff36","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.116048Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:add1a3a080c505fbdc5787fb34ec86c2f061e71f8fedba62d2504a153b2148e0","observation_id":"4c6232e1-13ef-464f-9272-7c5832301035","resolution":{"observed_at":"2026-08-07T05:35:42.457599Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.436335Z","title":"Ml test smell detection - online appendix,","venue":null,"work_id":"0e8df838-81f2-401c-b4a4-2e383fcf1a2d","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.121328Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:e2ef61b31d99201f7027c92a9307523a1286b93f683ac709e604cfcd9b21304f","observation_id":"17333b22-3370-4470-ada3-de15e04eaf6e","resolution":{"observed_at":"2026-08-07T05:35:42.441334Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06608","last_updated":"2025-02-26T18:59:01Z","snapshot_observed_at":"2026-08-13T23:11:34.204909Z","submitted_at":"2024-06-06T18:10:11Z","title":"The Prompt Report: A Systematic Survey of Prompt Engineering Techniques","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06608","snapshot_observed_at":"2026-08-07T05:35:41.127695Z","title":"The prompt report: A systematic survey of prompt engineering techniques,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.127695Z"},"links":{"cited_paper":"/paper/2406.06608","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:7d2489d3fe63df49e49185a06da042603dba6b9c53e331e46b1c17acfadb2244","observation_id":"979a075e-875f-4c4e-817a-651e9631ee66","resolution":{"observed_at":"2026-08-07T05:35:41.127695Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.11382","last_updated":"2023-02-21T12:42:44Z","snapshot_observed_at":"2026-07-06T14:54:37.559648Z","submitted_at":"2023-02-21T12:42:44Z","title":"A Prompt Pattern Catalog to Enhance Prompt Engineering with ChatGPT","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.11382","snapshot_observed_at":"2026-08-07T05:35:41.133174Z","title":"A prompt pattern catalog to enhance prompt engineering with chatgpt,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.133174Z"},"links":{"cited_paper":"/paper/2302.11382","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:35a4740abb7bd680f880889f3949f84fe5fee4aec46fed7987d637f0b5354b2b","observation_id":"9e1f8995-92c8-409a-9f5e-50e8760143c9","resolution":{"observed_at":"2026-08-07T05:35:41.133174Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2109.01652","last_updated":"2022-02-08T20:26:45Z","snapshot_observed_at":"2026-08-14T06:11:14.515796Z","submitted_at":"2021-09-03T17:55:52Z","title":"Finetuned Language Models Are Zero-Shot Learners","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.01652","snapshot_observed_at":"2026-08-07T05:35:41.138681Z","title":"Finetuned language models are zero-shot learners,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.138681Z"},"links":{"cited_paper":"/paper/2109.01652","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:61d5ff46d9105702e42d559ae3ff3ed294027319a257d6a2260651af5d4771a0","observation_id":"e92c0d03-5ce9-4438-8ffd-bad42eb089bf","resolution":{"observed_at":"2026-08-07T05:35:41.138681Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.143500Z","title":"Language models are few-shot learners,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.143500Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:3f769330b785c44b67c8c56eaf2731055da662636e3b45f1bef4decfb81937f6","observation_id":"bcf2c254-a14a-4662-9088-e3d66ed00209","resolution":{"observed_at":"2026-08-07T05:35:41.143500Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2201.11903","last_updated":"2023-01-10T23:07:57Z","snapshot_observed_at":"2026-08-13T07:04:41.220509Z","submitted_at":"2022-01-28T02:33:07Z","title":"Chain-of-Thought Prompting Elicits Reasoning in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2201.11903","snapshot_observed_at":"2026-08-07T05:35:41.154116Z","title":"Chain-of-thought prompting elicits reasoning in large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.154116Z"},"links":{"cited_paper":"/paper/2201.11903","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:6b1bb9d68e40cf8e6a310dcf4fbe326fbdc114c3492ba0eeb0475eb3f3f0ce4b","observation_id":"e8bcc4d9-42d0-42c7-bd7c-674bfb4baf9a","resolution":{"observed_at":"2026-08-07T05:35:41.154116Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.409666Z","title":"Enhancing zero-shot chain-of-thought reasoning in large language models through logic,","venue":null,"work_id":"dacbe487-1248-4864-aab7-5bcdb43d2e6a","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.159476Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:b8c268410b49ff57b95e319f27cfe7254d01bea5854a26c30dd9b266feae3e3e","observation_id":"f6690eac-0686-4d46-8b51-481bc0203440","resolution":{"observed_at":"2026-08-07T05:35:42.414785Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:41.164376Z","title":"Wilcoxon, Individual Comparisons by Ranking Methods","venue":null,"work_id":null,"year":1992},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.164376Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:57fff426e6a8e49ecad39f2adccbc9b752e11552b30b700a34fa57a3269662c6","observation_id":"c01ba966-08cf-47cb-af39-6179098f4152","resolution":{"observed_at":"2026-08-07T05:35:41.164376Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14261","last_updated":"2024-02-22T03:51:34Z","snapshot_observed_at":"2026-08-16T14:16:23.227477Z","submitted_at":"2024-02-22T03:51:34Z","title":"Copilot Evaluation Harness: Evaluating LLM-Guided Software Programming","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14261","snapshot_observed_at":"2026-08-07T05:35:41.169640Z","title":"Copilot evaluation harness: Evaluating llm-guided software programming,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.169640Z"},"links":{"cited_paper":"/paper/2402.14261","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:3547c2fb3480bf289d8fb06ab32835a5d325ba136554015c7334a18cabf69dcf","observation_id":"4b784f3d-ec10-4cef-ab7d-6c9848bebb76","resolution":{"observed_at":"2026-08-07T05:35:41.169640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.394234Z","title":"Towards effective validation and integration of llm-generated code,","venue":null,"work_id":"57ab2c31-5f7f-40d8-b7be-2324e9b730b3","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.174703Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:cfb01aaf0ae7612266b1a26ad2e030fcd5430b87ed3962547b5aee93f995c809","observation_id":"203e52a2-3eb4-4dd5-93e2-9ec74aabd054","resolution":{"observed_at":"2026-08-07T05:35:42.399120Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.378728Z","title":"Challenges and opportunities in integrating llms into con- tinuous integration/continuous deployment (ci/cd) pipelines,","venue":null,"work_id":"f3a55412-31fc-4066-bc74-1cff34b3e926","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.179595Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:ff83554c1c32d6f3bb38a53e6157c951739fb316db98d4ad2a6e2baa1a8da111","observation_id":"4d3c0078-843f-4abe-a90e-af4cd3796809","resolution":{"observed_at":"2026-08-07T05:35:42.383731Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.362357Z","title":"Next-generation refactoring: Combining llm insights and ide capabilities for extract method,","venue":null,"work_id":"16f7f52f-a69a-415b-9bd0-17dbc3165d21","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.184895Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:daca01f877200c52b237de6a84b3e6ebb93693d4493259f68a1a65d6ad770271","observation_id":"3866c2cb-de0c-4740-9325-25b67ecdb778","resolution":{"observed_at":"2026-08-07T05:35:42.367755Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.345946Z","title":"Llm-based multi-agent systems for software engineering: Literature review, vision and the road ahead,","venue":null,"work_id":"2bc82942-eb68-4592-a381-c780b41e26ad","year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.189548Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:9623dac20879d65374930d5b9d90fc2f92615d4f5802d962c31349549ba700b7","observation_id":"bcf1d462-afeb-49f6-8b97-81e8497ba694","resolution":{"observed_at":"2026-08-07T05:35:42.350942Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.329809Z","title":"Autorefactoring: A platform to build refactoring agents,","venue":null,"work_id":"c1542b1b-8ea6-4cea-af0a-9ab9eded7205","year":2015},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.194572Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:33f3086b5bca5b56d7b6227412390fbc74e468c8ef17afce241e64182628cf5c","observation_id":"ba98147c-6e57-4c7a-86bd-39625fad115e","resolution":{"observed_at":"2026-08-07T05:35:42.335039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.19437","last_updated":"2025-02-18T17:26:38Z","snapshot_observed_at":"2026-08-18T18:18:37.449517Z","submitted_at":"2024-12-27T04:03:16Z","title":"DeepSeek-V3 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.19437","snapshot_observed_at":"2026-08-07T05:35:41.200114Z","title":"Deepseek-v3 technical report,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.200114Z"},"links":{"cited_paper":"/paper/2412.19437","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:6e6cf409c25c3d8470317858c8d9948725fe54e2ad7d7037f8917a2149c5df5a","observation_id":"de96ca94-7c8f-43c9-8533-e95f24ffaad4","resolution":{"observed_at":"2026-08-07T05:35:41.200114Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-08-15T12:33:55.451951Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-07T05:35:41.204978Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.204978Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:6d3fd8d90fb45b5ec1d1bef7d71c9c001958da189f67e9167d135ab7b4635f01","observation_id":"d33ee6e0-9751-4ba5-9b5b-b0d791ab5ddc","resolution":{"observed_at":"2026-08-07T05:35:41.204978Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.313498Z","title":"Runeson, M","venue":null,"work_id":"8f6c9e2c-8a72-4f4b-94be-4e6d5daa2d1b","year":2012},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.210159Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:41aa9f99af126b401b2a403097c31ec4b860c92abe3b2ecd088bc554855af3b6","observation_id":"d9736758-fc93-45ee-819c-908bea04db40","resolution":{"observed_at":"2026-08-07T05:35:42.318400Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:35:42.297332Z","title":"Qualitative methods in empirical studies of software engineering,","venue":null,"work_id":"02caa27b-ba13-43bb-8a1a-f35cee23b65c","year":1999},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.214905Z"},"links":{"citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:89335ba8f9fcf5e8a3b63a67a318b345acd94e8d96f93260d0901659da575b40","observation_id":"9a477ee6-450c-4e12-8886-ead2dd72cacc","resolution":{"observed_at":"2026-08-07T05:35:42.302511Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2005.14165","last_updated":"2020-07-22T19:47:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-05-28T17:29:03Z","title":"Language Models are Few-Shot Learners","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.14165","snapshot_observed_at":"2026-08-07T05:35:41.149159Z","title":"Available: https://arxiv.org/abs/2005.14165","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-07T05:35:41.149159Z"},"links":{"cited_paper":"/paper/2005.14165","citing_paper":"/paper/2506.07594"},"observation_digest":"sha256:c7e701ca575a699bcfa761f7089fdb99f9e77eaaafe628f1358350bfb9b62bf3","observation_id":"a9133213-2008-441f-ab73-53fa2c8f35fb","resolution":{"observed_at":"2026-08-07T05:35:41.149159Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.07594","last_updated":"2025-06-09T09:46:41Z","latest_version":1,"primary_category":"cs.SE","snapshot_observed_at":"2026-08-11T13:20:33.453818Z","submitted_at":"2025-06-09T09:46:41Z","title":"Evaluating LLMs Effectiveness in Detecting and Correcting Test Smells: An Empirical Study"},"reference_resolution":{"displayed":53,"state_counts":{"malformed_identifier":0,"metadata_mismatch":2,"parse_uncertain":0,"unresolved":24,"verified_exact":0,"verified_fuzzy":27},"total_outbound_references":53},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 53 of 53 outbound references and 3 inbound Pith citation observations for arXiv:2506.07594."}