@inproceedings{raghunandan2023code, author = {Raghunandan, Deepthi and Roy, Aayushi and Shi, Shenzhi and Elmqvist, Niklas and Battle, Leilani}, title = {Code Code Evolution: Understanding How People Change Data Science Notebooks Over Time}, year = {2023}, isbn = {9781450394215}, publisher = {Association for Computing Machinery}, address = {New York, NY, USA}, url = {https://doi.org/10.1145/3544548.3580997}, doi = {10.1145/3544548.3580997}, abstract = {Sensemaking is the iterative process of identifying, extracting, and explaining insights from data, where each iteration is referred to as the “sensemaking loop.” However, little is known about how sensemaking behavior evolves from exploration and explanation during this process. This gap limits our ability to understand the full scope of sensemaking, which in turn inhibits the design of tools that support the process. We contribute the first mixed-method to characterize how sensemaking evolves within computational notebooks. We study 2,574 Jupyter notebooks mined from GitHub by identifying data science notebooks that have undergone significant iterations, presenting a regression model that automatically characterizes sensemaking activity, and using this regression model to calculate and analyze shifts in activity across GitHub versions. Our results show that notebook authors participate in various sensemaking tasks over time, such as annotation, branching analysis, and documentation. We use our insights to recommend extensions to current notebook environments.}, booktitle = {Proceedings of the 2023 CHI Conference on Human Factors in Computing Systems}, articleno = {863}, numpages = {12}, keywords = {data exploration, sensemaking, data science, machine learning, analysis., Computational notebooks}, location = {Hamburg, Germany}, series = {CHI '23} }