@inbook{08ce13e354854a458505f94921bbf19a,
title = "Comparative Analysis of Clustering Methodologies in DNA Storage",
abstract = "Owing to the significance of DNA storage technology in meeting exponential storage demands and longevity, the challenges caused by bio-molecular errors while reading/sequencing data from DNA molecules must be addressed. By reading redundant copies, data can be reconstructed but with associated cost of sequencing and decoding complexities. Hence, solutions for dealing with both errors and complexities are sought after. The main objective of this work is to study data reconstruction methods for processing sequence readouts at downstream stage of DNA data storage. We investigated applicability of three clustering tools -Starcode, Slidesort, MeShClust, and two algorithms - Majority Nucleotide Selection (MNS), Cooperative Sequence Clustering (CSC) by transforming them into suitable tools for storage application. We observed that for fixed redundancy of 6.3x to 8.6x based on the nature of the dataset, Starcode outperforms other tools with 1\% to 40\% higher recovery rate. However, it costs the highest decoding complexity whereas MNS and CSC provides the lowest decoding complexity. Moreover, the distribution of the cluster and clustering speed of each tool/method are compared. This is the first comparative analysis study of tools/methods for data reconstruction in DNA data storage.",
keywords = "Clustering, Data reconstruction, DNA data storage, Illumina sequencing",
author = "Subhasiny Sankar and Yixin Wang and Zhang Jiayu and Nur Sabrina and Erry Gunawan and Guan, \{Yong Liang\} and Md, \{Noor A.Rahim\} and Poh, \{Chueh Loo\}",
note = "Publisher Copyright: {\textcopyright} 2022 IEEE.; 26th International Computer Science and Engineering Conference, ICSEC 2022 ; Conference date: 21-12-2022 Through 23-12-2022",
year = "2022",
doi = "10.1109/ICSEC56337.2022.10049327",
language = "English",
series = "ICSEC 2022 - International Computer Science and Engineering Conference 2022",
publisher = "Institute of Electrical and Electronics Engineers Inc.",
pages = "269--274",
booktitle = "ICSEC 2022 - International Computer Science and Engineering Conference 2022",
address = "United States",
}