annotate naive_output.r @ 4:5ffd52fc35c4 draft

Uploaded
author davidvanzessen
date Mon, 12 Dec 2016 05:22:37 -0500
parents
children
Ignore whitespace changes - Everywhere: Within whitespace: At end of lines:
rev   line source
4
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
1 args <- commandArgs(trailingOnly = TRUE)
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
2
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
3 naive.file = args[1]
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
4 shm.file = args[2]
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
5 output.file.ca = args[3]
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
6 output.file.cg = args[4]
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
7 output.file.cm = args[5]
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
8
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
9 naive = read.table(naive.file, sep="\t", header=T, quote="", fill=T)
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
10 shm.merge = read.table(shm.file, sep="\t", header=T, quote="", fill=T)
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
11
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
12
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
13 final = merge(naive, shm.merge[,c("Sequence.ID", "best_match")], by.x="ID", by.y="Sequence.ID")
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
14 print(paste("nrow final:", nrow(final)))
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
15 names(final)[names(final) == "best_match"] = "Sample"
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
16 final.numeric = final[,sapply(final, is.numeric)]
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
17 final.numeric[is.na(final.numeric)] = 0
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
18 final[,sapply(final, is.numeric)] = final.numeric
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
19
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
20 final.ca = final[grepl("^ca", final$Sample),]
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
21 final.cg = final[grepl("^cg", final$Sample),]
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
22 final.cm = final[grepl("^cm", final$Sample),]
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
23
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
24 if(nrow(final.ca) > 0){
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
25 final.ca$Replicate = 1
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
26 }
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
27
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
28 if(nrow(final.cg) > 0){
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
29 final.cg$Replicate = 1
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
30 }
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
31
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
32 if(nrow(final.cm) > 0){
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
33 final.cm$Replicate = 1
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
34 }
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
35
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
36 #print(paste("nrow final:", nrow(final)))
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
37 #final2 = final
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
38 #final2$Sample = gsub("[0-9]", "", final2$Sample)
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
39 #final = rbind(final, final2)
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
40 #final$Replicate = 1
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
41
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
42 write.table(final.ca, output.file.ca, quote=F, sep="\t", row.names=F, col.names=T)
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
43 write.table(final.cg, output.file.cg, quote=F, sep="\t", row.names=F, col.names=T)
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
44 write.table(final.cm, output.file.cm, quote=F, sep="\t", row.names=F, col.names=T)
5ffd52fc35c4 Uploaded
davidvanzessen
parents:
diff changeset
45