view PDAUG_Peptide_Length_Distribution/PDAUG_Peptide_Length_Distribution.py @ 3:c8a33b76ce7c draft

"planemo upload for repository https://github.com/jaidevjoshi83/pdaug commit edb37634e419f75dd66292e712de51278746d883"
author jay
date Wed, 30 Dec 2020 03:25:54 +0000
parents a3a1d9bea1ad
children
line wrap: on
line source

import matplotlib.pyplot as plt
import Bio
from Bio import SeqIO
import os


def LegnthDestribution(InFile, OutFile):


    sizes = [len(rec.seq) for rec in SeqIO.parse(InFile, "fasta")]

    plt.hist(sizes, bins=20)
    plt.title("%i Negative bacteriocin sequences\nLengths %i to %i" \
                % (len(sizes),min(sizes),max(sizes)))
    plt.xlabel("Sequence length (bp)")
    plt.ylabel("Count")

    plt.savefig(OutFile)



if __name__=="__main__":

    import argparse
    
    parser = argparse.ArgumentParser()
    
    parser.add_argument("-I", "--InFile", required=True, default=None, help="Input file name")
    parser.add_argument("-O", "--OutFile", required=False, default="Out.png", help="Input file name")
    args = parser.parse_args()
    LegnthDestribution(args.InFile, args.OutFile)