annotate variant_effect_predictor/Bio/Tools/Profile.pm @ 3:d30fa12e4cc5 default tip

Merge heads 2:a5976b2dce6f and 1:09613ce8151e which were created as a result of a recently fixed bug.
author devteam <devteam@galaxyproject.org>
date Mon, 13 Jan 2014 10:38:30 -0500
parents 1f6dce3d34e0
children
Ignore whitespace changes - Everywhere: Within whitespace: At end of lines:
rev   line source
0
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
1 # BioPerl module for Bio::Tools::Profile
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
2 #
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
3 # Cared for by Balamurugan Kumarasamy
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
4 #
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
5 # You may distribute this module under the same terms as perl itself
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
6 # POD documentation - main docs before the code
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
7
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
8 =head1 NAME
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
9
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
10 Bio::Tools::Profile - parse Profile output
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
11
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
12 =head1 SYNOPSIS
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
13
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
14 use Bio::Tools::Profile;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
15 my $profile_parser = new Bio::Tools::Profile(-fh =>$filehandle );
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
16 while( my $profile_feat = $profile_parser->next_result ) {
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
17 push @profile_feat, $profile_feat;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
18 }
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
19
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
20 =head1 DESCRIPTION
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
21
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
22 Parser for Profile output
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
23
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
24 =head1 FEEDBACK
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
25
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
26 =head2 Mailing Lists
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
27
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
28 User feedback is an integral part of the evolution of this and other
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
29 Bioperl modules. Send your comments and suggestions preferably to
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
30 the Bioperl mailing list. Your participation is much appreciated.
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
31
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
32 bioperl-l@bioperl.org - General discussion
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
33 http://bioperl.org/MailList.shtml - About the mailing lists
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
34
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
35 =head2 Reporting Bugs
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
36
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
37 Report bugs to the Bioperl bug tracking system to help us keep track
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
38 of the bugs and their resolution. Bug reports can be submitted via
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
39 email or the web:
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
40
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
41 bioperl-bugs@bioperl.org
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
42 http://bugzilla.bioperl.org/
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
43 =head1 AUTHOR - Balamurugan Kumarasamy
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
44
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
45 Email: fugui@worf.fugu-sg.org
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
46
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
47 =head1 APPENDIX
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
48
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
49 The rest of the documentation details each of the object methods.
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
50 Internal methods are usually preceded with a _
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
51
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
52
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
53 =cut
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
54
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
55
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
56 package Bio::Tools::Profile;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
57 use vars qw(@ISA);
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
58 use strict;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
59
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
60 use Bio::Root::Root;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
61 use Bio::SeqFeature::FeaturePair;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
62 use Bio::Root::IO;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
63 use Bio::SeqFeature::Generic;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
64
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
65 @ISA = qw(Bio::Root::Root Bio::Root::IO );
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
66
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
67
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
68
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
69 =head2 new
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
70
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
71 Title : new
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
72 Usage : my $obj = new Bio::Tools::Profile();
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
73 Function: Builds a new Bio::Tools::Profile object
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
74 Returns : Bio::Tools::Profile
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
75 Args : -filename
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
76 -fh ($filehandle)
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
77
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
78 =cut
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
79
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
80 sub new {
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
81 my($class,@args) = @_;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
82
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
83 my $self = $class->SUPER::new(@args);
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
84 $self->_initialize_io(@args);
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
85
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
86 return $self;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
87 }
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
88
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
89 =head2 next_result
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
90
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
91 Title : next_result
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
92 Usage : my $feat = $profile_parser->next_result
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
93 Function: Get the next result set from parser data
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
94 Returns : L<Bio::SeqFeature::FeaturePair>
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
95 Args : none
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
96
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
97
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
98 =cut
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
99
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
100 sub next_result {
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
101 my ($self) = @_;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
102
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
103 my %printsac;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
104 my $line;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
105 my @features;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
106 while ($_=$self->_readline()) {
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
107 $line = $_;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
108 chomp $line;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
109 my ($nscore,$rawscore,$from,$to,$hfrom,$hto,$ac) = $line =~ /(\S+)\s+(\d+)\s*pos.\s+(\d*)\s*-\s+(\d*)\s*\[\s+(\d*),\s+(\S*)\]\s*(\w+)/;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
110 #for example in this output line
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
111 #38.435 2559 pos. 19958 - 20212 [ 1, -1] PS50011|PROTEIN_KINASE_DOM Protein kinase domain profile.
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
112 #$nscore = 38.435
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
113 #$rawscore = 2559
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
114 #$from = 19958
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
115 #$end = 20212
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
116 #$hfrom = 1
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
117 #$hto =-1
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
118 #$ac = PS50011
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
119 my $feat = "$ac,$from,$to,$hfrom,$hto,$nscore";
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
120 my $new_feat= $self->create_feature($feat);
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
121 return $new_feat
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
122
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
123 }
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
124 }
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
125
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
126
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
127 =head2 create_feature
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
128
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
129 Title : create_feature
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
130 Usage : my $feat= $profile_parser->create_feature($feature)
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
131 Function: creates a Bio::SeqFeature::FeaturePair object
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
132 Returns : L<Bio::SeqFeature::FeaturePair>
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
133 Args :
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
134
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
135
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
136 =cut
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
137
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
138 sub create_feature {
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
139 my ($self, $feat) = @_;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
140
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
141 my @f = split (/,/,$feat);
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
142
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
143
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
144 my $hto = $f[4];
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
145
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
146 if ($f[4] =~ /-1/) {
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
147
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
148 $hto = $f[2] - $f[1] + 1;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
149
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
150 }
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
151
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
152
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
153 my $feat1 = new Bio::SeqFeature::Generic ( -start => $f[1],
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
154 -end => $f[2],
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
155 -score => $f[5],
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
156 -source=>'pfscan',
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
157 -primary=>$f[0]);
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
158
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
159 my $feat2 = new Bio::SeqFeature::Generic (-start => $f[3],
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
160 -end => $hto,
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
161 );
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
162
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
163 my $feature = new Bio::SeqFeature::FeaturePair(-feature1 => $feat1,
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
164 -feature2 => $feat2);
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
165
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
166 return $feature;
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
167
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
168 }
1f6dce3d34e0 Uploaded
mahtabm
parents:
diff changeset
169 1;