FazBrowse GitHub Viewer
|
Trending
|
URL:
|
Home
Tools:
[Download Repo ZIP]
[View Raw Code]
[Original HTTPS Page]
python-mini-projects/projects/Split_File/split_files.py at master · Coder98321/python-mini-projects · GitHub
Coder98321
/
python-mini-projects
Public
forked from
Python-World/python-mini-projects
Notifications
You must be signed in to change notification settings
Fork
0
Star
0
Code
Pull requests
0
Actions
Projects
Security and quality
0
Insights
Additional navigation options
Code
Pull requests
Actions
Projects
Security and quality
Insights
Expand file tree
Breadcrumbs
python-mini-projects
/
projects
/
Split_File
/
split_files.py
Copy path
More file actions
More file actions
Latest commit
History
History
History
55 lines (50 loc) · 1.98 KB
Breadcrumbs
python-mini-projects
/
projects
/
Split_File
/
split_files.py
Copy path
File metadata and controls
55 lines (50 loc) · 1.98 KB
Raw
Copy raw file
Download raw file
Open symbols panel
Edit and raw actions
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
import
sys
import
os
import
shutil
import
pandas
as
pd
class
Split_Files
:
'''
Class file for split file program
'''
def
__init__
(
self
,
filename
,
split_number
):
'''
Getting the file name and the split index
Initializing the output directory, if present then truncate it.
Getting the file extension
'''
self
.
file_name
=
filename
self
.
directory
=
"file_split"
self
.
split
=
int
(
split_number
)
if
os
.
path
.
exists
(
self
.
directory
):
shutil
.
rmtree
(
self
.
directory
)
os
.
mkdir
(
self
.
directory
)
if
self
.
file_name
.
endswith
(
'.txt'
):
self
.
file_extension
=
'.txt'
else
:
self
.
file_extension
=
'.csv'
self
.
file_number
=
1
def
split_data
(
self
):
'''
spliting the input csv/txt file according to the index provided
'''
data
=
pd
.
read_csv
(
self
.
file_name
,
header
=
None
)
data
.
index
+=
1
split_frame
=
pd
.
DataFrame
()
output_file
=
f"
{
self
.
directory
}
/split_file
{
self
.
file_number
}
{
self
.
file_extension
}
"
for
i
in
range
(
1
,
len
(
data
)
+
1
):
split_frame
=
split_frame
.
append
(
data
.
iloc
[
i
-
1
])
if
i
%
self
.
split
==
0
:
output_file
=
f"
{
self
.
directory
}
/split_file
{
self
.
file_number
}
{
self
.
file_extension
}
"
if
self
.
file_extension
==
'.txt'
:
split_frame
.
to_csv
(
output_file
,
header
=
False
,
index
=
False
,
sep
=
' '
)
else
:
split_frame
.
to_csv
(
output_file
,
header
=
False
,
index
=
False
)
split_frame
.
drop
(
split_frame
.
index
,
inplace
=
True
)
self
.
file_number
+=
1
if
not
split_frame
.
empty
:
output_file
=
f"
{
self
.
directory
}
/split_file
{
self
.
file_number
}
{
self
.
file_extension
}
"
split_frame
.
to_csv
(
output_file
,
header
=
False
,
index
=
False
)
if
__name__
==
'__main__'
:
file
,
split_number
=
sys
.
argv
[
1
],
sys
.
argv
[
2
]
sp
=
Split_Files
(
file
,
split_number
)
sp
.
split_data
()
Back
|
FazBrowse Home
|
New Git URL