FazBrowse GitHub Viewer
|
Trending
|
URL:
|
Home
Tools:
[Download Repo ZIP]
[View Raw Code]
[Original HTTPS Page]
python-docs-samples/speech/api/speech_async_rest.py at master · ortgit/python-docs-samples · GitHub
ortgit
/
python-docs-samples
Public
forked from
GoogleCloudPlatform/python-docs-samples
Notifications
You must be signed in to change notification settings
Fork
0
Star
0
Code
Pull requests
0
Actions
Projects
Wiki
Security and quality
0
Insights
Additional navigation options
Code
Pull requests
Actions
Projects
Wiki
Security and quality
Insights
Expand file tree
Breadcrumbs
python-docs-samples
/
speech
/
api
/
speech_async_rest.py
Copy path
More file actions
More file actions
Latest commit
History
History
History
101 lines (84 loc) · 3.22 KB
Breadcrumbs
python-docs-samples
/
speech
/
api
/
speech_async_rest.py
Copy path
File metadata and controls
101 lines (84 loc) · 3.22 KB
Raw
Copy raw file
Download raw file
Open symbols panel
Edit and raw actions
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
75
76
77
78
79
80
81
82
83
84
85
86
87
88
89
90
91
92
93
94
95
96
97
98
99
100
101
#!/usr/bin/env python
# Copyright 2016 Google Inc. All Rights Reserved.
#
# Licensed under the Apache License, Version 2.0 (the "License");
# you may not use this file except in compliance with the License.
# You may obtain a copy of the License at
#
# http://www.apache.org/licenses/LICENSE-2.0
#
# Unless required by applicable law or agreed to in writing, software
# distributed under the License is distributed on an "AS IS" BASIS,
# WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
# See the License for the specific language governing permissions and
# limitations under the License.
"""Google Cloud Speech API sample application using the REST API for async
batch processing."""
# [START import_libraries]
import
argparse
import
base64
import
json
import
time
from
googleapiclient
import
discovery
import
httplib2
from
oauth2client
.
client
import
GoogleCredentials
# [END import_libraries]
# [START authenticating]
# Application default credentials provided by env variable
# GOOGLE_APPLICATION_CREDENTIALS
def
get_speech_service
():
credentials
=
GoogleCredentials
.
get_application_default
().
create_scoped
(
[
'https://www.googleapis.com/auth/cloud-platform'
])
http
=
httplib2
.
Http
()
credentials
.
authorize
(
http
)
return
discovery
.
build
(
'speech'
,
'v1beta1'
,
http
=
http
)
# [END authenticating]
def
main
(
speech_file
):
"""Transcribe the given audio file asynchronously.
Args:
speech_file: the name of the audio file.
"""
# [START construct_request]
with
open
(
speech_file
,
'rb'
)
as
speech
:
# Base64 encode the binary audio file for inclusion in the request.
speech_content
=
base64
.
b64encode
(
speech
.
read
())
service
=
get_speech_service
()
service_request
=
service
.
speech
().
asyncrecognize
(
body
=
{
'config'
: {
# There are a bunch of config options you can specify. See
# https://goo.gl/EPjAup for the full list.
'encoding'
:
'LINEAR16'
,
# raw 16-bit signed LE samples
'sampleRate'
:
16000
,
# 16 khz
# See https://goo.gl/DPeVFW for a list of supported languages.
'languageCode'
:
'en-US'
,
# a BCP-47 language tag
},
'audio'
: {
'content'
:
speech_content
.
decode
(
'UTF-8'
)
}
})
# [END construct_request]
# [START send_request]
response
=
service_request
.
execute
()
print
(
json
.
dumps
(
response
))
# [END send_request]
name
=
response
[
'name'
]
# Construct a GetOperation request.
service_request
=
service
.
operations
().
get
(
name
=
name
)
while
True
:
# Give the server a few seconds to process.
print
(
'Waiting for server processing...'
)
time
.
sleep
(
1
)
# Get the long running operation with response.
response
=
service_request
.
execute
()
if
'done'
in
response
and
response
[
'done'
]:
break
print
(
json
.
dumps
(
response
[
'response'
][
'results'
]))
# [START run_application]
if
__name__
==
'__main__'
:
parser
=
argparse
.
ArgumentParser
()
parser
.
add_argument
(
'speech_file'
,
help
=
'Full path of audio file to be recognized'
)
args
=
parser
.
parse_args
()
main
(
args
.
speech_file
)
# [END run_application]
Back
|
FazBrowse Home
|
New Git URL