2023-07-07 17:50:42 +08:00
|
|
|
import logging
|
|
|
|
|
2024-02-06 13:21:13 +08:00
|
|
|
from flask import request
|
|
|
|
from werkzeug.exceptions import InternalServerError
|
|
|
|
|
2023-07-07 17:50:42 +08:00
|
|
|
import services
|
|
|
|
from controllers.web import api
|
2024-02-06 13:21:13 +08:00
|
|
|
from controllers.web.error import (
|
|
|
|
AppUnavailableError,
|
|
|
|
AudioTooLargeError,
|
|
|
|
CompletionRequestError,
|
|
|
|
NoAudioUploadedError,
|
|
|
|
ProviderModelCurrentlyNotSupportError,
|
|
|
|
ProviderNotInitializeError,
|
|
|
|
ProviderNotSupportSpeechToTextError,
|
|
|
|
ProviderQuotaExceededError,
|
|
|
|
UnsupportedAudioTypeError,
|
|
|
|
)
|
2023-07-07 17:50:42 +08:00
|
|
|
from controllers.web.wraps import WebApiResource
|
2024-01-12 12:34:01 +08:00
|
|
|
from core.errors.error import ModelCurrentlyNotSupportError, ProviderTokenNotInitError, QuotaExceededError
|
2024-01-02 23:42:00 +08:00
|
|
|
from core.model_runtime.errors.invoke import InvokeError
|
2023-07-07 17:50:42 +08:00
|
|
|
from models.model import App, AppModelConfig
|
2024-01-12 12:34:01 +08:00
|
|
|
from services.audio_service import AudioService
|
2024-02-06 13:21:13 +08:00
|
|
|
from services.errors.audio import (
|
|
|
|
AudioTooLargeServiceError,
|
|
|
|
NoAudioUploadedServiceError,
|
|
|
|
ProviderNotSupportSpeechToTextServiceError,
|
|
|
|
UnsupportedAudioTypeServiceError,
|
|
|
|
)
|
2023-07-07 17:50:42 +08:00
|
|
|
|
|
|
|
|
|
|
|
class AudioApi(WebApiResource):
|
|
|
|
def post(self, app_model: App, end_user):
|
|
|
|
app_model_config: AppModelConfig = app_model.app_model_config
|
|
|
|
|
|
|
|
if not app_model_config.speech_to_text_dict['enabled']:
|
|
|
|
raise AppUnavailableError()
|
|
|
|
|
|
|
|
file = request.files['file']
|
|
|
|
|
|
|
|
try:
|
2024-01-24 01:05:37 +08:00
|
|
|
response = AudioService.transcript_asr(
|
2023-07-07 17:50:42 +08:00
|
|
|
tenant_id=app_model.tenant_id,
|
|
|
|
file=file,
|
2024-01-24 23:04:14 +08:00
|
|
|
end_user=end_user
|
2023-07-07 17:50:42 +08:00
|
|
|
)
|
|
|
|
|
|
|
|
return response
|
|
|
|
except services.errors.app_model_config.AppModelConfigBrokenError:
|
|
|
|
logging.exception("App model config broken.")
|
|
|
|
raise AppUnavailableError()
|
|
|
|
except NoAudioUploadedServiceError:
|
|
|
|
raise NoAudioUploadedError()
|
|
|
|
except AudioTooLargeServiceError as e:
|
|
|
|
raise AudioTooLargeError(str(e))
|
|
|
|
except UnsupportedAudioTypeServiceError:
|
|
|
|
raise UnsupportedAudioTypeError()
|
|
|
|
except ProviderNotSupportSpeechToTextServiceError:
|
|
|
|
raise ProviderNotSupportSpeechToTextError()
|
2023-07-17 00:14:19 +08:00
|
|
|
except ProviderTokenNotInitError as ex:
|
|
|
|
raise ProviderNotInitializeError(ex.description)
|
2023-07-07 17:50:42 +08:00
|
|
|
except QuotaExceededError:
|
|
|
|
raise ProviderQuotaExceededError()
|
|
|
|
except ModelCurrentlyNotSupportError:
|
|
|
|
raise ProviderModelCurrentlyNotSupportError()
|
2024-01-02 23:42:00 +08:00
|
|
|
except InvokeError as e:
|
2024-01-04 17:49:55 +08:00
|
|
|
raise CompletionRequestError(e.description)
|
2023-07-07 17:50:42 +08:00
|
|
|
except ValueError as e:
|
|
|
|
raise e
|
|
|
|
except Exception as e:
|
2024-02-15 22:41:18 +08:00
|
|
|
logging.exception(f"internal server error: {str(e)}")
|
2023-07-07 17:50:42 +08:00
|
|
|
raise InternalServerError()
|
|
|
|
|
2024-01-24 01:05:37 +08:00
|
|
|
|
|
|
|
class TextApi(WebApiResource):
|
|
|
|
def post(self, app_model: App, end_user):
|
2024-02-15 22:41:18 +08:00
|
|
|
app_model_config: AppModelConfig = app_model.app_model_config
|
|
|
|
|
|
|
|
if not app_model_config.text_to_speech_dict['enabled']:
|
|
|
|
raise AppUnavailableError()
|
|
|
|
|
2024-01-24 01:05:37 +08:00
|
|
|
try:
|
|
|
|
response = AudioService.transcript_tts(
|
|
|
|
tenant_id=app_model.tenant_id,
|
|
|
|
text=request.form['text'],
|
|
|
|
end_user=end_user.external_user_id,
|
2024-03-04 17:50:06 +08:00
|
|
|
voice=request.form['voice'] if request.form['voice'] else app_model.app_model_config.text_to_speech_dict.get('voice'),
|
2024-01-24 01:05:37 +08:00
|
|
|
streaming=False
|
|
|
|
)
|
|
|
|
|
|
|
|
return {'data': response.data.decode('latin1')}
|
|
|
|
except services.errors.app_model_config.AppModelConfigBrokenError:
|
|
|
|
logging.exception("App model config broken.")
|
|
|
|
raise AppUnavailableError()
|
|
|
|
except NoAudioUploadedServiceError:
|
|
|
|
raise NoAudioUploadedError()
|
|
|
|
except AudioTooLargeServiceError as e:
|
|
|
|
raise AudioTooLargeError(str(e))
|
|
|
|
except UnsupportedAudioTypeServiceError:
|
|
|
|
raise UnsupportedAudioTypeError()
|
|
|
|
except ProviderNotSupportSpeechToTextServiceError:
|
|
|
|
raise ProviderNotSupportSpeechToTextError()
|
|
|
|
except ProviderTokenNotInitError as ex:
|
|
|
|
raise ProviderNotInitializeError(ex.description)
|
|
|
|
except QuotaExceededError:
|
|
|
|
raise ProviderQuotaExceededError()
|
|
|
|
except ModelCurrentlyNotSupportError:
|
|
|
|
raise ProviderModelCurrentlyNotSupportError()
|
|
|
|
except InvokeError as e:
|
|
|
|
raise CompletionRequestError(e.description)
|
|
|
|
except ValueError as e:
|
|
|
|
raise e
|
|
|
|
except Exception as e:
|
2024-02-15 22:41:18 +08:00
|
|
|
logging.exception(f"internal server error: {str(e)}")
|
2024-01-24 01:05:37 +08:00
|
|
|
raise InternalServerError()
|
|
|
|
|
|
|
|
|
|
|
|
api.add_resource(AudioApi, '/audio-to-text')
|
|
|
|
api.add_resource(TextApi, '/text-to-audio')
|