Code
Hub
Workspaces
Following
Trending
Connect
MCP
copy
Create free account
hub
/
github.com/VITA-MLLM/VITA-Audio
/ functions
Functions
566 in github.com/VITA-MLLM/VITA-Audio
⨍
Functions
566
◇
Types & classes
121
↳
Endpoints
4
Method
inference
( self, data_in, data_lengths=None, key: list = ["wav_file_tmp_name"],
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/modeling_sensevoice.py:861
Function
init_weights
(m)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/resampler_projector.py:30
Method
is_contiguous
(self)
vita_audio/data/processor/audio_processor.py:71
Method
is_discrete
(self)
vita_audio/data/processor/audio_processor.py:67
Method
keys
(self)
evaluation/compute-wer.py:239
Method
load_cmvn
(self)
web/vad.py:150
Function
load_data_one
(data_file, output_dir)
vita_audio/data/dataset_base.py:326
Function
load_json_A
(data_file)
vita_audio/data/dataset_base.py:280
Function
load_json_B
(data_file)
vita_audio/data/dataset_base.py:289
Function
load_json_C
(data_file)
vita_audio/data/dataset_base.py:296
Function
make_inputs_require_grad
(module, input, output)
tools/finetune_sts_v4_48_3.py:533
Method
post_process_history
(self, history)
web/vad.py:180
Function
predict
(_chatbot, task_history, task)
web_demo.py:65
Method
prepare_for_tokenization
(self, text, **kwargs)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/tokenization_qwen2.py:339
Method
prepare_for_tokenization
(self, text, **kwargs)
vita_audio/models/qwen2_v4_48_3/tokenization_qwen2.py:337
Method
prepare_for_tokenization
(self, text, **kwargs)
vita_audio/models/qwen2_mtp_v4_48_3/tokenization_qwen2.py:339
Function
preprocess_logits_for_metrics
(logits, labels)
tools/finetune_sts_v4_48_3.py:655
Method
print_info
(self)
web/pool.py:53
Method
print_info
(self)
web/pool.py:92
Method
put
Receives tokens, decodes them, and prints them to stdout as soon as they form entire words.
tools/inference_sts.py:85
Method
put
Receives tokens, decodes them, and prints them to stdout as soon as they form entire words.
tools/inference_sts.py:140
Method
put
Add an item to the queue in a thread-safe manner. Parameters: - item (any): The item to be added to the queue. Retu
web/queue.py:100
Method
release
(self, obj)
web/pool.py:88
Method
release
(self)
web/parms.py:90
Function
reset_state
(task_history)
web_demo.py:246
Function
reset_user_input
()
web_demo.py:242
Function
safe_globals
()
tools/trainer_v4_48_3.py:274
Method
save_video_frames
(self, vid_path, max_fps=1, num_frames=8)
vita_audio/data/processor/image_processor.py:86
Method
save_vocabulary
(self, save_directory: str, filename_prefix: Optional[str] = None)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/tokenization_qwen2_fast.py:132
Method
save_vocabulary
(self, save_directory: str, filename_prefix: Optional[str] = None)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/tokenization_qwen2.py:310
Method
save_vocabulary
(self, save_directory: str, filename_prefix: Optional[str] = None)
vita_audio/models/qwen2_v4_48_3/tokenization_qwen2_fast.py:132
Method
save_vocabulary
(self, save_directory: str, filename_prefix: Optional[str] = None)
vita_audio/models/qwen2_v4_48_3/tokenization_qwen2.py:308
Method
save_vocabulary
(self, save_directory: str, filename_prefix: Optional[str] = None)
vita_audio/models/qwen2_mtp_v4_48_3/tokenization_qwen2_fast.py:132
Method
save_vocabulary
(self, save_directory: str, filename_prefix: Optional[str] = None)
vita_audio/models/qwen2_mtp_v4_48_3/tokenization_qwen2.py:310
Function
send_pcm
Sends PCM audio data to the dialogue system for processing. Parameters: - sid (str): The session ID of the user.
web_demo_stream.py:373
Method
set_decoder
(self, decoder)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/modeling_qwen2.py:866
Method
set_decoder
(self, decoder)
vita_audio/models/qwen2_v4_48_3/modeling_qwen2.py:758
Method
set_decoder
(self, decoder)
vita_audio/models/qwen2_mtp_v4_48_3/modeling_qwen2.py:813
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/modeling_qwen2.py:546
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/modeling_qwen2.py:857
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/modeling_qwen2.py:1378
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/modeling_qwen2.py:1481
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/modeling_qwen2.py:1563
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_v4_48_3/modeling_qwen2.py:498
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_v4_48_3/modeling_qwen2.py:749
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_v4_48_3/modeling_qwen2.py:882
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_v4_48_3/modeling_qwen2.py:985
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_v4_48_3/modeling_qwen2.py:1067
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_mtp_v4_48_3/modeling_qwen2.py:540
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_mtp_v4_48_3/modeling_qwen2.py:804
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_mtp_v4_48_3/modeling_qwen2.py:1321
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_mtp_v4_48_3/modeling_qwen2.py:1424
Method
set_input_embeddings
(self, value)
vita_audio/models/qwen2_mtp_v4_48_3/modeling_qwen2.py:1506
Method
set_output_embeddings
(self, new_embeddings)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/modeling_qwen2.py:863
Method
set_output_embeddings
(self, new_embeddings)
vita_audio/models/qwen2_v4_48_3/modeling_qwen2.py:755
Method
set_output_embeddings
(self, new_embeddings)
vita_audio/models/qwen2_mtp_v4_48_3/modeling_qwen2.py:810
Method
set_prompt
(self, prompt)
web/parms.py:53
Function
snac
()
tools/get_neural_audio_codecs.py:220
Function
sparktts
()
tools/get_neural_audio_codecs.py:300
Function
text_audio_interval_old
(input_ids, AUD_START_ID, AUD_END_ID, text_audio_interval_ratio)
vita_audio/data/dataset_deepseek.py:827
Function
text_audio_interval_old
(input_ids, AUD_START_ID, AUD_END_ID, text_audio_interval_ratio)
vita_audio/data/processor/audio_processor.py:146
Method
training_step
Perform a training step on a batch of inputs. Subclass and override to inject custom behavior. Args: model (`nn
tools/trainer_v4_48_3.py:1093
Method
vocab_size
(self)
vita_audio/models/qwen2_mtp_sensevoice_v4_48_3/tokenization_qwen2.py:213
Method
vocab_size
(self)
vita_audio/models/qwen2_v4_48_3/tokenization_qwen2.py:211
Method
vocab_size
(self)
vita_audio/models/qwen2_mtp_v4_48_3/tokenization_qwen2.py:213
Function
xcodec2
()
tools/get_neural_audio_codecs.py:77
← previous
501–566 of 566, ranked by callers