| any_to_any.html | 26.9 kB | | aad9a091 |
| any_to_any.md | 4.85 kB | | c50c9b05 |
| asr.html | 71.4 kB | | 3b88a461 |
| asr.md | 14.9 kB | | dd729979 |
| audio_classification.html | 65 kB | | cf977f37 |
| audio_classification.md | 12.3 kB | | ac71e47a |
| audio_text_to_text.html | 68.7 kB | | fdc8490b |
| audio_text_to_text.md | 12.2 kB | | ce312e20 |
| document_question_answering.html | 108 kB | | 88a844a9 |
| document_question_answering.md | 24.1 kB | | a69231b4 |
| image_captioning.html | 49.3 kB | | 51d3ea98 |
| image_captioning.md | 7.39 kB | | e9bab95b |
| image_classification.html | 60.7 kB | | 91533c08 |
| image_classification.md | 10.9 kB | | 4871166c |
| image_feature_extraction.html | 30.4 kB | | a53b4758 |
| image_feature_extraction.md | 4.49 kB | | d93a4916 |
| image_text_to_text.html | 60.7 kB | | 5718d746 |
| image_text_to_text.md | 16.6 kB | | d7f34e30 |
| instance_segmentation.html | 71.2 kB | | 5c6ca702 |
| instance_segmentation.md | 17.4 kB | | 7bb1fa77 |
| keypoint_detection.html | 25.9 kB | | 4ad4205b |
| keypoint_detection.md | 5.21 kB | | 57820119 |
| keypoint_matching.html | 23.7 kB | | 5786a213 |
| keypoint_matching.md | 4.4 kB | | d65ee050 |
| knowledge_distillation_for_image_classification.html | 31.6 kB | | ac479d72 |
| knowledge_distillation_for_image_classification.md | 7.79 kB | | 7e97937b |
| language_modeling.html | 62.4 kB | | 0f090910 |
| language_modeling.md | 14 kB | | 42f2364c |
| mask_generation.html | 79.9 kB | | 387e4b83 |
| mask_generation.md | 17.1 kB | | 43758a8a |
| masked_language_modeling.html | 63.3 kB | | bd69720a |
| masked_language_modeling.md | 13.6 kB | | abd1496e |
| monocular_depth_estimation.html | 29.5 kB | | f2a08d99 |
| monocular_depth_estimation.md | 5.82 kB | | 3b9de4fd |
| multiple_choice.html | 51.3 kB | | 27316d56 |
| multiple_choice.md | 9.42 kB | | 1f97e711 |
| object_detection.html | 91.9 kB | | bdf1af06 |
| object_detection.md | 22.8 kB | | 79c050c0 |
| prompting.html | 37 kB | | c193ae02 |
| prompting.md | 13.4 kB | | 5e705870 |
| question_answering.html | 54.7 kB | | bb8b1332 |
| question_answering.md | 11.4 kB | | 0c8ce6fc |
| semantic_segmentation.html | 114 kB | | 3546bdbb |
| semantic_segmentation.md | 23.7 kB | | 3874b60a |
| sequence_classification.html | 53.7 kB | | 3592b2e1 |
| sequence_classification.md | 10.4 kB | | a7fb4bf6 |
| summarization.html | 60.2 kB | | 757cbf07 |
| summarization.md | 17.7 kB | | 321bacfc |
| text-to-speech.html | 127 kB | | c75bdb79 |
| text-to-speech.md | 23.6 kB | | 5dcf6538 |
| token_classification.html | 73 kB | | dda9703c |
| token_classification.md | 13.9 kB | | 75b40819 |
| training_vision_backbone.html | 35.8 kB | | 597250f8 |
| training_vision_backbone.md | 8.7 kB | | 12ee3fa1 |
| translation.html | 54 kB | | ebb96a23 |
| translation.md | 10.4 kB | | c558e665 |
| video_classification.html | 89.4 kB | | 349f80ba |
| video_classification.md | 20.8 kB | | 5ac81a3b |
| video_text_to_text.html | 25.1 kB | | 79d75dbd |
| video_text_to_text.md | 6.01 kB | | f94a7484 |
| visual_document_retrieval.html | 26.9 kB | | da1518ce |
| visual_document_retrieval.md | 5.68 kB | | 5876079c |
| visual_question_answering.html | 69.7 kB | | cea7e849 |
| visual_question_answering.md | 14.8 kB | | 286a46d7 |
| zero_shot_image_classification.html | 29.6 kB | | e9fc123f |
| zero_shot_image_classification.md | 5.05 kB | | fa5af23d |
| zero_shot_object_detection.html | 57.3 kB | | f061658a |
| zero_shot_object_detection.md | 10.3 kB | | c0fbda06 |