huggingface-transformers/tests/test_modeling_tf_t5.py

# coding=utf-8
# Copyright 2018 Google T5 Authors and HuggingFace Inc. team.
#
# Licensed under the Apache License, Version 2.0 (the "License");
# you may not use this file except in compliance with the License.
# You may obtain a copy of the License at
#
#     http://www.apache.org/licenses/LICENSE-2.0
#
# Unless required by applicable law or agreed to in writing, software
# distributed under the License is distributed on an "AS IS" BASIS,
# WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
# See the License for the specific language governing permissions and
# limitations under the License.


import unittest

from transformers import T5Config, is_tf_available

from .test_configuration_common import ConfigTester
from .test_modeling_tf_common import TFModelTesterMixin, ids_tensor
from .utils import CACHE_DIR, require_tf, slow


if is_tf_available():
    from transformers.modeling_tf_t5 import TFT5Model, TFT5WithLMHeadModel


@require_tf
class TFT5ModelTest(TFModelTesterMixin, unittest.TestCase):

    is_encoder_decoder = True
    all_model_classes = (TFT5Model, TFT5WithLMHeadModel) if is_tf_available() else ()

    class TFT5ModelTester(object):
        def __init__(
            self,
            parent,
            batch_size=13,
            seq_length=7,
            is_training=True,
            use_input_mask=True,
            use_labels=True,
            vocab_size=99,
            n_positions=14,
            hidden_size=32,
            num_hidden_layers=5,
            num_attention_heads=4,
            d_ff=37,
            relative_attention_num_buckets=8,
            dropout_rate=0.1,
            initializer_factor=0.002,
            scope=None,
        ):
            self.parent = parent
            self.batch_size = batch_size
            self.seq_length = seq_length
            self.is_training = is_training
            self.use_input_mask = use_input_mask
            self.use_labels = use_labels
            self.vocab_size = vocab_size
            self.n_positions = n_positions
            self.hidden_size = hidden_size
            self.num_hidden_layers = num_hidden_layers
            self.num_attention_heads = num_attention_heads
            self.d_ff = d_ff
            self.relative_attention_num_buckets = relative_attention_num_buckets
            self.dropout_rate = dropout_rate
            self.initializer_factor = initializer_factor
            self.scope = scope

        def prepare_config_and_inputs(self):
            input_ids = ids_tensor([self.batch_size, self.seq_length], self.vocab_size)

            input_mask = None
            if self.use_input_mask:
                input_mask = ids_tensor([self.batch_size, self.seq_length], vocab_size=2)

            token_labels = None
            if self.use_labels:
                token_labels = ids_tensor([self.batch_size, self.seq_length], self.vocab_size)

            config = T5Config(
                vocab_size=self.vocab_size,
                n_positions=self.n_positions,
                d_model=self.hidden_size,
                d_ff=self.d_ff,
                d_kv=self.hidden_size // self.num_attention_heads,
                num_layers=self.num_hidden_layers,
                num_heads=self.num_attention_heads,
                relative_attention_num_buckets=self.relative_attention_num_buckets,
                dropout_rate=self.dropout_rate,
                initializer_factor=self.initializer_factor,
            )

            return (config, input_ids, input_mask, token_labels)

        def create_and_check_t5_model(self, config, input_ids, input_mask, token_labels):
            model = TFT5Model(config=config)
            inputs = {
                "encoder_input_ids": input_ids,
                "decoder_input_ids": input_ids,
                "decoder_attention_mask": input_mask,
            }
            encoder_output, decoder_output = model(inputs)

            encoder_output, decoder_output = model(
                input_ids, decoder_attention_mask=input_mask, encoder_input_ids=input_ids
            )

            result = {
                "encoder_output": encoder_output.numpy(),
                "decoder_output": decoder_output.numpy(),
            }
            self.parent.assertListEqual(
                list(result["encoder_output"].shape), [self.batch_size, self.seq_length, self.hidden_size]
            )
            self.parent.assertListEqual(
                list(result["decoder_output"].shape), [self.batch_size, self.seq_length, self.hidden_size]
            )

        def create_and_check_t5_with_lm_head(self, config, input_ids, input_mask, token_labels):
            model = TFT5WithLMHeadModel(config=config)
            inputs = {
                "encoder_input_ids": input_ids,
                "decoder_input_ids": input_ids,
                "decoder_attention_mask": input_mask,
            }
            prediction_scores, decoder_output = model(inputs)
            result = {
                "prediction_scores": prediction_scores.numpy(),
            }
            self.parent.assertListEqual(
                list(result["prediction_scores"].shape), [self.batch_size, self.seq_length, self.vocab_size]
            )

        def prepare_config_and_inputs_for_common(self):
            config_and_inputs = self.prepare_config_and_inputs()
            (config, input_ids, input_mask, token_labels) = config_and_inputs
            inputs_dict = {
                "encoder_input_ids": input_ids,
                "decoder_input_ids": input_ids,
                "decoder_attention_mask": input_mask,
            }
            return config, inputs_dict

    def setUp(self):
        self.model_tester = TFT5ModelTest.TFT5ModelTester(self)
        self.config_tester = ConfigTester(self, config_class=T5Config, d_model=37)

    def test_config(self):
        self.config_tester.run_common_tests()

    def test_t5_model(self):
        config_and_inputs = self.model_tester.prepare_config_and_inputs()
        self.model_tester.create_and_check_t5_model(*config_and_inputs)

    def test_with_lm_head(self):
        config_and_inputs = self.model_tester.prepare_config_and_inputs()
        self.model_tester.create_and_check_t5_with_lm_head(*config_and_inputs)

    @slow
    def test_model_from_pretrained(self):
        for model_name in ["t5-small"]:
            model = TFT5Model.from_pretrained(model_name, cache_dir=CACHE_DIR)
            self.assertIsNotNone(model)
adding tests and updating model 2019-11-06 13:52:50 +03:00			`# coding=utf-8`
			`# Copyright 2018 Google T5 Authors and HuggingFace Inc. team.`
			`#`
			`# Licensed under the Apache License, Version 2.0 (the "License");`
			`# you may not use this file except in compliance with the License.`
			`# You may obtain a copy of the License at`
			`#`
			`# http://www.apache.org/licenses/LICENSE-2.0`
			`#`
			`# Unless required by applicable law or agreed to in writing, software`
			`# distributed under the License is distributed on an "AS IS" BASIS,`
			`# WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.`
			`# See the License for the specific language governing permissions and`
			`# limitations under the License.`
Remove __future__ imports. 2019-12-22 18:20:32 +03:00
adding tests and updating model 2019-11-06 13:52:50 +03:00
Replace (TF)CommonTestCases for modeling with a mixin. I suspect the wrapper classes were created in order to prevent the abstract base class (TF)CommonModelTester from being included in test discovery and running, because that would fail. I solved this by replacing the abstract base class with a mixin. Code changes are just de-indenting and automatic reformattings performed by black to use the extra line space. 2019-12-22 16:57:20 +03:00			`import unittest`

Sort imports with isort. This is the result of: $ isort --recursive examples templates transformers utils hubconf.py setup.py 2019-12-21 17:57:32 +03:00			`from transformers import T5Config, is_tf_available`
adding tests and updating model 2019-11-06 13:52:50 +03:00
Switch test files to the standard test_*.py scheme. 2019-12-22 15:44:13 +03:00			`from .test_configuration_common import ConfigTester`
Replace (TF)CommonTestCases for modeling with a mixin. I suspect the wrapper classes were created in order to prevent the abstract base class (TF)CommonModelTester from being included in test discovery and running, because that would fail. I solved this by replacing the abstract base class with a mixin. Code changes are just de-indenting and automatic reformattings performed by black to use the extra line space. 2019-12-22 16:57:20 +03:00			`from .test_modeling_tf_common import TFModelTesterMixin, ids_tensor`
Take advantage of the cache when running tests. Caching models across test cases and across runs of the test suite makes slow tests somewhat more bearable. Use gettempdir() instead of /tmp in tests. This makes it easier to change the location of the cache with semi-standard TMPDIR/TEMP/TMP environment variables. Fix #2222. 2019-12-20 22:56:58 +03:00			`from .utils import CACHE_DIR, require_tf, slow`
adding tests and updating model 2019-11-06 13:52:50 +03:00

added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`if is_tf_available():`
Fix F401 flake8 warning (x88 / 116). This change is mostly autogenerated with: $ python -m autoflake --in-place --recursive --remove-all-unused-imports --ignore-init-module-imports examples templates transformers utils hubconf.py setup.py I made minor changes in the generated diff. 2019-12-21 23:54:07 +03:00			`from transformers.modeling_tf_t5 import TFT5Model, TFT5WithLMHeadModel`
adding tests and updating model 2019-11-06 13:52:50 +03:00

updating tests and TF 2.0 model 2019-12-10 17:11:07 +03:00			`@require_tf`
Replace (TF)CommonTestCases for modeling with a mixin. I suspect the wrapper classes were created in order to prevent the abstract base class (TF)CommonModelTester from being included in test discovery and running, because that would fail. I solved this by replacing the abstract base class with a mixin. Code changes are just de-indenting and automatic reformattings performed by black to use the extra line space. 2019-12-22 16:57:20 +03:00			`class TFT5ModelTest(TFModelTesterMixin, unittest.TestCase):`
adding tests and updating model 2019-11-06 13:52:50 +03:00
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`is_encoder_decoder = True`
			`all_model_classes = (TFT5Model, TFT5WithLMHeadModel) if is_tf_available() else ()`
adding tests and updating model 2019-11-06 13:52:50 +03:00
			`class TFT5ModelTester(object):`
Reformat source code with black. This is the result of: $ black --line-length 119 examples templates transformers utils hubconf.py setup.py There's a lot of fairly long lines in the project. As a consequence, I'm picking the longest widely accepted line length, 119 characters. This is also Thomas' preference, because it allows for explicit variable names, to make the code easier to understand. 2019-12-21 17:46:46 +03:00			`def __init__(`
			`self,`
			`parent,`
			`batch_size=13,`
			`seq_length=7,`
			`is_training=True,`
			`use_input_mask=True,`
			`use_labels=True,`
			`vocab_size=99,`
			`n_positions=14,`
			`hidden_size=32,`
			`num_hidden_layers=5,`
			`num_attention_heads=4,`
			`d_ff=37,`
			`relative_attention_num_buckets=8,`
			`dropout_rate=0.1,`
			`initializer_factor=0.002,`
			`scope=None,`
			`):`
adding tests and updating model 2019-11-06 13:52:50 +03:00			`self.parent = parent`
			`self.batch_size = batch_size`
			`self.seq_length = seq_length`
			`self.is_training = is_training`
			`self.use_input_mask = use_input_mask`
			`self.use_labels = use_labels`
			`self.vocab_size = vocab_size`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`self.n_positions = n_positions`
adding tests and updating model 2019-11-06 13:52:50 +03:00			`self.hidden_size = hidden_size`
			`self.num_hidden_layers = num_hidden_layers`
			`self.num_attention_heads = num_attention_heads`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`self.d_ff = d_ff`
			`self.relative_attention_num_buckets = relative_attention_num_buckets`
			`self.dropout_rate = dropout_rate`
			`self.initializer_factor = initializer_factor`
adding tests and updating model 2019-11-06 13:52:50 +03:00			`self.scope = scope`

			`def prepare_config_and_inputs(self):`
			`input_ids = ids_tensor([self.batch_size, self.seq_length], self.vocab_size)`

			`input_mask = None`
			`if self.use_input_mask:`
			`input_mask = ids_tensor([self.batch_size, self.seq_length], vocab_size=2)`

			`token_labels = None`
			`if self.use_labels:`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`token_labels = ids_tensor([self.batch_size, self.seq_length], self.vocab_size)`
adding tests and updating model 2019-11-06 13:52:50 +03:00
			`config = T5Config(`
update t5 tf 2019-12-16 11:59:36 +03:00			`vocab_size=self.vocab_size,`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`n_positions=self.n_positions,`
			`d_model=self.hidden_size,`
			`d_ff=self.d_ff,`
			`d_kv=self.hidden_size // self.num_attention_heads,`
			`num_layers=self.num_hidden_layers,`
			`num_heads=self.num_attention_heads,`
			`relative_attention_num_buckets=self.relative_attention_num_buckets,`
			`dropout_rate=self.dropout_rate,`
Reformat source code with black. This is the result of: $ black --line-length 119 examples templates transformers utils hubconf.py setup.py There's a lot of fairly long lines in the project. As a consequence, I'm picking the longest widely accepted line length, 119 characters. This is also Thomas' preference, because it allows for explicit variable names, to make the code easier to understand. 2019-12-21 17:46:46 +03:00			`initializer_factor=self.initializer_factor,`
			`)`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00
			`return (config, input_ids, input_mask, token_labels)`

			`def create_and_check_t5_model(self, config, input_ids, input_mask, token_labels):`
adding tests and updating model 2019-11-06 13:52:50 +03:00			`model = TFT5Model(config=config)`
Reformat source code with black. This is the result of: $ black --line-length 119 examples templates transformers utils hubconf.py setup.py There's a lot of fairly long lines in the project. As a consequence, I'm picking the longest widely accepted line length, 119 characters. This is also Thomas' preference, because it allows for explicit variable names, to make the code easier to understand. 2019-12-21 17:46:46 +03:00			`inputs = {`
			`"encoder_input_ids": input_ids,`
			`"decoder_input_ids": input_ids,`
			`"decoder_attention_mask": input_mask,`
			`}`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`encoder_output, decoder_output = model(inputs)`
adding tests and updating model 2019-11-06 13:52:50 +03:00
Reformat source code with black. This is the result of: $ black --line-length 119 examples templates transformers utils hubconf.py setup.py There's a lot of fairly long lines in the project. As a consequence, I'm picking the longest widely accepted line length, 119 characters. This is also Thomas' preference, because it allows for explicit variable names, to make the code easier to understand. 2019-12-21 17:46:46 +03:00			`encoder_output, decoder_output = model(`
			`input_ids, decoder_attention_mask=input_mask, encoder_input_ids=input_ids`
			`)`
adding tests and updating model 2019-11-06 13:52:50 +03:00
			`result = {`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`"encoder_output": encoder_output.numpy(),`
			`"decoder_output": decoder_output.numpy(),`
adding tests and updating model 2019-11-06 13:52:50 +03:00			`}`
			`self.parent.assertListEqual(`
Reformat source code with black. This is the result of: $ black --line-length 119 examples templates transformers utils hubconf.py setup.py There's a lot of fairly long lines in the project. As a consequence, I'm picking the longest widely accepted line length, 119 characters. This is also Thomas' preference, because it allows for explicit variable names, to make the code easier to understand. 2019-12-21 17:46:46 +03:00			`list(result["encoder_output"].shape), [self.batch_size, self.seq_length, self.hidden_size]`
			`)`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`self.parent.assertListEqual(`
Reformat source code with black. This is the result of: $ black --line-length 119 examples templates transformers utils hubconf.py setup.py There's a lot of fairly long lines in the project. As a consequence, I'm picking the longest widely accepted line length, 119 characters. This is also Thomas' preference, because it allows for explicit variable names, to make the code easier to understand. 2019-12-21 17:46:46 +03:00			`list(result["decoder_output"].shape), [self.batch_size, self.seq_length, self.hidden_size]`
			`)`
adding tests and updating model 2019-11-06 13:52:50 +03:00
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`def create_and_check_t5_with_lm_head(self, config, input_ids, input_mask, token_labels):`
adding tests and updating model 2019-11-06 13:52:50 +03:00			`model = TFT5WithLMHeadModel(config=config)`
Reformat source code with black. This is the result of: $ black --line-length 119 examples templates transformers utils hubconf.py setup.py There's a lot of fairly long lines in the project. As a consequence, I'm picking the longest widely accepted line length, 119 characters. This is also Thomas' preference, because it allows for explicit variable names, to make the code easier to understand. 2019-12-21 17:46:46 +03:00			`inputs = {`
			`"encoder_input_ids": input_ids,`
			`"decoder_input_ids": input_ids,`
			`"decoder_attention_mask": input_mask,`
			`}`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`prediction_scores, decoder_output = model(inputs)`
adding tests and updating model 2019-11-06 13:52:50 +03:00			`result = {`
			`"prediction_scores": prediction_scores.numpy(),`
			`}`
			`self.parent.assertListEqual(`
Reformat source code with black. This is the result of: $ black --line-length 119 examples templates transformers utils hubconf.py setup.py There's a lot of fairly long lines in the project. As a consequence, I'm picking the longest widely accepted line length, 119 characters. This is also Thomas' preference, because it allows for explicit variable names, to make the code easier to understand. 2019-12-21 17:46:46 +03:00			`list(result["prediction_scores"].shape), [self.batch_size, self.seq_length, self.vocab_size]`
			`)`
adding tests and updating model 2019-11-06 13:52:50 +03:00
			`def prepare_config_and_inputs_for_common(self):`
			`config_and_inputs = self.prepare_config_and_inputs()`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`(config, input_ids, input_mask, token_labels) = config_and_inputs`
Reformat source code with black. This is the result of: $ black --line-length 119 examples templates transformers utils hubconf.py setup.py There's a lot of fairly long lines in the project. As a consequence, I'm picking the longest widely accepted line length, 119 characters. This is also Thomas' preference, because it allows for explicit variable names, to make the code easier to understand. 2019-12-21 17:46:46 +03:00			`inputs_dict = {`
			`"encoder_input_ids": input_ids,`
			`"decoder_input_ids": input_ids,`
			`"decoder_attention_mask": input_mask,`
			`}`
adding tests and updating model 2019-11-06 13:52:50 +03:00			`return config, inputs_dict`

			`def setUp(self):`
			`self.model_tester = TFT5ModelTest.TFT5ModelTester(self)`
added TF2 model and tests - updated templates 2019-11-08 13:35:03 +03:00			`self.config_tester = ConfigTester(self, config_class=T5Config, d_model=37)`
adding tests and updating model 2019-11-06 13:52:50 +03:00
			`def test_config(self):`
			`self.config_tester.run_common_tests()`

			`def test_t5_model(self):`
			`config_and_inputs = self.model_tester.prepare_config_and_inputs()`
			`self.model_tester.create_and_check_t5_model(*config_and_inputs)`

			`def test_with_lm_head(self):`
			`config_and_inputs = self.model_tester.prepare_config_and_inputs()`
			`self.model_tester.create_and_check_t5_with_lm_head(*config_and_inputs)`

updating tests and TF 2.0 model 2019-12-10 17:11:07 +03:00			`@slow`
adding tests and updating model 2019-11-06 13:52:50 +03:00			`def test_model_from_pretrained(self):`
Reformat source code with black. This is the result of: $ black --line-length 119 examples templates transformers utils hubconf.py setup.py There's a lot of fairly long lines in the project. As a consequence, I'm picking the longest widely accepted line length, 119 characters. This is also Thomas' preference, because it allows for explicit variable names, to make the code easier to understand. 2019-12-21 17:46:46 +03:00			`for model_name in ["t5-small"]:`
Take advantage of the cache when running tests. Caching models across test cases and across runs of the test suite makes slow tests somewhat more bearable. Use gettempdir() instead of /tmp in tests. This makes it easier to change the location of the cache with semi-standard TMPDIR/TEMP/TMP environment variables. Fix #2222. 2019-12-20 22:56:58 +03:00			`model = TFT5Model.from_pretrained(model_name, cache_dir=CACHE_DIR)`
adding tests and updating model 2019-11-06 13:52:50 +03:00			`self.assertIsNotNone(model)`