<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE ep-patent-document PUBLIC "-//EPO//EP PATENT DOCUMENT 1.4//EN" "ep-patent-document-v1-4.dtd">
<ep-patent-document id="EP11159156B1" file="EP11159156NWB1.xml" lang="en" country="EP" doc-number="2369836" kind="B1" date-publ="20140423" status="n" dtd-version="ep-patent-document-v1-4">
<SDOBI lang="en"><B000><eptags><B001EP>ATBECHDEDKESFRGBGRITLILUNLSEMCPTIESILTLVFIRO..CY..TRBGCZEEHUPLSK....IS..MT..........................</B001EP><B005EP>J</B005EP><B007EP>DIM360 Ver 2.40 (30 Jan 2013) -  2100000/0</B007EP></eptags></B000><B100><B110>2369836</B110><B120><B121>EUROPEAN PATENT SPECIFICATION</B121></B120><B130>B1</B130><B140><date>20140423</date></B140><B190>EP</B190></B100><B200><B210>11159156.6</B210><B220><date>20070516</date></B220><B240><B241><date>20120614</date></B241><B242><date>20120926</date></B242></B240><B250>en</B250><B251EP>en</B251EP><B260>en</B260></B200><B300><B310>20060045184</B310><B320><date>20060519</date></B320><B330><ctry>KR</ctry></B330></B300><B400><B405><date>20140423</date><bnum>201417</bnum></B405><B430><date>20110928</date><bnum>201139</bnum></B430><B450><date>20140423</date><bnum>201417</bnum></B450><B452EP><date>20131111</date></B452EP></B400><B500><B510EP><classification-ipcr sequence="1"><text>H04N   7/00        20110101AFI20111104BHEP        </text></classification-ipcr><classification-ipcr sequence="2"><text>H04H  20/89        20080101ALI20111104BHEP        </text></classification-ipcr><classification-ipcr sequence="3"><text>H04S   7/00        20060101ALI20111104BHEP        </text></classification-ipcr></B510EP><B540><B541>de</B541><B542>Objektbasiertes dreidimensionales Audiodienstsystem mit im Voraus eingestellten Audioszenen</B542><B541>en</B541><B542>Object-based 3-dimensional audio service system using preset audio scenes</B542><B541>fr</B541><B542>Système de service audio tridimensionnel à base d'objet utilisant des scènes audio prédéfinies</B542></B540><B560><B561><text>US-B1- 6 665 318</text></B561><B562><text>KALVA H ET AL: "DELIVERING OBJECT-BASED AUDIO-VISUAL SERVICES", IEEE TRANSACTIONS ON CONSUMER ELECTRONICS, IEEE SERVICE CENTER, NEW YORK, NY, US LNKD- DOI:10.1109/30.809189, vol. 45, no. 4, 1 November 1999 (1999-11-01), pages 1108-1111, XP000928095, ISSN: 0098-3063</text></B562><B562><text>AVARO O ET AL: "The MPEG-4 systems and description languages: A way ahead in audio visual information representation", SIGNAL PROCESSING. IMAGE COMMUNICATION, ELSEVIER SCIENCE PUBLISHERS, AMSTERDAM, NL LNKD- DOI:10.1016/S0923-5965(97)00027-1, vol. 9, no. 4, 1 May 1997 (1997-05-01), pages 385-431, XP004075337, ISSN: 0923-5965</text></B562><B562><text>TAEJIN LEE, GI YOON PARK, INSEON JANG, KYEONGOK KANG: "An Object-based 3D Audio Broadcasting System for Interactive Service.", CONVENTION PAPER 6384 118TH CONVENTION AUDIO ENGINEERING SOCIETY, 31 May 2005 (2005-05-31), XP002577516, Retrieved from the Internet: URL:http://www.aes.org/tmpFiles/elib/20100 413/13100.pdf [retrieved on 2010-04-12]</text></B562></B560></B500><B600><B620><parent><pdoc><dnum><anum>07746543.3</anum><pnum>2022263</pnum></dnum><date>20070516</date></pdoc></parent></B620><B620EP><parent><cdoc><dnum><anum>12171933.0</anum><pnum>2501128</pnum></dnum><date>20120614</date></cdoc></parent></B620EP></B600><B700><B720><B721><snm>Lee, Yong-Ju</snm><adr><str>
No. 302, 148-6, Sinseong-dong, Yuseong-gu</str><city>305-345, Daejon</city><ctry>KR</ctry></adr></B721><B721><snm>Lee, Tae-Jin</snm><adr><str>No. 101-705 Cheongsol Apt., Songgang-dong, 
Yuseon-gu</str><city>305-752, Daejon</city><ctry>KR</ctry></adr></B721><B721><snm>Yoo, Jae-Hyoun</snm><adr><str>No. 1-219 ETRI Dormitory, 236-1 Gajeong-dong, 
Yuseong-gu</str><city>305-350, Daejon</city><ctry>KR</ctry></adr></B721><B721><snm>Kang, Kyeong-Ok</snm><adr><str>No. 101-605 Samsung Pureun Apt. Jeonmin-dong, 
Yuseong-gu</str><city>305-727, Daejon</city><ctry>KR</ctry></adr></B721><B721><snm>Hong, Jin-Woo</snm><adr><str>No. 130-702 Hanbit Apt., Eoeun-dong, 
Yuseong-gu</str><city>305-333, Daejon</city><ctry>KR</ctry></adr></B721><B721><snm>Jang, In-Seon</snm><adr><str>
148-4, Sinseong-dong, Yuseong-gu</str><city>305-345, Daejon</city><ctry>KR</ctry></adr></B721><B721><snm>Seo, Jeong-Il</snm><adr><str>No. 107-801 Sejong Apt., Jeonmin-dong, 
Yuseong-gu</str><city>305-728, Daejon</city><ctry>KR</ctry></adr></B721><B721><snm>Jang, Dae-Young</snm><adr><str>No. 904-1701 Yeolmae Maeul 9 Danji 
Noeun-dong, Yuseong-gu</str><city>305-768, Daejon</city><ctry>KR</ctry></adr></B721></B720><B730><B731><snm>Electronics and Telecommunications Research 
Institute</snm><iid>100115996</iid><irf>EMU 042 EP DIV</irf><adr><str>161 Gajeong-dong, 
Yuseong-gu</str><city>Daejeon 305-350</city><ctry>KR</ctry></adr></B731></B730><B740><B741><snm>Betten &amp; Resch</snm><iid>100060687</iid><adr><str>Theatinerstrasse 8</str><city>80333 München</city><ctry>DE</ctry></adr></B741></B740></B700><B800><B840><ctry>AT</ctry><ctry>BE</ctry><ctry>BG</ctry><ctry>CH</ctry><ctry>CY</ctry><ctry>CZ</ctry><ctry>DE</ctry><ctry>DK</ctry><ctry>EE</ctry><ctry>ES</ctry><ctry>FI</ctry><ctry>FR</ctry><ctry>GB</ctry><ctry>GR</ctry><ctry>HU</ctry><ctry>IE</ctry><ctry>IS</ctry><ctry>IT</ctry><ctry>LI</ctry><ctry>LT</ctry><ctry>LU</ctry><ctry>LV</ctry><ctry>MC</ctry><ctry>MT</ctry><ctry>NL</ctry><ctry>PL</ctry><ctry>PT</ctry><ctry>RO</ctry><ctry>SE</ctry><ctry>SI</ctry><ctry>SK</ctry><ctry>TR</ctry></B840><B880><date>20111214</date><bnum>201150</bnum></B880></B800></SDOBI>
<description id="desc" lang="en"><!-- EPO <DP n="1"> -->
<heading id="h0001"><b>TECHNICAL FIELD</b></heading>
<p id="p0001" num="0001">The present invention relates to an object-based three dimensional (3-D) audio service system using preset audio scenes and a method thereof; and, more particularly, to an object-based 3-D audio service system using preset audio scenes and a method thereof for providing an interactive service that enables a user or a viewer to directly form an audio scene using a 3-D audio related technology for providing realistic broadcasting to a user or a viewer.</p>
<heading id="h0002"><b>BACKGROUND ART</b></heading>
<p id="p0002" num="0002"><figref idref="f0001">Fig. 1</figref> is a diagram illustrating a conventional audio service system.</p>
<p id="p0003" num="0003">As shown in <figref idref="f0001">Fig. 1</figref>, the conventional audio service system includes an audio service providing apparatus 10 and an audio service reproducing apparatus 20. The audio service providing apparatus 10 includes an audio-capture unit 11 for capturing an audio signal such as sound, an editing/mixing unit 12 for editing and mixing the captured audio signal to transmit the audio signal to an audio service reproducing apparatus 20, and a storing/transmitting unit 13 for storing the mixed audio signal and transmitting the mixed audio signal to the audio service reproducing apparatus 20.</p>
<p id="p0004" num="0004">The audio service reproducing apparatus 20 includes a receiver 21 for receiving an audio signal transmitted from the audio service providing apparatus in, a controller 22 for controlling the received audio signal, and a reproducer 23 for reproducing an audio signal.</p>
<p id="p0005" num="0005">An audio signal, which is provided through<!-- EPO <DP n="2"> --> broadcasting services such as TV broadcasting, radio broadcasting, and Digital Multimedia Broadcasting (DMB) based on the conventional audio service system, is generally created by mixing a plurality of audio signals captured from various sound sources. For example, an audio signal provided through a soccer game broadcasting is created by mixing noises in a soccer stadium, yelling of a crowd, and a voice of an announcer.</p>
<p id="p0006" num="0006">Although a user or a viewer can control the volume of the overall audio signal, it is impossible to control the volume of each object such as the voice of an announcer, the yelling of a crowd, and the noises of the soccer stadium. It is because the audio signal is transmitted after a plurality of object audio signals are mixed into one audio signal in a general broadcasting service.</p>
<p id="p0007" num="0007">However, if a transmitter such as the audio service providing apparatus 10 independently transmits object audio signals of the sound sources without the object audio signals of the sound sources mixed to one audio signal, a receiver such as the audio service reproducing apparatus 20 can independently control the volumes of the object audio signals of the sound sources. An object-based audio service denotes such an audio service that allows a user or a viewer to control each of the object audio signals at a receiver by independently transmitting the object audio signals of the sound sources through a transmitter.</p>
<p id="p0008" num="0008">For example, if an audio signal of a soccer game broadcasting is provided based on an object-based 3-D audio service, a user or a viewer can control each of objects, such as the noises in the soccer stadium, the yelling of the crowd, and the voices of an announcer to obtain a desired audio setting. That is, a user or a viewer can control the noise of the soccer stadium loud,<!-- EPO <DP n="3"> --> the yelling of the crowd soft, and the voice of the announcer loud. Or, a viewer can control the audio signal to reproduce only the noises of the soccer stadium and the voice of an announcer without the yelling of the crowd reproduced.</p>
<p id="p0009" num="0009">Therefore, there is a great demand for developing a method for providing an object-based 3-D audio service that enables a user to control each of object audio signals of sound sources., which can be applied to all broadcasting services and multimedia services providing audio such as digital broadcasting, radio broadcasting, Digital Multimedia Broadcasting, internet broadcasting, digital movie, DVD, moving picture contents.</p>
<p id="p0010" num="0010">Although a conventional object-based 3-D audio system and a control method thereof was introduced in Korean Patent Publication No. <patcit id="pcit0001" dnum="KR1020040037437"><text>10-2004-0037437, published on May 7th 2004</text></patcit>, the conventional object-based 3-D audio system requires a user to control each of object audio signals of sound sources to set the audio signals according two user's preference. Therefore, it is very annoying to a user or a viewer.</p>
<p id="p0011" num="0011"><nplcit id="ncit0001" npl-type="s"><text>Kalva H. et al., "DELIVERING OBJECT-BASED AUDIO-VISUAL SERVICES", IEEE TRANSACTIONS ON CONSUMER ELECTRONICS, IEEE SERVICE CENTER, NEW YORK, US, vol. 45, no,4, 1 November 1999 (1999-11-01), pages 1108-1111</text></nplcit>, discloses a method for delivering audio-visual services based on MPEG-4. In particular, algorithms to schedule the delivery of object-based audio-visual presentations in general and MPEG-4 presentations are disclosed.</p>
<p id="p0012" num="0012">The document <patcit id="pcit0002" dnum="US6665A"><text>US 6,665</text></patcit>, <patcit id="pcit0003" dnum="US318B1"><text>318 B1</text></patcit> describes a video stream decoder which includes a demultiplexing unit for demultiplexing a video stream containing at least one or more object encoded visual or audio data and one or more scene descriptions which express scene contents by object encoded data; a decoder unit for decoding the object encoded visual data; a decoder unit for decoding the object encoded audio data; a visual synthesizing unit for synthesizing images corresponding to the object encoded visual data; an audio synthesizing unit for synthesizing sounds corresponding to the object encoded audio data; an analyzing unit for analyzing each scene description; and a selector for selecting one of at least two or more scene descriptions contained in the video stream.</p>
<heading id="h0003"><b>DISCLOSURE</b></heading>
<heading id="h0004"><b>TECHNICAL PROBLEM</b></heading>
<p id="p0013" num="0013">An embodiment of the present invention is directed to providing an object-based three dimensional (3-D) audio service system and a method thereof for enabling a user to easily and conveniently watch and listen an object-based 3-D audio service by eliminating inconvenience that requires a user to control each of object audio signals of sound sources.</p>
<p id="p0014" num="0014">Other objects and advantages of the present invention can be understood by the following description, and become apparent with reference to the embodiments of the present invention. Also, it is obvious to those<!-- EPO <DP n="4"> --><!-- EPO <DP n="5"> --> skilled in the art of the present invention that the objects and advantages of the present invention can be realized by the means as claimed and combinations thereof.</p>
<heading id="h0005"><b>TECHNICAL SOLUTION</b></heading>
<p id="p0015" num="0015">The present invention is defined in the independent claims. The dependent claims define embodiments thereof.</p>
<p id="p0016" num="0016">In accordance with an aspect of the present invention, there is provided an object-based three dimensional (3-D) audio service providing apparatus using preset audio scenes, including: audio input means for inputting an audio signal; preset audio scene generating means for extracting object audio signals from the audio signal inputted through the audio input means and generating more than one of 3-D audio scene information by arranging the extracted object audio signals in a 3-D space and editing features of each object; and encoding means for encoding and multiplexing the audio signal and the 3-D audio scene information for each object audio signal.</p>
<p id="p0017" num="0017">In accordance with another aspect of the present invention, there is provided an object-based 3-D audio service reproducing apparatus using preset audio scenes including: decoding means for de-multiplexing and decoding object-based 3-D audio contents; audio scene forming means for forming 3-D audio scene information according to one selected from a plurality of 3-D audio scene information in the de-multiplexed and decoded object-based 3-D audio contents by a user including a viewer; audio signal mixing means for controlling features of objects in an audio signal of the de-multiplexed and decoded object-based 3-D audio contents according to the formed 3-D audio scene information; and reproducing means for reproducing the audio signal with one of the features controlled.</p>
<p id="p0018" num="0018">In accordance with another aspect of the present<!-- EPO <DP n="6"> --> invention, there is provided a method for providing an object-based 3-D audio service using preset audio scenes, including the steps of: inputting an audio signal; extracting object audio signals from the inputted audio signal and generating more than one of 3-D audio scene information by arranging the extracted object audio signals in a 3-D space and editing features of each object; and encoding and multiplexing the audio signal and the 3-D audio scene information for each object audio signal.</p>
<p id="p0019" num="0019">In accordance with another aspect of the present invention, there is provided a method for reproducing object-based 3-D audio service using preset audio scenes including the steps of: de-multiplexing and decoding object-based 3-D audio contents; forming 3-D audio scene information according to one selected from a plurality of 3-D audio scene information in the de-multiplexed and decoded object-based 3-D audio contents by a user including a viewer; controlling features of objects in an audio signal of the de-multiplexed and decoded object-based 3-D audio contents according to the formed 3-D audio scene information; and reproducing the audio signal with one of the features controlled.</p>
<heading id="h0006"><b>ADVANTAGEOUS EFFECTS</b></heading>
<p id="p0020" num="0020">An object-based three dimensional (3-D) audio service system and a method thereof according to the present invention provides previously generated preset audio scenes to a user or a viewer with an object-based 3-D audio service applied to all broadcasting services and multimedia services providing audio, such as digital broadcasting, radio broadcasting. Digital Multimedia Broadcasting (DMB), Internet broadcasting, digital movies, Digital Video Disk (DVD), and moving picture contents. Therefore, the object-based 3-D audio service system and<!-- EPO <DP n="7"> --> a method thereof according to the present invention eliminates the inconvenience of a user to control each of object audio signals of sound sources and enables the user to easily and conveniently watch and listen the object-based 3-D audio service.</p>
<p id="p0021" num="0021">The present invention can be applied to broadcasting services and multimedia services providing audio, such as digital broadcasting, radio broadcasting, DMB, Internet broadcasting, digital movies, DVD, and moving picture contents, and the present invention is not limited to the types of mediums for transmitting and storing object-based audio contents for broadcasting and multimedia services providing audio.</p>
<heading id="h0007"><b>BRIEF DESCRIPTION OF THE DRAWINGS</b></heading>
<p id="p0022" num="0022">
<ul id="ul0001" list-style="none" compact="compact">
<li><figref idref="f0001">Fig. 1</figref> is a diagram illustrating a conventional audio service system.</li>
<li><figref idref="f0002">Fig. 2</figref> is a block diagram illustrating an object-based three-dimensional (3-D) audio service system using preset audio scenes in accordance with an embodiment of the present invention.</li>
<li><figref idref="f0003">Fig. 3</figref> is a flowchart illustrating a method for providing an object-based 3-D audio service using preset audio scenes in accordance with an embodiment of the present invention.</li>
<li><figref idref="f0004">Fig. 4</figref> is a flowchart illustrating a method for reproducing an object-based 3-D audio service using preset audio scenes in accordance with the embodiment of the present invention.</li>
</ul></p>
<heading id="h0008"><b>BEST MODE FOR THE INVENTION</b></heading>
<p id="p0023" num="0023">The advantages, features and aspects of the invention will become apparent from the following description of the embodiments with reference to the accompanying drawings, which is set forth hereinafter.<!-- EPO <DP n="8"> --></p>
<p id="p0024" num="0024"><figref idref="f0002">Fig. 2</figref> is a block diagram illustrating an object-based three-dimensional (3-D) audio service system using preset audio scenes in accordance with an embodiment of the present invention.</p>
<p id="p0025" num="0025">As shown in <figref idref="f0002">Fig. 2</figref>, the object-based 3-D audio service system includes an object-based 3-D audio service providing apparatus 30, a transmitting medium 50, and an object-based 3-D audio service reproducing apparatus 40. The 3-D service providing apparatus 30 receives an audio signal through various input devices, creates more than one of object-based 3-D audio scene information which can be selected by a user or a viewer, and transmits the created object-based 3-D audio scene information to the object-based 3-D audio service reproducing apparatus 40. The transmitting medium 50 is a medium such as a digital broadcasting network or an Internet network for connecting the object-based 3-D audio service providing apparatus 30 and the object-based 3-D audio service reproducing apparatus 40 through a network. The object-based 3-D audio service reproducing apparatus 40 generates more than one of 3-D audio scenes based on the object-based 3-D audio scene information transmitted from the object-based 3-D audio service providing apparatus 30.</p>
<p id="p0026" num="0026">Hereinafter, the constituent elements of the object-based 3-D audio service system using preset audio scenes according to the present embodiment will be described in detail.</p>
<p id="p0027" num="0027">The object-based 3-D audio service providing apparatus 30 includes an input unit 31, a preset audio scene generator 32, an encoder 33, and a transmitter 34. The input unit 31 receives audio signals through various input devices. The preset audio scene generator 32 extracts object-based audio signals (hereinafter, object audio signals) from the audio signal received through the input unit 31, arranges the extracted object audio<!-- EPO <DP n="9"> --> signals in a three dimensional space, and creates more than one of 3-D audio scene information by editing features such as a location, a size, a direction, and a sound field environment of each object. The encoder 33 encodes and multiplexes the audio signal inputted through the input unit 31 and the object-based 3-D audio scene information created by the preset audio scene generating unit 32 for transmitting the input audio signal and the generated preset audio scene information to the object-based 3-D audio service reproducing apparatus 40. For example, the input audio signals and the generated preset audio scene information are multiplexed to a moving picture experts group 4 (MPEG-4) file format in a digital broadcasting network. The transmitter 34 transforms the multiplexed object-based audio contents including the input audio signal and the created object-based 3-D audio scene information from the encoding unit 33 to a transport format. For example, the transmitter 34 transforms the multiplexed object-based audio contents to a MPEG-2 transport stream (TS) for a digital broadcasting network.</p>
<p id="p0028" num="0028">The transformed object-based audio contents including the input audio signal and the generated object-based 3-D audio scene information may be transmitted to the object-based 3-D audio reproducing apparatus 40 and may be stored in a storing medium.</p>
<p id="p0029" num="0029">The transmitter 34 may transmit the object-based audio contents including the input audio signal and the object-based 3-D audio scene information to the object-based 3-D audio reproducing apparatus 40 through a digital broadcasting network such as a terrestrial DMB channel 50.</p>
<p id="p0030" num="0030">If the sound source of the audio signal inputted to the input unit 31 is a mixed sound source, the preset audio scene generator 32 uses a Convolutive Blind Source<!-- EPO <DP n="10"> --> Separation technique to extract object audio signals. Especially, the preset audio scene generator 32 forms more than one of object-based 3-D audio scene information by controlling a ratio of each object-based the audio scene information of each object audio signal, which is set according to the control of a user such as an editor.</p>
<p id="p0031" num="0031">The object-based 3-D audio service reproducing apparatus 40 includes a decoder 42, an audio scene information forming unit 43, an audio signal mixer 44, and an audio signal reproducer 45. The decoder 42 de-multiplexes and decodes object-based audio contents including an audio signal and object-based 3-D audio scene information for reproducing. The audio scene information forming unit 43 provides the object-based 3-D audio scene information of the object-based 3-D audio contents, which is de-multiplexed and decoded by the decoder 42, to a user such as a viewer to select, and forms the object-based 3-D audio scene information according to the user selection. The audio signal mixer 44 mixes object audio signals of the audio signal of the de-multiplexed and decoded object-based 3-D audio contents from the decoder 42 by controlling features of each object, such as a location, a direction, a size, and a sound field of each object according to the object-based 3-D audio scene information formed by the audio scene information forming unit 43. The audio signal reproducer 45 reproduces the audio signal mixed to one object-based 3-D audio scenes by the audio signal mixer 44.</p>
<p id="p0032" num="0032">The object-based audio contents including the audio signal and the object-based 3-D audio scene information may be provided through a broadcasting service or a multimedia service such as digital broadcasting, radio broadcasting. Digital Multimedia Broadcasting (DMB), Internet broadcasting, digital movies. Digital Video Disk (DVD), and moving picture contents. Although the object-based<!-- EPO <DP n="11"> --> audio contents may be received through the receiver 41 in the present embodiment, the present invention is not limited thereto. That is, the object-based audio contents may be provided through a transmission medium or a storage medium that can provide a broadcasting service or a multimedia service that provides an audio.</p>
<p id="p0033" num="0033">The audio scene information forming unit 43 enables a user or a viewer to select features of objects such as a location, a direction, a volume, and a sound field environment of each object and forms new object-based 3-D audio scene information according to the features including a location, a direction, a volume, and a sound field environment of each object set by the user.</p>
<p id="p0034" num="0034">A user or a viewer can control features of a 3-D audio space by changing a reverberation time of a 3-D space through controlling a volume and a delay time of an initial reflected sound through the audio scene information forming unit 43.</p>
<p id="p0035" num="0035">That is, the object-based 3-D audio service system using the preset audio scene according to the present embodiment previously generates object-based 3-D audio scenes that are expected to be frequently used and provides the generated object-based 3-D audio scenes as preset audio scenes to a user or a viewer. That is, the object-based 3-D audio service system according to the present embodiment enables a user or a viewer to select one of the preset audio scenes in order to make a user to conveniently watch and listen a broadcasting program with the desired audio preference.</p>
<p id="p0036" num="0036">For example, noises of a soccer stadium, yelling of a crowd, a voice of an announcer are defined as audio objects for a soccer game broadcasting, and the defined audio objects are transmitted independently. With the audio objects, a. first audio scene having information about volume of the noises of soccer stadium, the yelling<!-- EPO <DP n="12"> --> of a crowd, and the voice of an announcer set to 1:1:1, a second audio scene having information about volume of the noises of a soccer stadium, the yelling of a crowd, and the voice of an announcer set to 1:0.5:1, and an audio scene halving information about volume of the noises of a soccer stadium, the yelling of a crowd, and the voice of an announcer set to 1:0:1 are transmitted as the preset audio scenes. Then, a user or a viewer selects one of the preset audio scenes to watch and listen the soccer game broadcasting with the desired audio preference.</p>
<p id="p0037" num="0037">A user may directly control each of the audio objects if the user cannot find a desired audio scene from the provided audio scenes. However, it is preferable to provide a large number of preset audio scenes to a user in order to enable the user to find a desired audio scene from the provided preset audio scenes.</p>
<p id="p0038" num="0038"><figref idref="f0003">Fig. 3</figref> is a flowchart illustrating a method for providing an object-based audio service using preset audio scenes in accordance with an embodiment of the present invention.</p>
<p id="p0039" num="0039">Referring to <figref idref="f0003">Fig. 3</figref>, the input unit 31 of the object-based 3-D audio service providing apparatus 30 receives an object-based audio signal through carious input device at step S301.</p>
<p id="p0040" num="0040">The preset audio scene generator 32 extracts object-based audio signals, that is, object audio signals, front the audio signal inputted through the input unit 31 at step S302. Then, the preset audio scene generator 32 generates more than one of object-based 3-D audio scene information at step s304 by arranging the extracted object: audio signals in a 3-D space and editing the features of each object audio signal such as a location, a direction, a volume, and a sound field environment of the audio object at step S303. The encoder 33 encodes and multiplexers the audio signal inputted through the<!-- EPO <DP n="13"> --> input unit 31 and the object-based 3-D audio scene information generated by the preset audio scene generator 32 at step S305. For example, the encoder 33 encodes and multiplexes the audio signal and the object-based 3-D audio scene information into MPEG-4 file format for a digital broadcasting network.</p>
<p id="p0041" num="0041">Then, the transmitter 34 transforms the multiplexed object-based audio contents including the audio signal and the object-based 3-D audio scene information to be proper to a transport format and transmits the transformed object-based audio contents at step S306. For example, the multiplexed object-based audio contents are transformed to a MFEG-2 TS in a digital broadcasting network.</p>
<p id="p0042" num="0042">For example, the transmitter 34 transmits the transformed object-based audio contents including the audio signal and the object-based 3-D audio scene information to the object-based 3-D audio reproducing apparatus 40 through a digital broadcasting network such as a terrestrial DMB channel. The transformed object-based audio contents including the audio signal and the object-based 3-D audio scene information may be stored in a storing medium.</p>
<p id="p0043" num="0043"><figref idref="f0004">Fig. 4</figref> is a flowchart illustrating a method for reproducing an object-based 3-D audio service using preset audio scenes in accordance with an embodiment of the present invention.</p>
<p id="p0044" num="0044">Referring to <figref idref="f0004">Fig. 4</figref>, the receiver 41 of the object-based 3-D audio service reproducing apparatus 40 receives the object-based audio contents including an audio signal and object-based 3-D audio information through, for example, a digital broadcasting network such as a terrestrial DMB channel 50 or the Internet network at step S401.</p>
<p id="p0045" num="0045">The receiver 41 may receive the object-based audio<!-- EPO <DP n="14"> --> contents through a transmission medium that can provide a broadcasting service or a multimedia service that provides an audio. Or, the object-based audio contents may be inputted through the storing medium.</p>
<p id="p0046" num="0046">The decoder 42 de-multiplexes and decodes the received or inputted object-based audio contents including the audio signal and the object-based 3-D audio scene information at step S402. The audio scene information forming unit 43 provides the object-based 3-D audio scene information of the de-multiplexed and decoded object-based 3-D audio contents to a user or a viewer to select, and forms object-based 3-D audio scene information according to the user selection at step S403.</p>
<p id="p0047" num="0047">Then, the audio signal mixer 44 mixers object audio signals by controlling features of objects in the audio signal of the de-multiplexed and decoded object-based 3-D audio contents, such as a location, a direction, a volume, and a sound field environment of each audio object, according to the object-based 3-D audio scene information formed by the audio scene information forming unit 43 at step S404. Finally, the audio signal reproducer 45 reproduces the audio signal mixed based on one of the object-based 3-D audio scenes by the audio signal mixer 44 at step S405.</p>
<p id="p0048" num="0048">The above described method according to the present invention can be embodied as a program and stored on a computer readable recording medium. The computer readable recording medium is any data storage device that can store data which can be thereafter read by the computer system. The computer readable recording medium includes a read-only memory (ROM), a random-access memory (RAM), a CD-ROM, a floppy disk, a hard disk and an optical magnetic disk.</p>
<p id="p0049" num="0049">While the present invention has been described with respect to certain preferred embodiments, it will be<!-- EPO <DP n="15"> --> apparent to those skilled in the art that various changes and modifications may be made without departing from the scope of the invention as defined in the following claims.</p>
</description>
<claims id="claims01" lang="en"><!-- EPO <DP n="16"> -->
<claim id="c-en-01-0001" num="0001">
<claim-text>An object-based audio service providing apparatus (30), comprising:
<claim-text>preset audio scene generating means (32) for generating a plurality of preset audio scenes, wherein each preset audio scene is based on 3D audio scene information which indicates features of a plurality of audio objects; and</claim-text>
<claim-text>encoding means (33) for encoding and multiplexing the formed audio objects and the plurality of preset audio scenes,</claim-text>
<claim-text>wherein the 3D audio scene information includes at least one of a location, a size, a volume, a direction, and a sound field environment of each of the plurality of audio objects in 3D space,</claim-text></claim-text></claim>
<claim id="c-en-01-0002" num="0002">
<claim-text>The object-based audio service providing apparatus (30) of claim 1, further comprising processing means (34) for processing the encoded and multiplexed audio objects.</claim-text></claim>
<claim id="c-en-01-0003" num="0003">
<claim-text>The object-based audio service providing apparatus of claim 1, wherein the processing means (34) transmits the encoded and multiplexed audio objects to an audio reproducing terminal (40) through a digital broadcasting network.</claim-text></claim>
<claim id="c-en-01-0004" num="0004">
<claim-text>An object-based audio service reproducing apparatus (40), comprising:
<claim-text>decoding means (42) for de-multiplexing and decoding object-based audio contents including a plurality of audio objects and a plurality of present audio scenes;</claim-text>
<claim-text>audio scene forming means (43) for forming one of said plurality of preset audio scenes , wherein each preset audio scene is based on 3D audio scene information which indicates features of said plurality of audio objects based on a selection of a user;</claim-text>
<claim-text>audio signal mixing means (44) for mixing the plurality of audio objects according to the formed preset audio scene ; and<!-- EPO <DP n="17"> --></claim-text>
<claim-text>reproducing means (45) for reproducing the mixed plurality of audio objects according to the formed preset audio scene,</claim-text>
<claim-text>wherein the 3D audio scene information includes at least one of a location, a size, a volume, a direction, and a sound field environment of each of the plurality of audio objects in 3D space.</claim-text></claim-text></claim>
<claim id="c-en-01-0005" num="0005">
<claim-text>An object-based audio service providing method, comprising:
<claim-text>generating (S304) a plurality of preset audio scenes, wherein each preset audio scene is based on 3D audio scene information which indicate features of a plurality of audio objects; and</claim-text>
<claim-text>encoding and multiplexing (S305) the plurality of audio objects and the plurality of preset audio scenes,</claim-text>
<claim-text>wherein the 3D audio scene information includes at least one of a location, a size, a volume, a direction, and a sound field environment of each of the plurality of audio objects in 3D space,</claim-text></claim-text></claim>
<claim id="c-en-01-0006" num="0006">
<claim-text>The method of claim 5, further comprising the step of:
<claim-text>processing the encoded and multiplexed audio objects.</claim-text></claim-text></claim>
<claim id="c-en-01-0007" num="0007">
<claim-text>The method of claim 6, wherein in the step of processing the audio objects, the encoded and multiplexed audio objects are transmitted through a digital broadcasting network.</claim-text></claim>
<claim id="c-en-01-0008" num="0008">
<claim-text>An object-based audio service reproducing method, comprising:
<claim-text>de-multiplexing and decoding (S402) object-based audio contents including a plurality of audio objects and a plurality of present audio scenes;</claim-text>
<claim-text>forming (S403) one of said plurality of preset audio scenes, wherein each preset audio scene is based on 3D audio scene information which indicates features of said plurality of audio objects, based on a selection of a user;</claim-text>
<claim-text>mixing (S404) the plurality of audio objects according to the formed preset audio scene; and<!-- EPO <DP n="18"> --></claim-text>
<claim-text>reproducing (S405) the mixed plurality of audio objects according to the formed preset audio scene,</claim-text>
<claim-text>wherein the 3D audio scene information includes at least one of a location, a size, a volume, a direction, and a sound field environment of each of the plurality of audio objects in the 3D space.</claim-text></claim-text></claim>
</claims>
<claims id="claims02" lang="de"><!-- EPO <DP n="19"> -->
<claim id="c-de-01-0001" num="0001">
<claim-text>Eine Vorrichtung (30) zum Liefern eines objektbasierten Audioservice, aufweisend:
<claim-text>Erzeugungsmiftel (32) für voreingestellte Audioszenen zum Erzeugen einer Mehrzahl von voreingestellten Audioszenen, wobei jede voreingestellte Audioszene basiert auf 3D-Audioszeneninformation, die Merkmale einer Mehrzahl von Audioobjekten angeben; und</claim-text>
<claim-text>Codiermittel (33) zum Codieren und Multiplexen der Mehrzahl von Audioobjekten und der Mehrzahl von voreingestellten Audioszenen;</claim-text>
<claim-text>wobei die 3D-Audioszeneninformation zumindest eines einschließt aus: Einer Lokation, einer Größe, einer Lautstärke, einer Richtung und einer Schallfeldumgebung von jedem der Mehrzahl von Audioobjekten im 3D-Raum.</claim-text></claim-text></claim>
<claim id="c-de-01-0002" num="0002">
<claim-text>Die Vorrichtung (30) zum Liefern von objektbasierten Audioservices nach Anspruch 1, ferner aufweisend: Verarbeitungsmittel (34) zum Verarbeiten der codierten und gemultiplexten Audioobjekte.</claim-text></claim>
<claim id="c-de-01-0003" num="0003">
<claim-text>Die Vorrichtung zum Liefern von objektbasierten Audioservices nach Anspruch 1, wobei die Verarbeitungsmittel (34) die codierten und gemultiplexten Audioobjekte an ein Audiowiedergabeterminal (40) über ein digitales Übertragungsnetzwerk übertragen.</claim-text></claim>
<claim id="c-de-01-0004" num="0004">
<claim-text>Eine Vorrichtung (40) zum Liefern objektbasierter Audioservices, aufweisend:
<claim-text>Decodiermittel (42) zum Demultiplexen und Decodieren objektbasierter Audioinhalte einschließend eine Mehrzahl von Audioobjekten und eine Mehrzahl von voreingestellten Audioszenen;</claim-text>
<claim-text>Audioszenen-Bildungsmittel (43) zum Bilden einer der mehrzahl von voreingestellten Audioszenen, wobei jede voreingestellte Audioszene basiert auf 3D-Audioszeneninformationen, die Merkmale der Mehrzahl von Audioobjekten angeben und auf einer Auswahl eines Benutzers basieren;<!-- EPO <DP n="20"> --></claim-text>
<claim-text>Audiosignal-Mischmittel (44) zum Mischen der Mehrzahl von Audioobjekten in Übereinstimmung mit der gebildeten voreingestellten Audioszene; und</claim-text>
<claim-text>Wiedergabemittel (45) zum Wiedergeben der gemischten Mehrzahl von Audioobjekten in Übereinstimmung mit der gebildeten voreingestellten Audioszene,</claim-text>
<claim-text>wobei die 3D-Audioszeneninformationen zumindest eines enschließen aus: Einer Lokation, einer Größe, einer Lautstärke, einer Richtung und einer Schallfeldumgebung von jedem der Mehrzahl von Audioobjekten im 3D-Raum.</claim-text></claim-text></claim>
<claim id="c-de-01-0005" num="0005">
<claim-text>Ein Verfahren zum Liefern eines objektbasierten Audioservices, aufweisend:
<claim-text>Erzeugen (S304) einer Mehrzahl von voreingestellten Audioszenen, wobei jede voreingestellte Audioszene basiert auf 3D-Audioszeneninformationen, die Merkmale einer Mehrzahl von Audioobjekten angeben; und</claim-text>
<claim-text>Codieren und Multiplexen (S305) der Mehrzahl von Audioobjekten und der Mehrzahl von voreingestellten Audioszenen,</claim-text>
<claim-text>wobei die 3D-Audioszeneninformationen zumindest eines enschließen aus: Einer Lokation, einer Größe, einer Lautstärke, einer Richtung und einer Schallfeldumgebung von jeder der Mehrzahl von Audioobjekten im 3D-Raum.</claim-text></claim-text></claim>
<claim id="c-de-01-0006" num="0006">
<claim-text>Das Verfahren nach Anspruch 5, ferner aufweisend den Schritt:
<claim-text>Verarbeiten der codierten und gemultiplexten Audioobjekte. Verarbeiten der codierten und gemultiplexten Audioobjekte</claim-text></claim-text></claim>
<claim id="c-de-01-0007" num="0007">
<claim-text>Das Verfahren nach Anspruch 6, wobei beim Schritt des Verarbeitens der Audioobjekte die codierten und gemultiplexten Audioobjekte übertragen werden durch ein digitales Übertragungsnetzwerk.</claim-text></claim>
<claim id="c-de-01-0008" num="0008">
<claim-text>Ein objektbasiertes Audioservice-Wiedergabeverfahren, aufweisend:
<claim-text>Demultiplexen und Decodieren (5402) objektbasierter Audioinhalte einschließend eine Mehrzahl von Audioobjekten und eine Mehrzahl von voreingestellten Audioszenen;</claim-text>
<claim-text>Bilden (S403) einer der Mehrzahl von Voreingestellten Audioszenen, wobei jede voreingestellte Audioszene basiert auf 3D-Audioszeneninformationen, die Merkmale der Mehrzahl von Audioobjekten basierend auf einer Auswahl eines Benutzers angeben;<!-- EPO <DP n="21"> --></claim-text>
<claim-text>Mischen (S404) der Mehrzahl von Audioobjekten in Übereinstimmung mit der gebildeten voreingestellten Audioszene; und</claim-text>
<claim-text>Wiedergeben (5405) der gemischten Mehrzahl von Audioobjekten in Übereinstimmung mit der gebildeten voreingestellten Audioszene;</claim-text>
<claim-text>wobei die 3D-Audioszeneninformationen zumindest eines emschließen von: Einer Lokation, einer Größe, einer Lautstärke, einer Richtung und einer Schallfeldumgebung jeder der Mehrzahl von Audioobjekten im 3D-Raum.</claim-text></claim-text></claim>
</claims>
<claims id="claims03" lang="fr"><!-- EPO <DP n="22"> -->
<claim id="c-fr-01-0001" num="0001">
<claim-text>Appareil fournissant un service audio basé sur des objets (30), comprenant :
<claim-text>des moyens de génération de scène audio préétablie (32) destinés à générer une pluralité de scènes audio préétablies, dans lequel chaque scène audio préétablie est basée sur des informations de scène audio 3D qui indique des particularités d'une pluralité d'objets audio ; et</claim-text>
<claim-text>des moyens de codage (33) destiné à coder et multiplexer la pluralité d'objets audio et la pluralité de scènes audio préétablies,</claim-text>
<claim-text>dans lequel les informations de scène audio 3D comprennent au moins l'un d'un emplacement, d'une taille, d'un volume, d'une direction et d'un environnement de champ sonore de chacun de la pluralité d'objets audio dans l'espace 3D.</claim-text></claim-text></claim>
<claim id="c-fr-01-0002" num="0002">
<claim-text>Appareil fournissant un service audio basé sur des objets (30) selon la revendication 1, comprenant en outre des moyens de traitement (34) destinés à traiter les objets audio codés et multiplexés.</claim-text></claim>
<claim id="c-fr-01-0003" num="0003">
<claim-text>Appareil fournissant un service audio basé sur des objets selon la revendication 1, dans lequel les moyens de traitement (34) transmettent les objets audio codés et multiplexés à un terminal de reproduction audio (40) par l'intermédiaire d'un réseau de radiodiffusion numérique.<!-- EPO <DP n="23"> --></claim-text></claim>
<claim id="c-fr-01-0004" num="0004">
<claim-text>Appareil de reproduction de service audio basé sur des objets (40) comprenant<br/>
des moyens de décodage (42) destinés à démultiplexer et décoder des contenus audio basés sur les objets comprenant une pluralité d'objets audio et une pluralité de scènes audio préétablies ;<br/>
des moyens de formation de scène audio (43) destinés à former l'une de ladite pluralité de scènes audio préétablies, dans lequel chaque scène audio préétablie est basée sur des informations de scène audio 3D qui indiquent des particularités de ladite pluralité d'objets audio basée sur une sélection d'un utilisateur ;<br/>
des moyens de mixage de signal audio (44) destiné à mixer la pluralité d'objets audio selon la scène audio préétablie formée ; et<br/>
des moyens de reproduction (45) destinés à reproduire la pluralité mixée d'objets audio selon la scène audio préétablie formée,<br/>
dans lequel les informations de scène audio 3D comprennent au moins l'un d'un emplacement, d'une taille, d'un volume, d'une direction et d'un environnement de champ sonore de chacun de la pluralité d'objets audio dans l'espace 3D.</claim-text></claim>
<claim id="c-fr-01-0005" num="0005">
<claim-text>Procédé de fourniture de service audio basé sur des objets, comprenant :
<claim-text>la génération (S304) d'une pluralité de scènes audio préétablies, où chaque scène audio préétablie est basée sur des informations de scène audio 3D qui<!-- EPO <DP n="24"> --> indiquent des particularités d'une pluralité d'objets audio ; et</claim-text>
<claim-text>le codage et le multiplexage (S305) de la pluralité d'objets audio et de la pluralité de scènes audio préétablies,</claim-text>
<claim-text>dans lequel les informations de scène audio 3D comprennent au moins l'un d'un emplacement, d'une taille, d'un volume, d'une direction et d'un environnement de champ sonore de chacun de la pluralité d'objets audio dans l'espace 3D.</claim-text></claim-text></claim>
<claim id="c-fr-01-0006" num="0006">
<claim-text>Procédé selon la revendication 5, comprenant en outre l'étape de :
<claim-text>traitement des objets audio codés et multiplexés.</claim-text></claim-text></claim>
<claim id="c-fr-01-0007" num="0007">
<claim-text>Procédé selon la revendication 6, dans lequel dans l'étape de traitement des objets audio, les objets audio codés et multiplexés sont transmis par l'intermédiaire d'un réseau de radiodiffusion numérique.</claim-text></claim>
<claim id="c-fr-01-0008" num="0008">
<claim-text>Procédé de reproduction de service audio basé sur des objets, comprenant :
<claim-text>le démultiplexage et le décodage (S402) de contenus audio basés sur les objets comprenant une pluralité d'objets audio et une pluralité de scènes audio préétablies ;</claim-text>
<claim-text>la formation (S403) de l'une de ladite pluralité de scènes audio préétablies, où chaque scène audio préétablie est basée sur des informations de scène audio 3D qui indiquent des particularités de ladite<!-- EPO <DP n="25"> --> pluralité d'objets audio basée sur une sélection d'un utilisateur ;</claim-text>
<claim-text>le mixage (S404) de la pluralité d'objets audio selon la scène audio préétablie formée ; et</claim-text>
<claim-text>la reproduction (S405) de la pluralité mixée d'objets audio selon la scène audio préétablie formée,</claim-text>
<claim-text>dans lequel les informations de scène audio 3D comprennent au moins l'un d'un emplacement, d'une taille, d'un volume, d'une direction et d'un environnement de champ sonore de chacun de la pluralité d'objets audio dans l'espace 3D.</claim-text></claim-text></claim>
</claims>
<drawings id="draw" lang="en"><!-- EPO <DP n="26"> -->
<figure id="f0001" num="1"><img id="if0001" file="imgf0001.tif" wi="132" he="222" img-content="drawing" img-format="tif"/></figure><!-- EPO <DP n="27"> -->
<figure id="f0002" num="2"><img id="if0002" file="imgf0002.tif" wi="141" he="233" img-content="drawing" img-format="tif"/></figure><!-- EPO <DP n="28"> -->
<figure id="f0003" num="3"><img id="if0003" file="imgf0003.tif" wi="127" he="171" img-content="drawing" img-format="tif"/></figure><!-- EPO <DP n="29"> -->
<figure id="f0004" num="4"><img id="if0004" file="imgf0004.tif" wi="120" he="149" img-content="drawing" img-format="tif"/></figure>
</drawings>
<ep-reference-list id="ref-list">
<heading id="ref-h0001"><b>REFERENCES CITED IN THE DESCRIPTION</b></heading>
<p id="ref-p0001" num=""><i>This list of references cited by the applicant is for the reader's convenience only. It does not form part of the European patent document. Even though great care has been taken in compiling the references, errors or omissions cannot be excluded and the EPO disclaims all liability in this regard.</i></p>
<heading id="ref-h0002"><b>Patent documents cited in the description</b></heading>
<p id="ref-p0002" num="">
<ul id="ref-ul0001" list-style="bullet">
<li><patcit id="ref-pcit0001" dnum="KR1020040037437"><document-id><country>KR</country><doc-number>1020040037437</doc-number><date>20040507</date></document-id></patcit><crossref idref="pcit0001">[0010]</crossref></li>
<li><patcit id="ref-pcit0002" dnum="US6665A"><document-id><country>US</country><doc-number>6665</doc-number><kind>A</kind></document-id></patcit><crossref idref="pcit0002">[0012]</crossref></li>
<li><patcit id="ref-pcit0003" dnum="US318B1"><document-id><country>US</country><doc-number>318</doc-number><kind>B1</kind></document-id></patcit><crossref idref="pcit0003">[0012]</crossref></li>
</ul></p>
<heading id="ref-h0003"><b>Non-patent literature cited in the description</b></heading>
<p id="ref-p0003" num="">
<ul id="ref-ul0002" list-style="bullet">
<li><nplcit id="ref-ncit0001" npl-type="s"><article><author><name>KALVA H. et al.</name></author><atl>DELIVERING OBJECT-BASED AUDIO-VISUAL SERVICES</atl><serial><sertitle>IEEE TRANSACTIONS ON CONSUMER ELECTRONICS</sertitle><pubdate><sdate>19991101</sdate><edate/></pubdate><vid>45</vid><ino>4</ino></serial><location><pp><ppf>1108</ppf><ppl>1111</ppl></pp></location></article></nplcit><crossref idref="ncit0001">[0011]</crossref></li>
</ul></p>
</ep-reference-list>
</ep-patent-document>
